{"person":{"slug":"jan-peters","name":"Jan Peters","university":"TU Darmstadt","role":"PI","profile_url":"https://aigude.ai/research/jan-peters"},"generated_at":"2026-09-13T23:11:31.696Z","count":261,"bibtex_url":"https://aigude.ai/api/research/feed/jan-peters?format=bibtex","papers":[{"id":"9b7b3643-782d-48a2-843b-cd3179152bf3","title":"The earlier you know, the smoother you act: anticipatory control in solo and dyadic juggling","year":2026,"date":"2026-05-12","venue":"Experimental Brain Research","venue_slug":null,"venue_type":null,"authors":null,"author_count":7,"doi":"10.1007/s00221-026-07311-z","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1007/s00221-026-07311-z","pdf_url":"https://link.springer.com/content/pdf/10.1007/s00221-026-07311-z.pdf","links":null,"keywords":null,"tldr":null,"cited_by_count":0,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"e10e7b73-607e-4543-9e01-78638cba2272","title":"Motion Planning Diffusion: Learning and Adapting Robot Motion Planning with Diffusion Models (Abstract Reprint)","year":2026,"date":"2026-03-14","venue":"AAAI","venue_slug":"aaai","venue_type":null,"authors":null,"author_count":5,"doi":"10.1609/aaai.v40i47.41373","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1609/aaai.v40i47.41373","pdf_url":"https://ojs.aaai.org/index.php/AAAI/article/download/41373/45334","links":null,"keywords":null,"tldr":null,"cited_by_count":0,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T14:20:40.392894+00:00"},{"id":"75d79a10-9d27-4a4e-ba6d-72b7fc09a56d","title":"A Survey on Deep Generative Models for Robot Learning From Multimodal Demonstrations","year":2025,"date":"2025-11-13","venue":"IEEE Transactions on Robotics","venue_slug":null,"venue_type":null,"authors":null,"author_count":8,"doi":"10.1109/tro.2025.3631816","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/tro.2025.3631816","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":1,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"22924310-fab8-470f-9985-3d27699a57b4","title":"The Earlier You Know, the Smoother You Act","year":2025,"date":"2025-11-10","venue":"bioRxiv (Cold Spring Harbor Laboratory)","venue_slug":null,"venue_type":null,"authors":null,"author_count":7,"doi":"10.1101/2025.11.08.687371","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1101/2025.11.08.687371","pdf_url":"https://www.biorxiv.org/content/biorxiv/early/2025/11/10/2025.11.08.687371.full.pdf","links":null,"keywords":null,"tldr":null,"cited_by_count":0,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"ddd97529-dc8f-47fd-8d4d-c5e799099250","title":"Bridge the Gap: Enhancing Quadruped Locomotion with Vertical Ground Perturbations","year":2025,"date":"2025-10-19","venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":8,"doi":"10.1109/iros60139.2025.11247549","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros60139.2025.11247549","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":0,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T14:20:44.062577+00:00"},{"id":"56da0f35-cf68-452c-83b0-0d166ebe49cd","title":"Context-Aware Deep Lagrangian Networks for Model Predictive Control","year":2025,"date":"2025-10-19","venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":3,"doi":"10.1109/iros60139.2025.11246292","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros60139.2025.11246292","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":3,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T14:20:42.714518+00:00"},{"id":"7243f985-63c2-4d17-a744-11364b7cce03","title":"FlowMP: Learning Motion Fields for Robot Planning with Conditional Flow Matching","year":2025,"date":"2025-10-19","venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":6,"doi":"10.1109/iros60139.2025.11246537","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros60139.2025.11246537","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":2,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T14:20:43.841224+00:00"},{"id":"47052f47-8626-441f-aac7-1f5373e47152","title":"Gait in Eight: Efficient On-Robot Learning for Omnidirectional Quadruped Locomotion","year":2025,"date":"2025-10-19","venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":5,"doi":"10.1109/iros60139.2025.11246447","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros60139.2025.11246447","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":1,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T14:20:43.528558+00:00"},{"id":"75e7899f-98fd-4e7d-8933-2e945cb8b0b5","title":"Learning Force Distribution Estimation for the GelSight Mini Optical Tactile Sensor Based on Finite Element Analysis","year":2025,"date":"2025-10-19","venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":5,"doi":"10.1109/iros60139.2025.11246486","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros60139.2025.11246486","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":0,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T14:20:43.018926+00:00"},{"id":"091cbbec-0e77-4d42-b235-1d28bb433355","title":"The Role of Embodiment in Intuitive Whole-Body Teleoperation for Mobile Manipulation","year":2025,"date":"2025-09-30","venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":7,"doi":"10.1109/humanoids65713.2025.11203105","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids65713.2025.11203105","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":0,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T14:20:40.768746+00:00"},{"id":"497f8a41-4997-45c8-9a69-8dad9f79f651","title":"Neuro-Symbolic Imitation Learning: Discovering Symbolic Abstractions for Skill Learning","year":2025,"date":"2025-05-19","venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":3,"doi":"10.1109/icra55743.2025.11127692","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra55743.2025.11127692","pdf_url":"https://arxiv.org/pdf/2503.21406","links":null,"keywords":null,"tldr":null,"cited_by_count":2,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T14:20:42.355404+00:00"},{"id":"262ea606-a2a2-47cd-8230-158de25048ac","title":"Adaptive Control Based Friction Estimation for Tracking Control of Robot Manipulators","year":2025,"date":"2025-01-15","venue":"IEEE Robotics and Automation Letters","venue_slug":null,"venue_type":null,"authors":null,"author_count":4,"doi":"10.1109/lra.2025.3530159","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/lra.2025.3530159","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":5,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"ff6d0782-958a-47f2-82dd-0933d752a691","title":"Fast and Robust Visuomotor Riemannian Flow Matching Policy","year":2025,"date":"2025-01-01","venue":"IEEE Transactions on Robotics","venue_slug":null,"venue_type":null,"authors":null,"author_count":4,"doi":"10.1109/tro.2025.3601293","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/tro.2025.3601293","pdf_url":"http://urn.kb.se/resolve?urn=urn:nbn:se:kth:diva-369864","links":null,"keywords":null,"tldr":null,"cited_by_count":3,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"113b5805-ff7e-4b98-a1e1-ea46275b8513","title":"Learning Multimodal Latent Dynamics for Human–Robot Interaction","year":2025,"date":"2025-01-01","venue":"IEEE Transactions on Robotics","venue_slug":null,"venue_type":null,"authors":null,"author_count":6,"doi":"10.1109/tro.2025.3582829","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/tro.2025.3582829","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":2,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"f6c7ceef-ce66-4f8c-a9fc-5c506d705276","title":"Motion Planning Diffusion: Learning and Adapting Robot Motion Planning With Diffusion Models","year":2025,"date":"2025-01-01","venue":"IEEE Transactions on Robotics","venue_slug":null,"venue_type":null,"authors":null,"author_count":5,"doi":"10.1109/tro.2025.3593109","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/tro.2025.3593109","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":5,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"ad30a189-b9df-48a5-9218-1d0182594496","title":"Reinforcement Learning for Robust Athletic Intel-ligence: Lessons Learned From the Second AI Olympics With RealAIGym Competition","year":2025,"date":"2025-01-01","venue":"IEEE Robotics & Automation Magazine","venue_slug":null,"venue_type":null,"authors":null,"author_count":21,"doi":"10.1109/mra.2025.3631571","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/mra.2025.3631571","pdf_url":"https://doi.org/10.1109/mra.2025.3631571","links":null,"keywords":null,"tldr":null,"cited_by_count":0,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"6e1698cc-4e73-4e38-bbcb-d5edd222804e","title":"Safe Reinforcement Learning on the Constraint Manifold: Theory and Applications","year":2025,"date":"2025-01-01","venue":"IEEE Transactions on Robotics","venue_slug":null,"venue_type":null,"authors":null,"author_count":4,"doi":"10.1109/tro.2025.3567477","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/tro.2025.3567477","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":2,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"34177767-d99f-4c9a-ad51-916564c391a0","title":"Adaptive Q-Network: On-the-fly Target Selection for Deep Reinforcement Learning.","year":2025,"date":null,"venue":"ICLR","venue_slug":"iclr","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/iclr/VincentW0BD25.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"481204fa-7c38-44f9-97ef-abc7e9d40e49","title":"DIME: Diffusion-Based Maximum Entropy Reinforcement Learning.","year":2025,"date":null,"venue":"ICML","venue_slug":"icml","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/icml/CelikLBLP0CN25.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"efd6e706-c01b-4a24-b6fa-d9db5dacd1ed","title":"Inverse decision-making using neural amortized Bayesian actors.","year":2025,"date":null,"venue":"ICLR","venue_slug":"iclr","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/iclr/StraubN0R25.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"88563f81-e5cc-4448-9be9-b722be256980","title":"Maximum Total Correlation Reinforcement Learning.","year":2025,"date":null,"venue":"ICML","venue_slug":"icml","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/icml/YouLL0A25.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"396483db-4eaf-4fc1-8c0e-25132db18bed","title":"Noise-conditioned Energy-based Annealed Rewards (NEAR): A Generative Framework for Imitation Learning from Observation.","year":2025,"date":null,"venue":"ICLR","venue_slug":"iclr","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/iclr/DiwanUK025.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"7702cc3f-8afa-4eeb-bc02-87a5c6356664","title":"QueryCAD: Grounded Question Answering for CAD Models.","year":2025,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra55743.2025.11128709","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra55743.2025.11128709","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"7b7bb357-f4a1-40dc-b5a2-fcc391351b11","title":"Exciting Action: Investigating Efficient Exploration for Learning Musculoskeletal Humanoid Locomotion","year":2024,"date":"2024-11-22","venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":5,"doi":"10.1109/humanoids58906.2024.10769835","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids58906.2024.10769835","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":3,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T14:20:45.846739+00:00"},{"id":"4ac88e03-0055-480d-b1bc-b964374807c4","title":"Unsupervised Skill Discovery for Robotic Manipulation through Automatic Task Generation","year":2024,"date":"2024-11-22","venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":4,"doi":"10.1109/humanoids58906.2024.10769879","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids58906.2024.10769879","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":3,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T14:20:46.098085+00:00"},{"id":"84fb229e-437a-4594-9aaa-3d976d41d552","title":"A Holistic Concept on AI Assistance for Robot-Supported Reconnaissance and Mitigation of Acute Radiation Hazard Situations","year":2024,"date":"2024-11-12","venue":null,"venue_slug":null,"venue_type":null,"authors":null,"author_count":14,"doi":"10.1109/ssrr62954.2024.10770059","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/ssrr62954.2024.10770059","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":2,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"db8d4c75-55a8-4ac5-9b0d-9fbaa372e1c0","title":"Beyond the Cascade: Juggling Vanilla Siteswap Patterns","year":2024,"date":"2024-10-14","venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":3,"doi":"10.1109/iros58592.2024.10801588","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros58592.2024.10801588","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":0,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T14:20:48.53849+00:00"},{"id":"9c6e867e-dd9f-4c07-b18f-0ec268e507a5","title":"Extended Tree Search for Robot Task and Motion Planning","year":2024,"date":"2024-10-14","venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":3,"doi":"10.1109/iros58592.2024.10802552","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros58592.2024.10802552","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":4,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T14:20:48.777389+00:00"},{"id":"6b2d4276-a88e-4eb7-bda8-fa9806c04d84","title":"Safe and Efficient Path Planning Under Uncertainty via Deep Collision Probability Fields","year":2024,"date":"2024-09-10","venue":"IEEE Robotics and Automation Letters","venue_slug":null,"venue_type":null,"authors":null,"author_count":6,"doi":"10.1109/lra.2024.3457208","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/lra.2024.3457208","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":3,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"972a9958-b1a0-4ffe-82f9-d3ae3f472faf","title":"A Comparison of Imitation Learning Algorithms for Bimanual Manipulation","year":2024,"date":"2024-08-19","venue":"IEEE Robotics and Automation Letters","venue_slug":null,"venue_type":null,"authors":null,"author_count":7,"doi":"10.1109/lra.2024.3445630","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/lra.2024.3445630","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":17,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"f00fabd5-70db-4296-a680-fb9406c37a35","title":"Advancing Sustainable Construction: Discrete Modular Systems &amp; Robotic Assembly","year":2024,"date":"2024-08-04","venue":"Sustainability","venue_slug":null,"venue_type":null,"authors":null,"author_count":8,"doi":"10.3390/su16156678","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.3390/su16156678","pdf_url":"https://www.mdpi.com/2071-1050/16/15/6678/pdf?version=1722947710","links":null,"keywords":null,"tldr":null,"cited_by_count":3,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"47068049-a83f-4dca-9a37-879e95e01987","title":"ROSCOM: Robust Safe Reinforcement Learning on Stochastic Constraint Manifolds","year":2024,"date":"2024-07-31","venue":"IEEE Transactions on Automation Science and Engineering","venue_slug":null,"venue_type":null,"authors":null,"author_count":6,"doi":"10.1109/tase.2024.3431530","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/tase.2024.3431530","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":1,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"92d83823-6603-43b1-b7b7-55c02a943eab","title":"Learning Energy-Efficient Trajectory Planning for Robotic Manipulators Using Bayesian Optimization","year":2024,"date":"2024-06-25","venue":null,"venue_slug":null,"venue_type":null,"authors":null,"author_count":4,"doi":"10.23919/ecc64448.2024.10590756","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.23919/ecc64448.2024.10590756","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":7,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"ab722b23-d9ed-4a51-9ccf-69a9ba950178","title":"What Matters for Active Texture Recognition With Vision-Based Tactile Sensors","year":2024,"date":"2024-05-13","venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":9,"doi":"10.1109/icra57147.2024.10610274","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra57147.2024.10610274","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":7,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T14:20:47.868755+00:00"},{"id":"ba5364b6-812e-4cc4-a926-23f36992458c","title":"MoVEInt: Mixture of Variational Experts for Learning Human–Robot Interactions From Demonstrations","year":2024,"date":"2024-05-01","venue":"IEEE Robotics and Automation Letters","venue_slug":null,"venue_type":null,"authors":null,"author_count":6,"doi":"10.1109/lra.2024.3396074","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/lra.2024.3396074","pdf_url":"https://arxiv.org/pdf/2407.07636","links":null,"keywords":null,"tldr":null,"cited_by_count":8,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"74486e04-6dc9-40ed-a46d-f851f3b1c12b","title":"Autonomous underwater vehicle link alignment control in unknown environments using reinforcement learning","year":2024,"date":"2024-04-23","venue":"Journal of Field Robotics","venue_slug":null,"venue_type":null,"authors":null,"author_count":8,"doi":"10.1002/rob.22348","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1002/rob.22348","pdf_url":"https://onlinelibrary.wiley.com/doi/pdfdirect/10.1002/rob.22348","links":null,"keywords":null,"tldr":null,"cited_by_count":4,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"add9aea6-a769-4127-b3ba-7be323fbd77d","title":"On the Benefit of Optimal Transport for Curriculum Reinforcement Learning","year":2024,"date":"2024-04-16","venue":"IEEE Transactions on Pattern Analysis and Machine Intelligence","venue_slug":null,"venue_type":null,"authors":null,"author_count":4,"doi":"10.1109/tpami.2024.3390051","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/tpami.2024.3390051","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":7,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"fbabee66-2f30-45f2-a467-f77a5b84997e","title":"Kinematically Constrained Human-like Bimanual Robot-to-Human Handovers","year":2024,"date":"2024-03-10","venue":null,"venue_slug":null,"venue_type":null,"authors":null,"author_count":7,"doi":"10.1145/3610978.3640670","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1145/3610978.3640670","pdf_url":"https://arxiv.org/pdf/2402.14525","links":null,"keywords":null,"tldr":null,"cited_by_count":2,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"9ca210bf-6c97-43aa-ad6e-0bb2430bc71a","title":"Transition State Clustering for Interaction Segmentation and Learning","year":2024,"date":"2024-03-10","venue":null,"venue_slug":null,"venue_type":null,"authors":null,"author_count":7,"doi":"10.1145/3610978.3640738","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1145/3610978.3640738","pdf_url":"https://arxiv.org/pdf/2402.14548","links":null,"keywords":null,"tldr":null,"cited_by_count":1,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"b31c8ccd-e574-4c54-9fb0-ba85bd3ca3b8","title":"Sharing Knowledge in Multi-Task Deep Reinforcement Learning","year":2024,"date":"2024-01-17","venue":"arXiv (Cornell University)","venue_slug":null,"venue_type":null,"authors":null,"author_count":5,"doi":"10.48550/arxiv.2401.09561","arxiv_id":null,"openreview_id":null,"landing_url":"http://arxiv.org/abs/2401.09561","pdf_url":"https://arxiv.org/pdf/2401.09561","links":null,"keywords":null,"tldr":null,"cited_by_count":62,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"bce7797c-9e63-4fcb-a7a1-6fbd5aa06856","title":"Evetac: An Event-Based Optical Tactile Sensor for Robotic Manipulation","year":2024,"date":"2024-01-01","venue":"IEEE Transactions on Robotics","venue_slug":null,"venue_type":null,"authors":null,"author_count":5,"doi":"10.1109/tro.2024.3428430","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/tro.2024.3428430","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":45,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"ffe459ca-bcfe-4659-ac43-3e92397877d8","title":"Grip Stabilization through Independent Finger Tactile Feedback Control","year":2024,"date":"2024-01-01","venue":"TUbilio (Technical University of Darmstadt)","venue_slug":null,"venue_type":null,"authors":null,"author_count":3,"doi":"10.26083/tuprints-00016296","arxiv_id":null,"openreview_id":null,"landing_url":"http://tubiblio.ulb.tu-darmstadt.de/142379/","pdf_url":"https://tuprints.ulb.tu-darmstadt.de/16296/1/sensors-20-01748.pdf","links":null,"keywords":null,"tldr":null,"cited_by_count":0,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"f7330bc2-44f1-4c67-a96f-6b53d1098f40","title":"Hierarchical Tactile-Based Control Decomposition of Dexterous In-Hand Manipulation Tasks","year":2024,"date":"2024-01-01","venue":"TUbilio (Technical University of Darmstadt)","venue_slug":null,"venue_type":null,"authors":null,"author_count":3,"doi":"10.26083/tuprints-00016159","arxiv_id":null,"openreview_id":null,"landing_url":"http://tubiblio.ulb.tu-darmstadt.de/143655/","pdf_url":"https://tuprints.ulb.tu-darmstadt.de/16159/13/frobt-07-521448.pdf","links":null,"keywords":null,"tldr":null,"cited_by_count":0,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"799db5b0-793c-43f6-a40a-9d2caa234906","title":"Learning Sequential Force Interaction Skills","year":2024,"date":"2024-01-01","venue":"TUbilio (Technical University of Darmstadt)","venue_slug":null,"venue_type":null,"authors":null,"author_count":4,"doi":"10.26083/tuprints-00016992","arxiv_id":null,"openreview_id":null,"landing_url":"https://tuprints.ulb.tu-darmstadt.de/16992","pdf_url":"https://tuprints.ulb.tu-darmstadt.de/16992","links":null,"keywords":null,"tldr":null,"cited_by_count":0,"lab_submitted":false,"register_url":null,"updated_at":"2026-06-24T07:26:35.218175+00:00"},{"id":"af70206a-b4d4-4de9-9f2b-5e42d23642c1","title":"A Retrospective on the Robot Air Hockey Challenge: Benchmarking Robust, Reliable, and Safe Learning Techniques for Real-world Robotics.","year":2024,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/LiuGFGCBJMCOOZL24.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"8eb8c718-e0b4-4369-885c-7b6fbd008706","title":"Bridging the gap between Learning-to-plan, Motion Primitives and Safe Reinforcement Learning.","year":2024,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/KickiTLG0W24.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"1f143c94-9c0e-4fb7-88d8-6572d0f12490","title":"CrossQ: Batch Normalization in Deep Reinforcement Learning for Greater Sample Efficiency and Simplicity.","year":2024,"date":null,"venue":"ICLR","venue_slug":"iclr","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/iclr/0001PBAAB024.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"b078396b-c044-440e-8c8c-80765fd146c9","title":"Domain Randomization via Entropy Maximization.","year":2024,"date":null,"venue":"ICLR","venue_slug":"iclr","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/iclr/TiboniK0TDC24.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"5b272ad0-d4bb-45c7-8bc8-26bb6fe57e31","title":"Handling Long-Term Safety and Uncertainty in Safe Reinforcement Learning.","year":2024,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/GunsterL0T24.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"ac5f7f0b-e635-4b2a-a589-1948345a766e","title":"Multi-Task Reinforcement Learning with Mixture of Orthogonal Experts.","year":2024,"date":null,"venue":"ICLR","venue_slug":"iclr","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/iclr/Hendawy0D24.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"580c0040-1edf-4c51-9483-0038c2641f52","title":"One Policy to Run Them All: an End-to-end Learning Approach to Multi-Embodiment Locomotion.","year":2024,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/BohlingerCKKW0T24.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"5a994e0d-0cdd-412e-9085-2ed47d20e609","title":"Open X-Embodiment: Robotic Learning Datasets and RT-X Models : Open X-Embodiment Collaboration.","year":2024,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra57147.2024.10611477","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra57147.2024.10611477","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"bdd4d73d-8efd-4186-8da3-c38dc3403db7","title":"Parameterized Projected Bellman Operator.","year":2024,"date":null,"venue":"AAAI","venue_slug":"aaai","venue_type":null,"authors":null,"author_count":null,"doi":"10.1609/aaai.v38i14.29465","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1609/aaai.v38i14.29465","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"1f67aee6-af3a-4ebc-97f7-51dca0da1e6c","title":"Peer Learning: Learning Complex Policies in Groups from Scratch via Action Recommendations.","year":2024,"date":null,"venue":"AAAI","venue_slug":"aaai","venue_type":null,"authors":null,"author_count":null,"doi":"10.1609/aaai.v38i10.29061","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1609/aaai.v38i10.29061","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"37127172-e9d5-48d1-a0b9-7990ecfe3986","title":"PianoMime: Learning a Generalist, Dexterous Piano Player from Internet Demonstrations.","year":2024,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/QianUZ024.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"a8f8adf8-957f-4c6a-b262-1d2a2d871c05","title":"Reinforcement Learning for Athletic Intelligence: Lessons from the 1st &#34;AI Olympics with RealAIGym&#34; Competition.","year":2024,"date":null,"venue":"IJCAI","venue_slug":"ijcai","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/ijcai/WiebeTLZVVGCRSZ24.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"7cffd128-371d-41dd-a160-164a9934a27d","title":"Robust Adversarial Reinforcement Learning via Bounded Rationality Curricula.","year":2024,"date":null,"venue":"ICLR","venue_slug":"iclr","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/iclr/ReddiT0CD24.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"bf57618a-1dad-4b89-8681-a2b2face6948","title":"Structure-Aware E(3)-Invariant Molecular Conformer Aggregation Networks.","year":2024,"date":null,"venue":"ICML","venue_slug":"icml","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/icml/NguyenL00NH0SZN24.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"70bda4e6-bfd0-49c7-82a8-8eb20cc7ef37","title":"Time-Efficient Reinforcement Learning with Stochastic Stateful Policies.","year":2024,"date":null,"venue":"ICLR","venue_slug":"iclr","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/iclr/Al-HafezZ0T24.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"07587d01-6216-42f5-a48f-b263c1033886","title":"Zero-Shot Transfer of a Tactile-based Continuous Force Control Policy from Simulation to Robot.","year":2024,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros58592.2024.10802386","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros58592.2024.10802386","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"0c55e1fe-2a69-4083-87aa-69fbbcbb9afc","title":"Accelerating Motion Planning via Optimal Transport.","year":2023,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/0001CBP23.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"379e4fa2-b545-45ae-b49f-ca35c0348949","title":"Clustering of Motion Trajectories by a Distance Measure Based on Semantic Features.","year":2023,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids57100.2023.10375228","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids57100.2023.10375228","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"f108425c-755b-43d6-bd14-060a5359fd8d","title":"Diminishing Return of Value Expansion Methods in Model-Based Reinforcement Learning.","year":2023,"date":null,"venue":"ICLR","venue_slug":"iclr","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/iclr/PalenicekLC023.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"7053cc11-e727-4688-86c5-072938b597cc","title":"Hierarchical Policy Blending as Inference for Reactive Robot Control.","year":2023,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra48891.2023.10161374","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra48891.2023.10161374","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"9a21a3a6-1f68-4353-bddf-f3a3542f09bc","title":"Improved Algorithms for Stochastic Linear Bandits Using Tail Bounds for Martingale Mixtures.","year":2023,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/FlynnRKP23.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"658e4fe0-8f98-49eb-920a-623ddc456fb2","title":"LS-IQ: Implicit Reward Regularization for Inverse Reinforcement Learning.","year":2023,"date":null,"venue":"ICLR","venue_slug":"iclr","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/iclr/Al-HafezTAZ023.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"7a3e674a-0a40-4147-9ccd-84ffdc61e55f","title":"Model-Based Uncertainty in Value Functions.","year":2023,"date":null,"venue":"AISTATS","venue_slug":"aistats","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/aistats/LuisBVB023.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"d991ffe9-59ef-498d-bd99-d0f40a4fad16","title":"Motion Planning Diffusion: Learning and Planning of Robot Motions with Diffusion Models.","year":2023,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros55552.2023.10342382","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros55552.2023.10342382","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"29159505-6a13-4f12-be5d-edaf0e34bd59","title":"Placing by Touching: An Empirical Study on the Importance of Tactile Sensing for Precise Object Placing.","year":2023,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros55552.2023.10342340","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros55552.2023.10342340","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"b59db895-ad31-4ffe-b941-13eaa40bc430","title":"Pseudo-Likelihood Inference.","year":2023,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/GrunerBMPP23.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"a93f8a65-c233-4898-ba50-95be0eea7fb6","title":"Safe Reinforcement Learning of Dynamic High-Dimensional Robotic Tasks: Navigation, Manipulation, Interaction.","year":2023,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra48891.2023.10161548","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra48891.2023.10161548","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"0a60572e-abce-4912-8408-626bbfbaffd8","title":"SE(3)-DiffusionFields: Learning smooth cost functions for joint grasp and motion optimization through diffusion.","year":2023,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra48891.2023.10161569","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra48891.2023.10161569","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"232b2f98-9a93-4500-9bac-5afb20ae0b61","title":"Start State Selection for Control Policy Learning from Optimal Trajectories.","year":2023,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra48891.2023.10160978","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra48891.2023.10160978","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"c55985b5-c86f-45f9-a8b9-6563b98bd786","title":"Active Exploration for Robotic Manipulation.","year":2022,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros47612.2022.9982061","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros47612.2022.9982061","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"3b82eeeb-8af3-4f8a-8f94-ff569f439b8a","title":"Adapting Object-Centric Probabilistic Movement Primitives with Residual Reinforcement Learning.","year":2022,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids53995.2022.10000148","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids53995.2022.10000148","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"ea2f9948-bb14-4b05-9636-9ef88ca61616","title":"Boosted Curriculum Reinforcement Learning.","year":2022,"date":null,"venue":"ICLR","venue_slug":"iclr","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/iclr/KlinkD0P22.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"6f9b3c6a-7b48-4899-bf45-7cae7551fcce","title":"Controlling the Cascade: Kinematic Planning for N-ball Toss Juggling.","year":2022,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros47612.2022.9981678","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros47612.2022.9981678","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"0007a003-6c12-4174-aece-f280c6013cdb","title":"Curriculum Reinforcement Learning via Constrained Optimal Transport.","year":2022,"date":null,"venue":"ICML","venue_slug":"icml","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/icml/KlinkYD0P22.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"2462850d-04a3-429e-8ffb-ae0d010882cc","title":"Dimensionality Reduction and Prioritized Exploration for Policy Search.","year":2022,"date":null,"venue":"AISTATS","venue_slug":"aistats","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/aistats/MemmelLT022.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"715df005-5649-45e7-bcbb-8c1276d9ccd7","title":"Graph-based Reinforcement Learning meets Mixed Integer Programs: An application to 3D robot assembly discovery.","year":2022,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros47612.2022.9981784","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros47612.2022.9981784","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"768c1c73-d6cb-41cf-afe5-5e27f0730d22","title":"Improving Sample Efficiency of Example-Guided Deep Reinforcement Learning for Bipedal Walking.","year":2022,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids53995.2022.10000068","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids53995.2022.10000068","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"e00d2523-4756-4c8b-bee2-04a9de0c60df","title":"Inferring Smooth Control: Monte Carlo Posterior Policy Iteration with Gaussian Processes.","year":2022,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/Watson022.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"2e23712a-e63f-485f-bec5-c96c291ac22b","title":"Information-Theoretic Safe Exploration with Gaussian Processes.","year":2022,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/BotteroLVB022.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"63252238-5776-4e1e-a919-68719bd387c5","title":"Integrated Bi-Manual Motion Generation and Control shaped for Probabilistic Movement Primitives.","year":2022,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids53995.2022.10000149","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids53995.2022.10000149","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"322ee7ca-f3eb-445d-b148-0b765e71790a","title":"Learning Implicit Priors for Motion Optimization.","year":2022,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros47612.2022.9981264","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros47612.2022.9981264","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"30b8634e-8376-49cd-ac67-22eea076aca2","title":"MILD: Multimodal Interactive Latent Dynamics for Learning Human-Robot Interaction.","year":2022,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids53995.2022.10000239","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids53995.2022.10000239","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"a6fa7403-738b-4145-9918-aead23016e9b","title":"Regularized Deep Signed Distance Fields for Reactive Motion Generation.","year":2022,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros47612.2022.9981456","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros47612.2022.9981456","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"967686f1-bafa-45a5-a6b7-6c1efcccbf43","title":"A Variational Infinite Mixture for Probabilistic Inverse Dynamics Learning.","year":2021,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra48506.2021.9560832","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra48506.2021.9560832","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"77719f89-a298-489f-96ae-9fcb1a447b7c","title":"Composable Energy Policies for Reactive Motion Generation and Reinforcement Learning.","year":2021,"date":null,"venue":"RSS","venue_slug":"rss","venue_type":null,"authors":null,"author_count":null,"doi":"10.15607/rss.2021.xvii.052","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.15607/rss.2021.xvii.052","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"c8b4d152-30ff-4781-839d-2f4ce22fde8e","title":"Contextual Latent-Movements Off-Policy Optimization for Robotic Manipulation Skills.","year":2021,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra48506.2021.9561870","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra48506.2021.9561870","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"54755de5-663c-41ae-91e6-0ec0319899ed","title":"Convex Regularization in Monte-Carlo Tree Search.","year":2021,"date":null,"venue":"ICML","venue_slug":"icml","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/icml/DamD0P21.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"4e7affaa-d546-45ea-837b-510b2f12ddbb","title":"Differentiable Physics Models for Real-world Offline Model-based Reinforcement Learning.","year":2021,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra48506.2021.9561805","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra48506.2021.9561805","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"ba0dd551-599d-4573-aa33-fd49a6ef41e5","title":"Directed Acyclic Graph Neural Network for Human Motion Prediction.","year":2021,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra48506.2021.9561540","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra48506.2021.9561540","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"1ef9c193-6590-4252-9676-2c6cd3da1131","title":"Efficient and Reactive Planning for High Speed Robot Air Hockey.","year":2021,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros51168.2021.9636263","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros51168.2021.9636263","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"f3595c9f-a795-4c32-9662-a67dc8fe6c1f","title":"Latent Derivative Bayesian Last Layer Networks.","year":2021,"date":null,"venue":"AISTATS","venue_slug":"aistats","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/aistats/WatsonLKP021.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"6a8a3e38-cb5f-4095-a863-8f08ba5677cb","title":"Learn2Assemble with Structured Representations and Search for Robotic Architectural Construction.","year":2021,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/FunkCB021.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"f8f77523-9641-4378-8643-fa8b5d430db2","title":"Learning Human-like Hand Reaching for Human-Robot Handshaking.","year":2021,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra48506.2021.9560746","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra48506.2021.9560746","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"fdec368b-f740-43e0-8270-54104c6e0dbd","title":"Model Predictive Actor-Critic: Accelerating Robot Skill Acquisition with Deep Reinforcement Learning.","year":2021,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra48506.2021.9561298","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra48506.2021.9561298","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"672bedbd-1f53-4e82-b4f1-a16a1378de6f","title":"Neural Posterior Domain Randomization.","year":2021,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/MuratoreGWBG021.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"1d12dacf-3e9a-4cc0-b32b-1119d7ab4d87","title":"Real Robot Challenge: A Robotics Competition in the Cloud.","year":2021,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/BauerWWBSGSAJBA21.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"9801d66b-6fc7-4cd9-9410-2cf814b7f773","title":"Robot Reinforcement Learning on the Constraint Manifold.","year":2021,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/LiuTB021.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"a62d59a0-cbb5-40b1-be83-d789a5aa851e","title":"Robust Value Iteration for Continuous Control Tasks.","year":2021,"date":null,"venue":"RSS","venue_slug":"rss","venue_type":null,"authors":null,"author_count":null,"doi":"10.15607/rss.2021.xvii.007","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.15607/rss.2021.xvii.007","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"a3ff15e6-1fe2-46ef-a446-d904515f5840","title":"Value Iteration in Continuous Actions, States and Time.","year":2021,"date":null,"venue":"ICML","venue_slug":"icml","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/icml/LutterM0FG21.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"18ab5656-df41-4514-b8d6-db8fae4e7502","title":"A Nonparametric Off-Policy Policy Gradient.","year":2020,"date":null,"venue":"AISTATS","venue_slug":"aistats","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/aistats/TosattoCA020.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"8cb2d177-74ea-498c-8523-217ef4cf1e3e","title":"Bayesian Online Prediction of Change Points.","year":2020,"date":null,"venue":"UAI","venue_slug":"uai","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/uai/Agudelo-EspanaG20.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"a04c1961-6676-4554-9e9e-be371c57ecc2","title":"Deep Adversarial Reinforcement Learning for Object Disentangling.","year":2020,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros45743.2020.9341578","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros45743.2020.9341578","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"ca0d6237-fe45-4e74-9ace-7456c1bd8a8c","title":"Generalized Mean Estimation in Monte-Carlo Tree Search.","year":2020,"date":null,"venue":"IJCAI","venue_slug":"ijcai","venue_type":null,"authors":null,"author_count":null,"doi":"10.24963/ijcai.2020/332","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.24963/ijcai.2020/332","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"e35fea5b-9a82-41ee-9dfa-bd36c8e7142c","title":"High Acceleration Reinforcement Learning for Real-World Juggling with Binary Rewards.","year":2020,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/PloegerL020.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"b4e67c67-5922-4780-b362-e5467848fe3a","title":"ImitationFlow: Learning Deep Stable Stochastic Dynamic Systems by Normalizing Flows.","year":2020,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros45743.2020.9341035","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros45743.2020.9341035","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"bcf2dbb5-9d92-4cbe-a017-609cd3be8e5b","title":"Learning Control Policies from Optimal Trajectories.","year":2020,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra40945.2020.9196791","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra40945.2020.9196791","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"1845da13-0973-4124-9e6f-af7e2352c129","title":"Learning Hierarchical Acquisition Functions for Bayesian Optimization.","year":2020,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros45743.2020.9341335","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros45743.2020.9341335","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"1d456818-96cb-417c-814c-95c66c7ead6a","title":"Model-Based Quality-Diversity Search for Efficient Robot Learning.","year":2020,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros45743.2020.9340794","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros45743.2020.9340794","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"343f76ac-8b56-486b-81dc-b16370c9bef3","title":"Redundancy resolution under hard joint constraints: a generalized approach to rank updates.","year":2020,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros45743.2020.9341581","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros45743.2020.9341581","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"98e7ed00-1392-4ea3-bd36-c535b6f9541a","title":"Self-Paced Deep Reinforcement Learning.","year":2020,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/KlinkD0P20.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"4df14c5b-8735-4516-94e7-d5242abe9c88","title":"Sharing Knowledge in Multi-Task Deep Reinforcement Learning.","year":2020,"date":null,"venue":"ICLR","venue_slug":"iclr","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/iclr/DEramoTBR020.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"d26dabac-a382-4fc6-ade5-e14daa427136","title":"Underactuated Waypoint Trajectory Optimization for Light Painting Photography.","year":2020,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra40945.2020.9196516","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra40945.2020.9196516","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"e712cec6-f559-4cae-bf2f-2c622c85e4e6","title":"Building a Library of Tactile Skills Based on FingerVision.","year":2019,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids43949.2019.9035000","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids43949.2019.9035000","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"ad6f72ba-044b-4213-98ce-32af97da7294","title":"Chance-Constrained Trajectory Optimization for Non-linear Systems with Unknown Stochastic Dynamics.","year":2019,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros40897.2019.8967794","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros40897.2019.8967794","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"2188fa54-9813-4cc6-bf27-3123caba316e","title":"Deep Lagrangian Networks for end-to-end learning of energy-based control for under-actuated systems.","year":2019,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros40897.2019.8968268","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros40897.2019.8968268","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"77d97f5d-aedf-4490-b6a8-30e3cfb783f5","title":"Deep Lagrangian Networks: Using Physics as Model Prior for Deep Learning.","year":2019,"date":null,"venue":"ICLR","venue_slug":"iclr","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/iclr/LutterRP19.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"f92c7988-7c17-4738-b567-ca50a9b6b87f","title":"Entropic Risk Measure in Policy Search.","year":2019,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros40897.2019.8967699","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros40897.2019.8967699","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"692f389f-6bcb-44ef-9e31-0d16f4a80d6e","title":"Experience Reuse with Probabilistic Movement Primitives.","year":2019,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros40897.2019.8968545","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros40897.2019.8968545","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"0f562159-bb4b-4315-9480-4b54b66dae27","title":"Generalized Multiple Correlation Coefficient as a Similarity Measurement between Trajectories.","year":2019,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros40897.2019.8967884","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros40897.2019.8967884","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"6a505d1f-f295-4a8e-a8cc-a70cdca27af3","title":"HJB Optimal Feedback Control with Deep Differential Value Functions and Action Constraints.","year":2019,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/LutterBLC019.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"ea24d1d3-9d67-456e-b815-6abadb77213b","title":"Local Online Motor Babbling: Learning Motor Abundance of a Musculoskeletal Robot Arm*.","year":2019,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros40897.2019.8967791","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros40897.2019.8967791","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"ae24b7c3-7025-44d2-8f7d-52f8b58866f1","title":"Multimodal Uncertainty Reduction for Intention Recognition in Human-Robot Interaction.","year":2019,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros40897.2019.8968171","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros40897.2019.8968171","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"5e7390ba-b205-4469-946b-90f34a61f5c4","title":"Projections for Approximate Policy Iteration Algorithms.","year":2019,"date":null,"venue":"ICML","venue_slug":"icml","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/icml/AkrourP0N19.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"847bdd7b-50e9-4a2d-b8ca-34c32151842c","title":"Receding Horizon Curiosity.","year":2019,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/SchultheisBA019.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"d29487c4-a82f-4951-a68b-880dcacafb5e","title":"Reinforcement Learning of Trajectory Distributions: Applications in Assisted Teleoperation and Motion Planning.","year":2019,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros40897.2019.8967856","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros40897.2019.8967856","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"dd3d9c14-19dc-45aa-b945-2985fc5bcac5","title":"Self-Paced Contextual Reinforcement Learning.","year":2019,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/KlinkAB019.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"78ad6557-a448-4920-8527-ff49d221faea","title":"Stochastic Optimal Control as Approximate Input Inference.","year":2019,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/WatsonA019.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"ad5b009d-7490-46f1-bc56-8c2f444abe43","title":"Switching Linear Dynamics for Variational Bayes Filtering.","year":2019,"date":null,"venue":"ICML","venue_slug":"icml","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/icml/Becker-Ehmck0S19.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"72860161-ea86-42d7-afca-cb8f40b9b916","title":"Domain Randomization for Simulation-Based Policy Optimization with Transferability Assessment.","year":2018,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/MuratoreTG018.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"9a65ab62-99b2-4d2e-83bb-3af3afd52632","title":"Inducing Probabilistic Context-Free Grammars for the Sequencing of Movement Primitives.","year":2018,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2018.8460190","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2018.8460190","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"c74b91a5-cd27-49ce-8bbe-9ab46f8036b0","title":"Learning Coupled Forward-Inverse Models with Combined Prediction Errors.","year":2018,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2018.8460675","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2018.8460675","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"9ef60866-3b0d-494f-9910-90b645e23e0d","title":"Online Learning of an Open-Ended Skill Library for Collaborative Tasks.","year":2018,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2018.8625031","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2018.8625031","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"8c67855f-c085-4a85-9b2c-e6a8a4f53ace","title":"PIPPS: Flexible Model-Based Policy Search Robust to the Curse of Chaos.","year":2018,"date":null,"venue":"ICML","venue_slug":"icml","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/icml/ParmasR0D18.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"621a20f0-dfac-4b93-a659-7e957dfa5f0b","title":"Regularizing Reinforcement Learning with State Abstraction.","year":2018,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2018.8594201","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2018.8594201","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"11a85d92-cded-42cc-aa40-bd9d3cb1bc44","title":"Sample and Feedback Efficient Hierarchical Reinforcement Learning from Human Preferences.","year":2018,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2018.8460907","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2018.8460907","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"5edd29c3-7073-4104-9607-0cb0e70705a9","title":"Utilizing Human Feedback in POMDP Execution and Specification.","year":2018,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2018.8625022","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2018.8625022","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"18d617f1-9273-43d7-9342-291ac4898802","title":"A comparison of distance measures for learning nonparametric motor skill libraries.","year":2017,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2017.8246937","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2017.8246937","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"42310874-bc4c-41cf-8a58-e024f9c31e9b","title":"A learning-based shared control architecture for interactive task execution.","year":2017,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2017.7989042","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2017.7989042","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"985af9f6-70f6-468c-96f7-b53bca32d3c2","title":"Active Incremental Learning of Robot Movement Primitives.","year":2017,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/MaedaEOB017.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"f62da587-dcaa-4ba6-9ff4-2fc997ea873f","title":"Context-driven movement primitive adaptation.","year":2017,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2017.7989396","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2017.7989396","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"062d72a0-ac48-49e0-be0f-1f190c258d01","title":"Efficient online adaptation with stochastic recurrent neural networks.","year":2017,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2017.8246875","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2017.8246875","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"b5d61bae-5774-4474-883e-5e5dbeef41ca","title":"Empowered skills.","year":2017,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2017.7989760","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2017.7989760","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"598c7548-584b-4f5b-bcf0-f581a6cd5e8b","title":"Goal-driven dimensionality reduction for reinforcement learning.","year":2017,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2017.8206334","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2017.8206334","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"2a41db79-ed43-4c6e-b582-292bd18cc078","title":"Hybrid control trajectory optimization under uncertainty.","year":2017,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2017.8206460","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2017.8206460","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"94ca6b9e-9eb6-4f75-960f-6c1328c29ea7","title":"Layered direct policy search for learning hierarchical skills.","year":2017,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2017.7989761","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2017.7989761","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"b506bc02-b883-49d1-aeda-fe86002eca1c","title":"Learning inverse dynamics models in O(n) time with LSTM networks.","year":2017,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2017.8246965","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2017.8246965","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"09d2bd6a-3ed2-43dc-babd-9d609b7089b6","title":"Local Bayesian Optimization of Motor Skills.","year":2017,"date":null,"venue":"ICML","venue_slug":"icml","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/icml/AkrourS0N17.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"f706d411-d590-4365-8830-2ccbf27a62d0","title":"Online Learning with Stochastic Recurrent Neural Networks using Intrinsic Motivation Signals.","year":2017,"date":null,"venue":"CoRL","venue_slug":"corl","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/corl/Tanneberg0R17.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"b7ed7bc1-9be0-475b-a22c-9546e0ded812","title":"Policy Search with High-Dimensional Context Variables.","year":2017,"date":null,"venue":"AAAI","venue_slug":"aaai","venue_type":null,"authors":null,"author_count":null,"doi":"10.1609/aaai.v31i1.10911","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1609/aaai.v31i1.10911","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"f18ba036-088b-49f7-9c4a-35182ce9c57f","title":"A lightweight robotic arm with pneumatic muscles for robot learning.","year":2016,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2016.7487599","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2016.7487599","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"e6f2b021-7efc-41f9-a060-eefd8f05641a","title":"A new trajectory generation framework in robotic table tennis.","year":2016,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2016.7759552","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2016.7759552","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"7693d3d8-63c8-4a3c-991e-a4cb609c11cc","title":"Active tactile object exploration with Gaussian processes.","year":2016,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2016.7759723","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2016.7759723","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"c4f5cdcf-2005-43d6-83b8-769f21c69bac","title":"Catching heuristics are optimal control policies.","year":2016,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/BelousovNRP16.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"0587b1fa-2d56-496c-88b0-bf5e7a9906d5","title":"Deep spiking networks for model-based planning in humanoids.","year":2016,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2016.7803344","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2016.7803344","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"7249fac6-1842-49d3-83c1-649b13841aa4","title":"Demonstration based trajectory optimization for generalizable robot motions.","year":2016,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2016.7803324","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2016.7803324","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"0d243ba3-d48f-49dd-a8e5-106c2c3938ad","title":"Incremental imitation learning of context-dependent motor skills.","year":2016,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2016.7803300","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2016.7803300","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"435a9439-b058-485e-952f-bbf44cc6e24e","title":"Jointly learning trajectory generation and hitting point prediction in robot table tennis.","year":2016,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2016.7803343","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2016.7803343","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"29c91cd2-3f25-4514-82af-1a557cf494ab","title":"Learning soft task priorities for control of redundant robots.","year":2016,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2016.7487137","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2016.7487137","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"89e99847-46c4-48c7-9a08-8347df0b45c4","title":"Movement primitives with multiple phase parameters.","year":2016,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2016.7487134","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2016.7487134","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"04809b1e-dbe7-423d-8d28-4f06c01fe8b9","title":"Probabilistic decomposition of sequential force interaction tasks into Movement Primitives.","year":2016,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2016.7759577","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2016.7759577","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"9a397049-abf3-4759-a941-50720994c677","title":"Stability of Controllers for Gaussian Process Forward Models.","year":2016,"date":null,"venue":"ICML","venue_slug":"icml","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/icml/VinogradskaBNRS16.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"764cbd19-1f97-4f4e-bdbd-1bf0bbaea99b","title":"Stable reinforcement learning with autoencoders for tactile and visual data.","year":2016,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2016.7759578","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2016.7759578","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"d135af4f-b61b-4ede-8a27-4a024c72d46b","title":"Using probabilistic movement primitives for striking movements.","year":2016,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2016.7803322","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2016.7803322","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"50bc11db-e270-4906-be76-92b71773534c","title":"A comparison of contact distribution representations for learning to predict object interactions.","year":2015,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2015.7363435","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2015.7363435","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"0fa41cf9-5c68-4470-9b42-8820753c06bc","title":"Combined pose-wrench and state machine representation for modeling Robotic Assembly Skills.","year":2015,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2015.7353471","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2015.7353471","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"93bb47f1-2dd1-4089-be2b-83d6d5c1313c","title":"Evaluation of tactile feature extraction for interactive object recognition.","year":2015,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2015.7363560","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2015.7363560","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"86fee168-3f78-4889-ad63-3990bd3278cc","title":"Extracting low-dimensional control variables for movement primitives.","year":2015,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2015.7139390","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2015.7139390","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"2cfc31be-13c3-4f5f-bad4-b59359dceab9","title":"First-person tele-operation of a humanoid robot.","year":2015,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2015.7363475","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2015.7363475","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"824f7c07-ffc5-4bab-a90c-6f558b7f5306","title":"Learning inverse dynamics models with contacts.","year":2015,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2015.7139638","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2015.7139638","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"8d055bcf-ddc6-4db0-bc23-3110a924f869","title":"Learning motor skills from partially observed movements executed at different speeds.","year":2015,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2015.7353412","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2015.7353412","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"de1802a1-00e0-4127-abd3-c1ff2d377183","title":"Learning multiple collaborative tasks with a mixture of Interaction Primitives.","year":2015,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2015.7139393","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2015.7139393","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"4341851b-519b-41ca-bcbb-7e421eefe6c7","title":"Learning of Non-Parametric Control Policies with High-Dimensional State Features.","year":2015,"date":null,"venue":"AISTATS","venue_slug":"aistats","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/aistats/Hoof0N15.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"d2dfeee9-57ad-401b-b941-9b2a1d49eded","title":"Learning optimal striking points for a ping-pong playing robot.","year":2015,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2015.7354030","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2015.7354030","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"157d407c-98b2-4465-aa3c-27028b934a2f","title":"Learning robot in-hand manipulation with tactile features.","year":2015,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2015.7363524","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2015.7363524","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"bd6260e6-270d-4ef5-9ea3-e2bfad7bfabe","title":"Learning torque control in presence of contacts using tactile sensing from robot skin.","year":2015,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2015.7363429","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2015.7363429","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"80e9d68b-b66a-425d-b02e-45f5ae3ed21a","title":"Model-Based Relative Entropy Stochastic Search.","year":2015,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/AbdolmalekiLPLR15.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"5238c876-b9fc-4388-9b8d-3737f4ccbdae","title":"Model-free Probabilistic Movement Primitives for physical interaction.","year":2015,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2015.7353771","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2015.7353771","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"8553dc0c-d807-4c62-8fae-3b01c8cb663b","title":"Optimizing robot striking movement primitives with Iterative Learning Control.","year":2015,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2015.7363535","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2015.7363535","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"dcda5daa-f94c-41f1-839a-47b4bc55b28f","title":"Probabilistic progress prediction and sequencing of concurrent movement primitives.","year":2015,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2015.7353411","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2015.7353411","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"406b628a-8788-45ad-86e3-b51b64b20a16","title":"Probabilistic segmentation applied to an assembly task.","year":2015,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2015.7363584","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2015.7363584","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"f23d17dd-59cb-4b2e-923b-7a7af9d9b8b9","title":"Reinforcement learning vs human programming in tetherball robot games.","year":2015,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2015.7354296","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2015.7354296","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"4f5e20de-9787-4573-99d5-33914fefdd41","title":"Stabilizing novel objects by learning to predict tactile slip.","year":2015,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2015.7354090","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2015.7354090","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"b78e0712-c4fb-4aa1-bd67-8dead7011144","title":"Towards learning hierarchical skills for multi-phase manipulation tasks.","year":2015,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2015.7139389","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2015.7139389","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"62848bb5-e6fb-4ad8-92aa-1672d93ee1a4","title":"Active Reward Learning.","year":2014,"date":null,"venue":"RSS","venue_slug":"rss","venue_type":null,"authors":null,"author_count":null,"doi":"10.15607/rss.2014.x.031","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.15607/rss.2014.x.031","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"1edac320-d8f8-4ef8-a1b8-5608f7e4995d","title":"An experimental comparison of Bayesian optimization for bipedal locomotion.","year":2014,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2014.6907117","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2014.6907117","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"c82a6861-d2d9-4a92-857c-b0e3eb42f57f","title":"Dimensionality reduction for probabilistic movement primitives.","year":2014,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2014.7041454","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2014.7041454","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"097bcd63-5ef0-449d-bc0b-6004809266a6","title":"Generalizing pouring actions between objects using warped parameters.","year":2014,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2014.7041426","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2014.7041426","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"96e3bcc4-f241-4e38-8a9d-b6e5d080d1a4","title":"Interaction primitives for human-robot cooperation tasks.","year":2014,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2014.6907265","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2014.6907265","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"6a0910bb-90f6-4d91-9e59-04e5df792477","title":"Latent space policy search for robotics.","year":2014,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2014.6942745","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2014.6942745","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"222387b6-bbb7-43ca-871e-94ccdb7a0364","title":"Learning interaction for collaborative tasks with probabilistic movement primitives.","year":2014,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2014.7041413","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2014.7041413","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"786287b9-9dab-4e63-a60c-d0571ca7bd0b","title":"Learning robot tactile sensing for object manipulation.","year":2014,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2014.6943031","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2014.6943031","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"02bd946d-a495-4bf1-a067-4012ab29b083","title":"Learning to predict phases of manipulation tasks as hidden states.","year":2014,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2014.6907441","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2014.6907441","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"a68e244a-a69f-4b73-bee0-d0b5db867de4","title":"Learning to sequence movement primitives from demonstrations.","year":2014,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2014.6943187","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2014.6943187","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"9d56c345-7bc1-4c65-bb38-f3fd9b9e52a0","title":"Multi-task policy search for robotics.","year":2014,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2014.6907421","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2014.6907421","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"810515e6-fe6a-4585-b0af-81814ff188fd","title":"Policy search for learning robot control using sparse data.","year":2014,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2014.6907422","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2014.6907422","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"59f7e3dd-82cc-4140-9d82-a20c04614d3e","title":"Predicting object interactions from contact distributions.","year":2014,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2014.6943030","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2014.6943030","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"a1473b7d-874e-466c-bf79-2057223f4473","title":"Robust policy updates for stochastic optimal control.","year":2014,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2014.7041389","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2014.7041389","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"fac97b3a-e00c-470b-b0d4-f04b4c1ddf60","title":"Sample-based informationl-theoretic stochastic optimal control.","year":2014,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2014.6907424","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2014.6907424","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"4a6cbdd5-3e81-4274-8d96-65dcb89e5e26","title":"Tools for simulating humanoid robot dynamics: A survey based on user feedback.","year":2014,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2014.7041462","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2014.7041462","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"50b762f8-3684-46e0-9a4e-100b88eddc93","title":"A probabilistic approach to robot trajectory generation.","year":2013,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2013.7030017","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2013.7030017","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"3f672065-7b77-4f12-bccf-c96f7b9ca2b1","title":"Data-Efficient Generalization of Robot Skills with Contextual Policy Search.","year":2013,"date":null,"venue":"AAAI","venue_slug":"aaai","venue_type":null,"authors":null,"author_count":null,"doi":"10.1609/aaai.v27i1.8546","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1609/aaai.v27i1.8546","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"635fd76a-5de2-436c-9101-e9b651462822","title":"Feedback error learning for rhythmic motor primitives.","year":2013,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2013.6630741","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2013.6630741","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"bc44444b-4ab7-4ba9-bb19-f3d9a249ee10","title":"Learning responsive robot behavior by imitation.","year":2013,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2013.6696819","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2013.6696819","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"7e081c12-9973-4d30-90c1-567de49965e8","title":"Learning sequential motor tasks.","year":2013,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2013.6630937","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2013.6630937","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"1d56930a-413d-45b0-bff1-92c0bf1235e7","title":"Model-based imitation learning by probabilistic trajectory matching.","year":2013,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2013.6630832","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2013.6630832","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"55afba49-78c7-494c-9623-1bc64aacf29c","title":"Probabilistic interactive segmentation for anthropomorphic robots in cluttered environments.","year":2013,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2013.7029972","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2013.7029972","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"cef52add-ce90-4cb4-83d4-a7e54bbaf8da","title":"Probabilistic Movement Primitives.","year":2013,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/ParaschosDPN13.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"a15c24f8-bf93-4f77-af16-0b16c06483ff","title":"A brain-robot interface for studying motor learning after stroke.","year":2012,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2012.6385646","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2012.6385646","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"61a02806-468b-4da4-a191-06514600fcaf","title":"A kernel-based approach to direct action perception.","year":2012,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2012.6224957","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2012.6224957","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"894ca968-361c-4511-a392-a7e96ea8e33e","title":"Algorithms for Learning Markov Field Policies.","year":2012,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/BoulariasKP12.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"3ef996ae-8956-420a-b78d-42d75760b3c9","title":"Generalization of human grasping for multi-fingered robot hands.","year":2012,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2012.6386072","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2012.6386072","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"13b64ab1-3158-492b-b12f-4ef511177834","title":"Learning concurrent motor skills in versatile solution spaces.","year":2012,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2012.6386047","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2012.6386047","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"5c5c7e1e-a1c1-4201-88ff-fe9edc949e03","title":"Learning throwing and catching skills.","year":2012,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2012.6386267","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2012.6386267","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"0678b4d9-df1a-449c-9728-3a2dd574d8a8","title":"Learning tracking control with forward models.","year":2012,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2012.6224831","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2012.6224831","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"9cd4d2e2-44d4-481c-bdaa-37a4b9ccf230","title":"Maximally informative interaction learning for scene exploration.","year":2012,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2012.6386008","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2012.6386008","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"5e36fb53-17a3-4921-9658-bc560b3686b7","title":"Point cloud completion using extrusions.","year":2012,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/humanoids.2012.6651593","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/humanoids.2012.6651593","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"dc819ae3-1343-433f-b364-d9cad4e4f465","title":"Probabilistic Modeling of Human Movements for Intention Inference.","year":2012,"date":null,"venue":"RSS","venue_slug":"rss","venue_type":null,"authors":null,"author_count":null,"doi":"10.15607/rss.2012.viii.055","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.15607/rss.2012.viii.055","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"58df6cb8-7922-4bca-bcf4-eea4146febfa","title":"Toward fast policy search for learning legged locomotion.","year":2012,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2012.6385955","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2012.6385955","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"e59cd145-dc1a-4e2a-9075-0feec776ea0a","title":"A flexible hybrid framework for modeling complex manipulation tasks.","year":2011,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2011.5980237","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2011.5980237","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"c2dfb911-71cd-4d85-a0bc-3052391f6f2d","title":"A Non-Parametric Approach to Dynamic Programming.","year":2011,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/KroemerP11.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"93a8b7a0-165f-4b64-8916-da10025ad5a6","title":"Balancing Safety and Exploitability in Opponent Modeling.","year":2011,"date":null,"venue":"AAAI","venue_slug":"aaai","venue_type":null,"authors":null,"author_count":null,"doi":"10.1609/aaai.v25i1.7981","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1609/aaai.v25i1.7981","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"24b7827d-a868-4def-a075-dc228e2dabf7","title":"Learning anticipation policies for robot table tennis.","year":2011,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2011.6094892","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2011.6094892","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"f99abb55-1a7a-42d5-b7d3-2b0d1844c78d","title":"Learning elementary movements jointly with a higher level task.","year":2011,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2011.6094834","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2011.6094834","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"2cb1d27d-d705-4b71-a4d6-3940eca6c50f","title":"Learning inverse kinematics with structured prediction.","year":2011,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2011.6094666","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2011.6094666","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"907b67ef-1376-458b-b0bd-881943f53128","title":"Learning robot grasping from 3-D images with Markov Random Fields.","year":2011,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2011.6094888","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2011.6094888","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"e808953a-00db-42b9-9404-f70a9c97f22a","title":"Learning task-space tracking control with kernels.","year":2011,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2011.6094428","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2011.6094428","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"1f1d9aed-9337-46c9-bf33-b8dab5e8bce5","title":"Modeling Opponent Actions for Table-Tennis Playing Robot.","year":2011,"date":null,"venue":"AAAI","venue_slug":"aaai","venue_type":null,"authors":null,"author_count":null,"doi":"10.1609/aaai.v25i1.8051","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1609/aaai.v25i1.8051","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"ff5a5708-0999-4b05-b202-35d33eaa6cef","title":"Reinforcement Learning to Adjust Robot Movements to New Situations.","year":2011,"date":null,"venue":"IJCAI","venue_slug":"ijcai","venue_type":null,"authors":null,"author_count":null,"doi":"10.5591/978-1-57735-516-8/ijcai11-441","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.5591/978-1-57735-516-8/ijcai11-441","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"c5cabfd3-027d-43d6-b6d2-da01072ec3bf","title":"Trajectory planning for optimal robot catching in real-time.","year":2011,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/icra.2011.5980114","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/icra.2011.5980114","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"fbafc2ba-00d9-497b-abea-25577179c724","title":"A biomimetic approach to robot table tennis.","year":2010,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2010.5650305","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2010.5650305","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"ef6ee98c-8b6f-40c1-95f3-9a22358a2bf3","title":"Learning probabilistic discriminative models of grasp affordances under limited supervision.","year":2010,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2010.5650088","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2010.5650088","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"0458e914-e591-44f8-a0f0-c8fbf9ede946","title":"Learning table tennis with a Mixture of Motor Primitives.","year":2010,"date":null,"venue":"Humanoids","venue_slug":"humanoids","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/ichr.2010.5686298","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/ichr.2010.5686298","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"4897002e-60b5-406b-967e-0a4ac07e58b7","title":"Movement extraction by detecting dynamics switches and repetitions.","year":2010,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/ChiappaP10.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"7ac0eca1-5db5-4e75-8dc6-b9c03362c3f3","title":"Movement templates for learning of hitting and batting.","year":2010,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/robot.2010.5509672","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/robot.2010.5509672","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"f5c06475-0874-4ae0-9eef-d6b11149362e","title":"Relative Entropy Policy Search.","year":2010,"date":null,"venue":"AAAI","venue_slug":"aaai","venue_type":null,"authors":null,"author_count":null,"doi":"10.1609/aaai.v24i1.7727","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1609/aaai.v24i1.7727","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"057b06f3-8bcd-4223-ba38-57d3aab30715","title":"Switched Latent Force Models for Movement Segmentation.","year":2010,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/AlvarezPSL10.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"4c67de69-de80-496e-8da2-4377cc1cb032","title":"Using model knowledge for learning inverse dynamics.","year":2010,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/robot.2010.5509858","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/robot.2010.5509858","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"e10e1ac2-5134-4a76-8f9d-7314c4a509d3","title":"Active learning using mean shift optimization for robot grasping.","year":2009,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2009.5354345","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2009.5354345","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"309b83f3-546c-4235-bc8c-5903b4252a2e","title":"Learning complex motions by sequencing simpler motion templates.","year":2009,"date":null,"venue":"ICML","venue_slug":"icml","venue_type":null,"authors":null,"author_count":null,"doi":"10.1145/1553374.1553471","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1145/1553374.1553471","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"1124198c-56f4-4f5c-9284-73f5f9e604b8","title":"Learning motor primitives for robotics.","year":2009,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/robot.2009.5152577","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/robot.2009.5152577","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"7e17aa3c-8c84-4acc-bf8e-b26f3f88d954","title":"Sparse online model learning for robot control with support vector regression.","year":2009,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2009.5354609","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2009.5354609","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"c402fec2-3b27-473e-8d04-58eff0c70afc","title":"Adaptive Importance Sampling with Automatic Model Selection in Value Function Approximation.","year":2008,"date":null,"venue":"AAAI","venue_slug":"aaai","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/aaai/HachiyaASP08.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"c69acec2-6a46-49da-85ce-0dd53f27ffc1","title":"Fitted Q-iteration by Advantage Weighted Regression.","year":2008,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/NeumannP08.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"d230c366-271b-430c-8871-53cf9d5f420e","title":"Learning perceptual coupling for motor primitives.","year":2008,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2008.4650953","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2008.4650953","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"07f625cb-cabc-4690-8fcc-299dd4c48285","title":"Local Gaussian Process Regression for Real Time Online Model Learning.","year":2008,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/Nguyen-TuongSP08.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"529de248-76eb-4013-a30a-7e95c78842c1","title":"Local Gaussian process regression for real-time model-based robot control.","year":2008,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2008.4650850","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2008.4650850","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"9d725b2c-bb90-44a3-a223-c810b957c952","title":"Policy Search for Motor Primitives in Robotics.","year":2008,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/KoberP08.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"a2130976-d17d-4c36-862b-75e8296d3029","title":"Real-time learning of resolved velocity control on a Mitsubishi PA-10.","year":2008,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/robot.2008.4543645","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/robot.2008.4543645","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"0a496a31-0fc7-435d-a5ca-dd625e2c02a9","title":"Using Bayesian Dynamical Systems for Motion Template Libraries.","year":2008,"date":null,"venue":"NeurIPS","venue_slug":"neurips","venue_type":null,"authors":null,"author_count":null,"doi":null,"arxiv_id":null,"openreview_id":null,"landing_url":"https://dblp.org/rec/conf/nips/ChiappaKP08.html","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"b9326d91-20ff-4027-9ef7-1fce675e9fce","title":"Reinforcement learning by reward-weighted regression for operational space control.","year":2007,"date":null,"venue":"ICML","venue_slug":"icml","venue_type":null,"authors":null,"author_count":null,"doi":"10.1145/1273496.1273590","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1145/1273496.1273590","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"01c53e49-7424-4a2b-b38e-736953c274a2","title":"Reinforcement Learning for Operational Space Control.","year":2007,"date":null,"venue":"ICRA","venue_slug":"icra","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/robot.2007.363633","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/robot.2007.363633","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"cf0c2116-224c-4824-a7bd-57eb05bf66c9","title":"Towards compliant humanoids-an experimental assessment of suitable task space position/orientation controllers.","year":2007,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2007.4399562","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2007.4399562","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"58fb86cc-365d-470c-93c3-3f6994d0d13a","title":"A Bayesian Approach to Nonlinear Parameter Identification for Rigid Body Dynamics.","year":2006,"date":null,"venue":"RSS","venue_slug":"rss","venue_type":null,"authors":null,"author_count":null,"doi":"10.15607/rss.2006.ii.032","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.15607/rss.2006.ii.032","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"a6e25d39-1eb1-4c40-bf0d-5138b6e3782d","title":"Learning Operational Space Control.","year":2006,"date":null,"venue":"RSS","venue_slug":"rss","venue_type":null,"authors":null,"author_count":null,"doi":"10.15607/rss.2006.ii.033","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.15607/rss.2006.ii.033","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"d7cb1025-1a30-4665-b9c6-1978cc0d7f8c","title":"Policy Gradient Methods for Robotics.","year":2006,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2006.282564","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2006.282564","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"098baae8-5414-469d-80cc-bec83585ad3e","title":"A unifying methodology for the control of robotic systems.","year":2005,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2005.1545516","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2005.1545516","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"},{"id":"ebebecfa-46f9-4a31-8360-dc6f4d109744","title":"Comparative experiments on task space control with redundancy resolution.","year":2005,"date":null,"venue":"IROS","venue_slug":"iros","venue_type":null,"authors":null,"author_count":null,"doi":"10.1109/iros.2005.1545203","arxiv_id":null,"openreview_id":null,"landing_url":"https://doi.org/10.1109/iros.2005.1545203","pdf_url":null,"links":null,"keywords":null,"tldr":null,"cited_by_count":null,"lab_submitted":false,"register_url":null,"updated_at":"2026-07-01T20:27:36.040674+00:00"}]}