{"total":25,"items":[{"citing_arxiv_id":"2607.08742","ref_index":16,"ref_count":1,"confidence":0.98,"is_internal_anchor":true,"paper_title":"ContactMimic: Humanoid Object Interaction via Contact Control","primary_cat":"cs.RO","submitted_at":"2026-07-09T17:42:15+00:00","verdict":"CONDITIONAL","verdict_confidence":"MODERATE","novelty_score":6.0,"formal_verification":"none","one_line_summary":"A humanoid tracking policy is trained with contact-following rewards and trajectory augmentation to decouple physical contact from keypoint geometry, enabling runtime contact control.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2607.06052","ref_index":4,"ref_count":1,"confidence":0.98,"is_internal_anchor":true,"paper_title":"ThorArena: Benchmarking Humanoid Physical Interaction with Human Motion-Force Demonstrations","primary_cat":"cs.RO","submitted_at":"2026-07-07T09:26:51+00:00","verdict":"CONDITIONAL","verdict_confidence":"MODERATE","novelty_score":6.0,"formal_verification":"none","one_line_summary":"A force-aware humanoid benchmark pairs synchronized human motion-force data with simulation-based force replay to evaluate whole-body control policies under realistic physical disturbances.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2606.29940","ref_index":10,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"WARP: Whole-Body Retargeting for Learning from Offline Human Demonstrations","primary_cat":"cs.RO","submitted_at":"2026-06-29T08:17:31+00:00","verdict":null,"verdict_confidence":null,"novelty_score":null,"formal_verification":null,"one_line_summary":null,"context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2606.27676","ref_index":22,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"CWI: Composite Humanoid Whole-Body Imitation System for Loco-manipulation","primary_cat":"cs.RO","submitted_at":"2026-06-26T03:14:52+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":6.0,"formal_verification":"none","one_line_summary":"CWI decouples MoCap data for upper-body manipulation and lower-body locomotion, using dual discriminators and multi-critic training plus distillation to produce a policy that works from hand poses and velocity commands alone.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2606.26201","ref_index":1,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"OmniContact: Chaining Meta-Skills via Contact Flow for Generalizable Humanoid Loco-Manipulation","primary_cat":"cs.RO","submitted_at":"2026-06-24T16:28:23+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":6.0,"formal_verification":"none","one_line_summary":"OmniContact introduces contact flow as a shared representation of body trajectories and contact signals to learn and chain loco-manipulation meta-skills, reporting 98.7% success on box carrying and 76.5% on push-stack tasks.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2606.23680","ref_index":46,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"CoorDex: Coordinating Body and Hand Priors for Continuous Dexterous Humanoid Loco-Manipulation","primary_cat":"cs.RO","submitted_at":"2026-06-22T17:59:20+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":6.0,"formal_verification":"none","one_line_summary":"CoorDex distills privileged body and hand motion teachers into proprioceptive latent priors and composes them via shared-context residual RL heads to enable continuous high-DoF dexterous loco-manipulation.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2606.22174","ref_index":5,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"OpenHLM: An Empirical Recipe for Whole-Body Humanoid Loco-Manipulation","primary_cat":"cs.RO","submitted_at":"2026-06-20T18:02:50+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":6.0,"formal_verification":"none","one_line_summary":"OpenHLM is an empirical recipe yielding a whole-body humanoid VLA model that outperforms GR00T N1.6 and Ψ0 baselines on long-horizon tasks using less than half the demonstration time.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2606.18772","ref_index":11,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"HALOMI: Learning Humanoid Loco-Manipulation with Active Perception from Human Demonstrations","primary_cat":"cs.RO","submitted_at":"2026-06-17T07:33:37+00:00","verdict":"UNVERDICTED","verdict_confidence":"UNKNOWN","novelty_score":5.0,"formal_verification":"none","one_line_summary":"HALOMI extends UMI with egocentric sensing and a manifold-constrained controller plus alignment adaptations to learn loco-manipulation on humanoids from human demos, reporting 85% average success on three real-world tasks.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2606.17511","ref_index":53,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"MagicSim: A Unified Infrastructure for Executable Embodied Interaction","primary_cat":"cs.RO","submitted_at":"2026-06-16T04:42:43+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":5.0,"formal_verification":"none","one_line_summary":"MagicSim is a unified embodied interaction infrastructure built on a deterministic batched runtime and shared MDP that supports diverse world construction, execution, task evaluation, automatic rollout generation, and interactive agent interfaces.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2606.17446","ref_index":32,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"AnnotateAnything: Automatic Annotation of 3D Assets for Robot Manipulation","primary_cat":"cs.RO","submitted_at":"2026-06-16T03:00:58+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":6.0,"formal_verification":"none","one_line_summary":"AnnotateAnything converts passive 3D assets into manipulation-ready assets by combining vision-language reasoning for semantics with parallel physics pipelines for executable action annotations such as grasps and articulations.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2606.09215","ref_index":5,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"MotionWAM: Towards Foundation World Action Models for Real-Time Humanoid Loco-Manipulation","primary_cat":"cs.RO","submitted_at":"2026-06-08T08:50:14+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":6.0,"formal_verification":"none","one_line_summary":"MotionWAM conditions a policy on intermediate features from a video world model to predict unified whole-body motion tokens, enabling real-time humanoid loco-manipulation that outperforms VLA baselines by over 30% on nine Unitree G1 tasks.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2606.07934","ref_index":1,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"X-OP: Cross-Morphology Whole-Body Teleoperation via MPC Retargeting","primary_cat":"cs.RO","submitted_at":"2026-06-06T01:50:59+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":6.0,"formal_verification":"none","one_line_summary":"MPC-based retargeting framework enables cross-morphology whole-body teleoperation from a single XR device via dynamic feasibility optimization, state synchronization, and SLAM feedback, with reported gains in simulation and real-world tests.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2606.05880","ref_index":2,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"TAGA: Terrain-aware Active Gaze Learning for Generalizable Agile Humanoid Locomotion","primary_cat":"cs.RO","submitted_at":"2026-06-04T08:52:56+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":6.0,"formal_verification":"none","one_line_summary":"TAGA learns terrain-aware active gaze behaviors for humanoid robots via RL alone, enabling generalizable locomotion with 1.2m real-world gap traversal.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2606.05160","ref_index":138,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"GRAIL: Generating Humanoid Loco-Manipulation from 3D Assets and Video Priors","primary_cat":"cs.RO","submitted_at":"2026-06-03T17:57:45+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":5.0,"formal_verification":"none","one_line_summary":"GRAIL creates over 20,000 synthetic loco-manipulation sequences from known 3D configurations and video priors, then trains policies that achieve 84% pick-up and 90% stair-climbing success on a real Unitree G1 humanoid using only the generated data.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2605.27724","ref_index":5,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"HumanoidMimicGen: Data Generation for Loco-Manipulation via Whole-Body Planning","primary_cat":"cs.RO","submitted_at":"2026-05-26T21:57:11+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":5.0,"formal_verification":"none","one_line_summary":"HumanoidMimicGen automatically generates large loco-manipulation datasets from few source demonstrations using whole-body planning, enabling visuomotor policies that outperform real-data-only training by 20% on a new nine-task benchmark.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2605.24592","ref_index":49,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"MuGen: Multi-Skill Generative Locomotion Controller for Humanoid Robots","primary_cat":"cs.RO","submitted_at":"2026-05-23T14:06:06+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":5.0,"formal_verification":"none","one_line_summary":"MuGen learns a generative latent representation of multi-skill humanoid locomotion from heterogeneous human data using VQ-VAEs and RL, then distills a deployable policy that tracks unseen motions and reuses the latent space.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2605.22272","ref_index":45,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"Imagine2Real: Towards Zero-shot Humanoid-Object Interaction via Video Generative Priors","primary_cat":"cs.RO","submitted_at":"2026-05-21T10:15:39+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":6.0,"formal_verification":"none","one_line_summary":"Imagine2Real enables zero-shot humanoid-object interaction by unifying motions as 4D point trajectories, tracking only base/hands/object keypoints inside a BFM latent space, and training with progressive simple rewards for mocap deployment.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2605.21133","ref_index":1,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"Humanoid Whole-Body Manipulation via Active Spatial Brain and Generalizable Action Cerebellum","primary_cat":"cs.RO","submitted_at":"2026-05-20T13:05:31+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":5.0,"formal_verification":"none","one_line_summary":"A multi-agent large-model framework (Active Spatial Brain + Generalizable Action Cerebellum) enables spatial-aware humanoid whole-body manipulation without task-specific real-robot data.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2605.01518","ref_index":24,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"VOFA: Visual Object Goal Pushing with Force-Adaptive Control for Humanoids","primary_cat":"cs.RO","submitted_at":"2026-05-02T16:16:23+00:00","verdict":"CONDITIONAL","verdict_confidence":"MODERATE","novelty_score":6.0,"formal_verification":"none","one_line_summary":"VOFA combines a depth-image visuomotor policy with a force-adaptive whole-body controller to push objects of unknown mass to arbitrary goals on a humanoid.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2604.21355","ref_index":6,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"RPG: Robust Policy Gating for Smooth Multi-Skill Transitions in Humanoid Fighting","primary_cat":"cs.RO","submitted_at":"2026-04-23T07:14:35+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":4.0,"formal_verification":"none","one_line_summary":"RPG trains a unified humanoid robot policy using motion and temporal randomization to achieve smooth, stable transitions between fighting skills and locomotion.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2604.21351","ref_index":39,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"Learn Weightlessness: Imitate Non-Self-Stabilizing Motions on Humanoid Robot","primary_cat":"cs.RO","submitted_at":"2026-04-23T07:10:05+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":6.0,"formal_verification":"none","one_line_summary":"A weightlessness mechanism enables humanoid robots to dynamically relax joints for stable, contact-rich motions across diverse environments without task-specific tuning.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2604.13015","ref_index":32,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"Learning Versatile Humanoid Manipulation with Touch Dreaming","primary_cat":"cs.RO","submitted_at":"2026-04-14T17:54:17+00:00","verdict":null,"verdict_confidence":null,"novelty_score":null,"formal_verification":null,"one_line_summary":null,"context_count":1,"top_context_role":"background","top_context_polarity":"background","context_text":"terous whole-body behaviors [20], [21]. Related systems also combine learned whole-body control with specialized teleoperation hardware or tracking modules for more precise loco-manipulation [30], [31]. Another line instead seeks unified controllers that directly coordinate locomotion and manipulation within a single whole-body tracking frame- work [32], [33]. Complementary teleoperation and motion- tracking systems further improve the practicality and scal- ability of commanding humanoids through RGB- or pose- based shadowing, immersive VR interfaces, portable mocap- free setups, and closed-loop long-horizon tracking [1], [5]- [9], [23], [25], [26], [34]. Building on this line of work, our system combines an RL-based whole-body controller"},{"citing_arxiv_id":"2602.03205","ref_index":5,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"HUSKY: Humanoid Skateboarding System via Physics-Aware Whole-Body Control","primary_cat":"cs.RO","submitted_at":"2026-02-03T07:18:01+00:00","verdict":"CONDITIONAL","verdict_confidence":"MODERATE","novelty_score":6.0,"formal_verification":"none","one_line_summary":"HUSKY combines humanoid-skateboard dynamics modeling with adversarial motion priors and physics-guided lean-to-steer strategies to achieve real-world stable skateboarding on a humanoid robot.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2511.22963","ref_index":3,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"Commanding Humanoid by Free-form Language: A Large Language Action Model with Unified Motion Vocabulary","primary_cat":"cs.RO","submitted_at":"2025-11-28T08:11:24+00:00","verdict":"UNVERDICTED","verdict_confidence":"LOW","novelty_score":6.0,"formal_verification":"none","one_line_summary":"Humanoid-LLA converts unconstrained natural language commands into stable whole-body motions for humanoid robots using a unified motion vocabulary and two-stage supervised-plus-reinforcement fine-tuning.","context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null},{"citing_arxiv_id":"2511.07820","ref_index":4,"ref_count":1,"confidence":0.9,"is_internal_anchor":false,"paper_title":"SONIC: Supersizing Motion Tracking for Natural Humanoid Whole-Body Control","primary_cat":"cs.RO","submitted_at":"2025-11-11T04:37:40+00:00","verdict":null,"verdict_confidence":null,"novelty_score":null,"formal_verification":null,"one_line_summary":null,"context_count":0,"top_context_role":null,"top_context_polarity":null,"context_text":null}],"limit":50,"offset":0}