{"as_of":"2026-08-07T23:08:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:e628a89e5364d25509eadf78dc942f5a566c4b5a29a801362aad413a7f466bfa","coverage":[{"denominator":35,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":35,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T11:33:50.317068Z","state":"measured"},{"denominator":36,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":36,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-01T22:35:37.040638Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.02206","snapshot_observed_at":"2026-08-01T22:35:37.040638Z","title":"Reinforcement learning with data bootstrapping for dynamic subgoal pursuit in humanoid robot navigation,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.15701","last_updated":"2026-07-17T07:23:54Z","snapshot_observed_at":"2026-08-07T17:58:06.995714Z","submitted_at":"2026-07-17T07:23:54Z","title":"RAVEN: Reinforcement-Adaptive Visibility-Graph Planning for Robust Humanoid Navigation with Collision-Free MPC","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-01T22:35:37.040638Z"},"links":{"cited_paper":"/paper/2506.02206","citing_paper":"/paper/2607.15701"},"observation_digest":"sha256:7d90a020313b8684c13979b55368a55e1156a8bf5f79a3e54593fadc01524848","observation_id":"b3964702-4ab5-47bb-ba95-e3886eecbf7c","resolution":{"observed_at":"2026-08-01T22:35:37.040638Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2506.02206/citation-record","integrity":"/paper/2506.02206/integrity","json":"/paper/2506.02206/citation-record.json","paper":"/paper/2506.02206"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:57.455223Z","title":"Introduction of the Foot Placement Estimator: A Dynamic Measure of Balance for Bipedal Robotics,","venue":null,"work_id":"59c1a50b-ec63-4425-a7c3-da4b21802db1","year":2007},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:45.982721Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:7370fff9cbdac0131170425cf1708469b96e3531927a9691307143fcec444856","observation_id":"5e55a5b1-3656-4b35-a091-6a474d898f24","resolution":{"observed_at":"2026-08-07T11:33:57.551834Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:57.185280Z","title":"Navigation planning for legged robots in challenging terrain,","venue":null,"work_id":"61b4a866-30f6-4b8a-a1a6-0e5af178622b","year":2016},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:46.067955Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:676f8ab14fcbe87529040a35ec8b8ac94659b503c79dc6dcfce8cb099ede65a6","observation_id":"e184e266-55c1-446f-a646-fa74e451b99a","resolution":{"observed_at":"2026-08-07T11:33:57.306213Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:46.171453Z","title":"Optimization-based locomotion planning, estimation, and control design for the atlas humanoid robot,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:46.171453Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:8500e30c7c822b2d64179c48f67609a273766cd1a1beecb5776c2b4783e2c2aa","observation_id":"787cc79f-2b84-4667-857e-08cc2e5eca42","resolution":{"observed_at":"2026-08-07T11:33:46.171453Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:56.917037Z","title":"Exact cell decomposition of arrangements used for path planning in robotics,","venue":null,"work_id":"746ffa80-b0bd-438a-a625-9cb40d41b0cb","year":1999},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:46.263030Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:4b922b7027f92cae9a35fbfdb1d521231df957cc76b97d71645a45753acabbd2","observation_id":"43d458b8-00df-4988-b1d6-6210be1fec8c","resolution":{"observed_at":"2026-08-07T11:33:57.032791Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:56.668222Z","title":"An overview of autonomous mobile robot path planning algorithms,","venue":null,"work_id":"c1f94a8f-0a59-4ca4-bea9-edf242ed6dd9","year":2006},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:46.395990Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:f439c4e4d8314c87b5d19cdfcf56eb5e54c49e53f3e58d42d6cb7b84b2554cb4","observation_id":"91759f8a-8b67-4a89-a046-423966bdbc46","resolution":{"observed_at":"2026-08-07T11:33:56.783635Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:56.380269Z","title":"Path planning and trajectory planning algorithms: A general overview,","venue":null,"work_id":"903c2952-4da9-4f08-b0e7-69d88f7f37b6","year":2015},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:46.547721Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:6dd73e548d23e3d724c99529aee38b33e7d5d8af79e8d89f69919d42178dfb7b","observation_id":"d231c70e-f21f-46b6-ad9c-740947232dfc","resolution":{"observed_at":"2026-08-07T11:33:56.536460Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:56.049237Z","title":"Confidence random tree- based algorithm for mobile robot path planning considering the path length and safety,","venue":null,"work_id":"864298f5-406d-4431-a954-01f1d923d07b","year":2019},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:46.715004Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:1a35ca2cd3df5842f79291d1ca27dace4318a44e228f3544d071789a2bd90477","observation_id":"506751ae-76fb-4abb-bdf9-83af3c55f446","resolution":{"observed_at":"2026-08-07T11:33:56.205520Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:55.823344Z","title":"Optimization-based motion planning for legged robots,","venue":null,"work_id":"3db66bc9-67a4-4ef9-aa64-40e88c1b51b1","year":2018},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:46.882468Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:c6bddc832e4df3107717a21f3642e164f5eb8583f4f83f05664b0771df7ffaff","observation_id":"cccc64f7-cb5b-4dcd-a7eb-6f09f0d3a155","resolution":{"observed_at":"2026-08-07T11:33:55.930491Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:55.535636Z","title":"Integrated task and motion planning for safe legged navigation in partially observable environments,","venue":null,"work_id":"afed0956-e20d-4e06-8652-ede437975c28","year":2023},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:46.984997Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:07c03bc2c0b9a6102f3dfe7f0078820ea610d93dd9cd69cfa88ba1f743242a7f","observation_id":"1fe19592-9ca4-4e45-b327-3d3130da7bb5","resolution":{"observed_at":"2026-08-07T11:33:55.644239Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.17347","last_updated":"2024-03-26T03:15:14Z","snapshot_observed_at":"2026-07-06T17:50:39.208397Z","submitted_at":"2024-03-26T03:15:14Z","title":"Unified Path and Gait Planning for Safe Bipedal Robot Navigation","version":1},"cited_work":{"arxiv_id":"2403.17347","doi":null,"metadata_source":"pith","pith_arxiv_id":"2403.17347","snapshot_observed_at":"2026-08-07T11:33:50.527082Z","title":"Unified Path and Gait Planning for Safe Bipedal Robot Navigation","venue":"cs.RO","work_id":"81c8cb71-52ad-43fe-86cd-c250bfc99b65","year":2024},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:47.137767Z"},"links":{"cited_paper":"/paper/2403.17347","citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:528dda46ae977474be638e2275de2a252237d7a49f2004f944ec8c592eafe8f1","observation_id":"185c0dd2-a53c-4d03-9b02-a89681136225","resolution":{"observed_at":"2026-08-07T11:33:50.614740Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:55.289792Z","title":"Fast direct multiple shooting algorithms for optimal robot control,","venue":null,"work_id":"1a972612-4547-48ff-94d1-23fcf4f9e35f","year":2006},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:47.253223Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:44344691397bfe2e58057036f67c8e12404644da6113071c75d84ce65ae21576","observation_id":"ee89f1e1-1319-4112-901e-3f9ca352ce98","resolution":{"observed_at":"2026-08-07T11:33:55.421021Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:55.008695Z","title":"Using optimization to create self-stable human-like running,","venue":null,"work_id":"c5c938a0-5cb9-48f2-a65a-029868c54fe5","year":2009},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:47.401679Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:422ef112d4865bd3a78c2bdd2578b191ff9043b45da6c6363ef5dd327e976c1a","observation_id":"4a83a47b-1afa-4501-b94b-438e51b99606","resolution":{"observed_at":"2026-08-07T11:33:55.146851Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:54.709142Z","title":"Whole-body motion planning with centroidal dynamics and full kinematics,","venue":null,"work_id":"0bd90b9d-f5a4-4091-8c9c-63202cbd0e2d","year":2014},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:47.499247Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:9222ed31a63ac0262d84143a77313015b2504e00bef7b0e670418eed1b12202c","observation_id":"4073e0d1-07d9-4e0d-81eb-8374028d13bf","resolution":{"observed_at":"2026-08-07T11:33:54.878671Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:54.455784Z","title":"The 3d linear inverted pendulum mode: A simple modeling for a biped walking pattern generation,","venue":null,"work_id":"9d1ae947-a7c2-4875-9d88-4711efc15bc4","year":2001},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:47.670448Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:8c84f7bae13500bed9668f6ce2846f38257019f4ba7fdbef1067f7e72624994a","observation_id":"ad922e4f-c540-40af-b183-2aefc9c91cae","resolution":{"observed_at":"2026-08-07T11:33:54.561134Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:54.168581Z","title":"Bipedal walking control based on capture point dynamics,","venue":null,"work_id":"fa45419c-cb8c-4beb-ad93-1237c6e99e1c","year":2011},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:47.823455Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:891e6a60f83886bc0e8ef1d20572279a0d55a6cafa6cef2b8f9da5802484402b","observation_id":"339963da-d9e3-4b9f-8617-f27af015526f","resolution":{"observed_at":"2026-08-07T11:33:54.298431Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:53.937074Z","title":"Nonlinear model predictive control for rough-terrain robot hopping,","venue":null,"work_id":"ea2ffa79-8d90-4023-adda-e6c221e9d7e5","year":2012},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:48.004964Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:a435e827fac2b03064ccdac4b95f01d50def9316c7c29bfca81bcb49c5f34e5f","observation_id":"a56aac47-df3a-4a49-aeea-6703e77f7eff","resolution":{"observed_at":"2026-08-07T11:33:54.040334Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:48.116002Z","title":"Perceptive locomotion through nonlinear model-predictive control,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:48.116002Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:bc1f4b6794e8f13686376f88df85e74a7a32571b4504b3585bb67531eeda7d05","observation_id":"dbcb9620-946d-467a-9667-d332fcf5a736","resolution":{"observed_at":"2026-08-07T11:33:48.116002Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:53.652579Z","title":"Model predictive control for dynamic footstep adjustment using the divergent component of motion,","venue":null,"work_id":"03b60288-f61a-472f-8201-79d6c198e6a6","year":2016},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:48.254270Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:b8b3d7911804b9fc6fb09579cbb228d00eae1440de8004d618ea9f3314cdb08d","observation_id":"c1f7cd1c-0a89-480b-b5ec-f3d9243b81f6","resolution":{"observed_at":"2026-08-07T11:33:53.786744Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:48.362816Z","title":"A sequential mpc approach to reactive planning for bipedal robots using safe corridors in highly cluttered environments,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:48.362816Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:a3c4de0c3e6274ac0ce9f355a374f850d6068c9c73bb6e6d5084adea7c3ca397","observation_id":"566d52c9-0baa-45f1-ae10-292bcd563abf","resolution":{"observed_at":"2026-08-07T11:33:48.362816Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:53.390426Z","title":"Real-time safe bipedal robot navigation using linear discrete control barrier functions,","venue":null,"work_id":"4c3902ac-b330-4f22-b26b-37ce9ccdd85d","year":2025},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:48.481563Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:b72000e48d0e8b88d79b597d88abba0b4b42f0f8e3a6f7a52d616f5e03d04a99","observation_id":"4a85c47e-894c-43b0-a189-b32e9e7d9491","resolution":{"observed_at":"2026-08-07T11:33:53.514967Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:48.595799Z","title":"Apprenticeship learning via inverse rein- forcement learning,","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:48.595799Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:d7125bd7885040753350fa82ff77db2be410a1621ce1e0efa1ac0fb1a3fce2cb","observation_id":"4076a78d-f047-4957-b732-9d468378f1f2","resolution":{"observed_at":"2026-08-07T11:33:48.595799Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:53.163026Z","title":"Apprenticeship learning using linear programming,","venue":null,"work_id":"331f959b-6085-4c54-93e5-967de37b32b4","year":2008},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:48.692116Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:8fc4be212980df15121d083e8dd4e76906ea83ba98df20a616fa82d8f1c3dab9","observation_id":"0a5c7f4f-a9f7-4c4a-bd3d-8e2a1b241ac1","resolution":{"observed_at":"2026-08-07T11:33:53.238377Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:48.843022Z","title":"Generative adversarial imitation learning,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:48.843022Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:c547fffdc886ea5bc02a97d7a1d577c8fef4240cefdb38bbdfd00864fac68b1a","observation_id":"eaa33a93-53ab-4b61-918f-3256a4af3882","resolution":{"observed_at":"2026-08-07T11:33:48.843022Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:52.859177Z","title":"Goal-oriented obstacle avoid- ance with deep reinforcement learning in continuous action space,","venue":null,"work_id":"5b79700e-6391-4076-859a-0e00862791bc","year":2020},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:48.937422Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:41409c05748cec0a355d46064eade69a51af235c1735585695232f0eb3b96e5c","observation_id":"a9d0184d-8bb5-4e04-be9a-4bd07c948efb","resolution":{"observed_at":"2026-08-07T11:33:52.999533Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:49.072122Z","title":"Where to go next: Learning a subgoal recommendation policy for navigation in dynamic environments,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:49.072122Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:38da9e2b2de762abaceaa7d84892b7dc6c5e1f7600286a0ff08f5c63548a20b4","observation_id":"32780fab-6113-4503-bb13-19ae90ccba2e","resolution":{"observed_at":"2026-08-07T11:33:49.072122Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:52.642516Z","title":"Robot navigation in constrained pedestrian environments using reinforcement learning,","venue":null,"work_id":"5a214949-4393-4649-bb4b-fde9103593b4","year":2021},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:49.185856Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:c85e2ff4695d97c219ae9d48b930e46f05e081c3fea448d01097a63910046f3b","observation_id":"5cb83d8d-a174-43ae-88c1-8d1898170a7e","resolution":{"observed_at":"2026-08-07T11:33:52.724840Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:52.390014Z","title":"A hierarchical deep reinforcement learning framework with high efficiency and generalization for fast and safe navigation,","venue":null,"work_id":"68e6828d-0411-49c8-906b-8e29160505ff","year":2022},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:49.305115Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:f89a877267def3d3fd31fe5353f8d8fe045279bc9f87e6660e8eeec6c3b9a50c","observation_id":"072de955-c836-4511-9ba6-135d96e9f8a8","resolution":{"observed_at":"2026-08-07T11:33:52.486999Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:52.064974Z","title":"Drl-vo: Learning to navigate through crowded dynamic scenes using velocity obstacles,","venue":null,"work_id":"f70a26f3-8c4d-4acd-957e-0a6c64944b40","year":2023},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:49.427620Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:a1bcd424a3f0c34691d5d8ebe61e4f33cdcb808d89f9774780802cf8573f52cc","observation_id":"b2b485ca-9249-46b1-8473-6464a8bded44","resolution":{"observed_at":"2026-08-07T11:33:52.209570Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:51.831966Z","title":"Efficient deep reinforcement learning with imitative expert priors for autonomous driving,","venue":null,"work_id":"370cc653-03e1-445d-979a-0f4f3d7aac95","year":2022},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:49.525766Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:f01d2d61a3c5fd7916d2c3715a069e365c4c67aceb5107d674c8c187d32b23f5","observation_id":"40c9287a-bc88-4cbb-8fe5-282da52f7ce6","resolution":{"observed_at":"2026-08-07T11:33:51.950891Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:51.572860Z","title":"Pre-training goal-based models for sample-efficient reinforcement learning,","venue":null,"work_id":"570891b5-3be6-4158-9000-4b69fbc0f66d","year":2024},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:49.640866Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:d125d3458fe601b535a89a1c45e6060ad571675a58adcde8228403bc0216bdaa","observation_id":"7a35d118-85a1-466f-952a-eb6c61418154","resolution":{"observed_at":"2026-08-07T11:33:51.691359Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:51.292026Z","title":"Deep q- learning from demonstrations,","venue":null,"work_id":"897085f5-f448-4737-bafd-f2014f501ee3","year":2018},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:49.819870Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:25536847a7f0c4cb1bd6e9bc20451656c9928269912a47ae5d23a5df8df691ff","observation_id":"80730197-6df9-4edf-a685-df9ee676e76f","resolution":{"observed_at":"2026-08-07T11:33:51.413355Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:49.972886Z","title":"Soft actor-critic: Off- policy maximum entropy deep reinforcement learning with a stochastic actor,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:49.972886Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:77d0a6255d9ea47ccc5ab80ab59afeea01e880c281b96f3470a33c2ea3f57f0b","observation_id":"9c54604b-f17c-43a5-8cb0-8778ed0b5ecc","resolution":{"observed_at":"2026-08-07T11:33:49.972886Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:51.026719Z","title":"Template model inspired task space learning for robust bipedal locomotion,","venue":null,"work_id":"1940e5c7-5690-4a6a-b5da-8f6231f40c63","year":2023},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:50.057744Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:2982f41fa5f30101e047dc71d45b31cf959ea2f7463b26df88ff899254b6294e","observation_id":"5b7a965f-8aab-4297-bc93-4da4edbc2279","resolution":{"observed_at":"2026-08-07T11:33:51.141854Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:33:50.738584Z","title":"Conservative q- learning for offline reinforcement learning,","venue":null,"work_id":"42452e19-d0f4-4546-a33b-e18ce57d86a2","year":2020},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:50.168382Z"},"links":{"citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:cfd0e97e97f8b3b044d871a75b7ebf7f1aecae1c7aa19b850dbacdc8bad7c345","observation_id":"e4b73840-65c4-4e0a-94c3-20e761fdb8b3","resolution":{"observed_at":"2026-08-07T11:33:50.895121Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-07T11:33:50.317068Z","title":"Proximal policy optimization algorithms,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:50.317068Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2506.02206"},"observation_digest":"sha256:fbdf96ce8f048b3744f81e897ad7bd0f283b9b0bcc1489402632c22d6d6b5924","observation_id":"a809ba34-adc7-45de-8f63-eb05135ab05c","resolution":{"observed_at":"2026-08-07T11:33:50.317068Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.02206","last_updated":"2025-06-02T19:45:29Z","latest_version":1,"primary_category":"cs.RO","snapshot_observed_at":"2026-08-07T11:26:18.757346Z","submitted_at":"2025-06-02T19:45:29Z","title":"Reinforcement Learning with Data Bootstrapping for Dynamic Subgoal Pursuit in Humanoid Robot Navigation"},"reference_resolution":{"displayed":35,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":8,"verified_exact":1,"verified_fuzzy":26},"total_outbound_references":35},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 35 of 35 outbound references and 1 inbound Pith citation observation for arXiv:2506.02206."}