{"as_of":"2026-08-05T04:40:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:c3f69a977b31cdff50d84913119e16cc18b072698e9ae09b845925e26d5b6405","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":34,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":34,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-04T06:34:03.388597+00:00","state":"measured"},{"denominator":34,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":34,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T10:38:16.451292Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"pith","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":1313,"observed_at":"2026-08-05T02:28:24.338817Z","source":"pith"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"1606.06565","last_updated":"2016-07-25T17:23:29Z","snapshot_observed_at":"2026-07-06T05:00:46.434335Z","submitted_at":"2016-06-21T13:37:05Z","title":"Concrete Problems in AI Safety","version":2},"reference_index":129,"source":"pdf_text","source_observed_at":"2026-05-11T05:16:31.970352Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/1606.06565"},"observation_digest":"sha256:96e1a6cfc48e5e1d1240d59233ce121fab75e2c82ae248ca39b8d93fc6f2415d","observation_id":"154771cb-5ad7-4438-ab61-3e40057a11ea","resolution":{"observed_at":"2026-05-11T05:16:33.017036Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"1910.07113","last_updated":"2019-10-16T00:59:05Z","snapshot_observed_at":"2026-08-02T15:37:37.200292Z","submitted_at":"2019-10-16T00:59:05Z","title":"Solving Rubik's Cube with a Robot Hand","version":1},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-05-15T09:38:28.621842Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/1910.07113"},"observation_digest":"sha256:1809bcb973d05d24c3226c2468652a577b7a185a1c26e8184bfe57ee5c4dc2be","observation_id":"1954308c-9a5f-4f5a-aa59-1df9ae89675c","resolution":{"observed_at":"2026-05-15T09:38:28.814290Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2403.01823","last_updated":"2024-06-01T01:54:54Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-03-04T08:16:11Z","title":"RT-H: Action Hierarchies Using Language","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-17T06:53:27.642020Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2403.01823"},"observation_digest":"sha256:f8366c6806173f6621b85d327bd21bb3d38b61a8fb3ed91b60e2d7ecb6c5e6ae","observation_id":"2ad9ed37-4e32-4a28-b3d9-287c131c16fd","resolution":{"observed_at":"2026-05-17T06:53:27.796205Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2411.05174","last_updated":"2026-04-27T19:23:31Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-11-07T20:27:29Z","title":"Bayesian Inverse Transition Learning: Learning Dynamics From Near-Optimal Trajectories","version":2},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-05-23T17:10:04.386499Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2411.05174"},"observation_digest":"sha256:10ccddb62180deccd48a24f8dfe2cf10979ebe2d9157b2fabd27cadf009131fc","observation_id":"3fcf81a1-d13d-4781-a4eb-06512c46e120","resolution":{"observed_at":"2026-05-23T17:13:14.104665Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-04T10:38:16.451292Z","title":"A reduction of imitation learning and structured prediction to no-regret online learning,","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2510.09497","last_updated":"2026-07-11T05:48:58Z","snapshot_observed_at":"2026-08-04T14:05:40.414745Z","submitted_at":"2025-10-10T15:57:09Z","title":"Toward Autonomous Soft Robotic Endovascular Navigation via Imitation Learning","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-04T10:38:16.451292Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2510.09497"},"observation_digest":"sha256:7d65d3e7469aabceda10978469a0f0e4bfd1b862b74aca11cfcaf37227adcecc","observation_id":"51febd5d-11e9-474f-881c-c6357ba8f007","resolution":{"observed_at":"2026-08-04T10:38:16.451292Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-04T09:17:29.817666Z","title":"Richard S","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.16462","last_updated":"2026-06-02T08:32:59Z","snapshot_observed_at":"2026-08-04T09:17:25.948063Z","submitted_at":"2025-10-18T12:03:15Z","title":"Buzz, Choose, Forget: A Meta-Bandit Framework for Bee-Like Decision Making","version":3},"reference_index":2011,"source":"pdf_text","source_observed_at":"2026-08-04T09:17:29.817666Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2510.16462"},"observation_digest":"sha256:76a95b74650bd208c8b9bfc9699e4e1554d2aff809cdcbeb06a67c69b8a9494d","observation_id":"bca3f45d-d9ff-4829-a942-c11a7d93992d","resolution":{"observed_at":"2026-08-04T09:17:29.817666Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2510.17640","last_updated":"2026-04-10T05:06:03Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-10-20T15:21:12Z","title":"RESample: A Robust Data Augmentation Framework via Exploratory Sampling for Robotic Manipulation","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-18T06:10:47.309028Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2510.17640"},"observation_digest":"sha256:44f119ab6a2d251875c3e806a54e7e04c02c919f70f140cb819e8a1184909909","observation_id":"84794da0-134a-459a-9d0f-3b236f35fd2d","resolution":{"observed_at":"2026-05-18T06:10:57.807717Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2512.15692","last_updated":"2025-12-19T18:30:30Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-12-17T18:47:31Z","title":"mimic-video: Video-Action Models for Generalizable Robot Control Beyond VLAs","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-05-15T10:41:00.142543Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2512.15692"},"observation_digest":"sha256:6017975b00c3bcbeee81e2c646f684693d6aa200fa659f6f5dbe211e9b8018c5","observation_id":"d9bde1b3-1677-46f3-9f72-ccf150271699","resolution":{"observed_at":"2026-05-15T10:41:00.225346Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2604.03023","last_updated":"2026-04-03T13:13:47Z","snapshot_observed_at":"2026-07-06T22:52:15.263364Z","submitted_at":"2026-04-03T13:13:47Z","title":"Behavior-Constrained Reinforcement Learning with Receding-Horizon Credit Assignment for High-Performance Control","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-13T19:30:23.447901Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2604.03023"},"observation_digest":"sha256:f44ec1e38c502578f8523d2a618c004595f68f1f1b68e39f55456c119d80134e","observation_id":"87c9f919-3637-4899-813c-ba9e402b14f2","resolution":{"observed_at":"2026-05-13T19:33:09.999619Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2604.07745","last_updated":"2026-04-09T03:03:06Z","snapshot_observed_at":"2026-08-02T14:47:37.509813Z","submitted_at":"2026-04-09T03:03:06Z","title":"The Cartesian Cut in Agentic AI","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-05-10T17:54:45.333038Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2604.07745"},"observation_digest":"sha256:9060959f671f49e852685b18f06562ed95ffd2b7c1ef14f224a3f251f4c928b2","observation_id":"fbcb8e0f-7af0-4df0-ab2d-6d9cddc75c20","resolution":{"observed_at":"2026-05-11T05:51:10.797338Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2604.09118","last_updated":"2026-06-30T16:12:26Z","snapshot_observed_at":"2026-08-02T13:56:03.921867Z","submitted_at":"2026-04-10T08:58:07Z","title":"Efficient Uniform Feasible-Set Sampling for Approximate Linear MPC","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-10T16:43:06.373842Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2604.09118"},"observation_digest":"sha256:afb57e7d35a5ee97cb6815e63747f535fc3bb9b091e62253926f7216e21e5fbc","observation_id":"f20b228c-ab15-4a9d-b617-56345cbe2011","resolution":{"observed_at":"2026-05-11T08:20:58.201575Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2605.03846","last_updated":"2026-05-05T15:13:15Z","snapshot_observed_at":"2026-07-06T23:16:44.876986Z","submitted_at":"2026-05-05T15:13:15Z","title":"SigLoMa: Learning Open-World Quadrupedal Loco-Manipulation from Ego-Centric Vision","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-07T15:41:18.805962Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2605.03846"},"observation_digest":"sha256:add537765bbaab7ababfb95f87537cd10c7fdee6bf18f52c4c9696f8b00cbd47","observation_id":"1f3e9531-6906-4212-9113-1153f7cd3f0d","resolution":{"observed_at":"2026-05-12T00:11:17.789484Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2605.21606","last_updated":"2026-05-20T18:14:03Z","snapshot_observed_at":"2026-07-06T23:32:01.202542Z","submitted_at":"2026-05-20T18:14:03Z","title":"When Are Teacher Tokens Reliable? Position-Weighted On-Policy Self-Distillation for Reasoning","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-22T09:25:37.960991Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2605.21606"},"observation_digest":"sha256:878dc0e0bbd523ea9cb8e00b8928ea786497e4f3547f6da22db13255d200a86c","observation_id":"f4579156-58ab-4351-9302-ea3c8fe35ef8","resolution":{"observed_at":"2026-05-22T09:26:20.631095Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2605.28775","last_updated":"2026-05-27T17:37:00Z","snapshot_observed_at":"2026-07-06T23:38:19.096349Z","submitted_at":"2026-05-27T17:37:00Z","title":"Learn from Weaknesses: Automated Domain Specialization for Small Computer-Use Agents","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-06-29T14:15:55.180284Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2605.28775"},"observation_digest":"sha256:98e1cd66c73abab5eda9da7034ee14b409dd815e6d479c6d36f0640bf38865b9","observation_id":"7f5bd78a-461c-4213-97a9-37a5265fe118","resolution":{"observed_at":"2026-06-29T14:23:30.906824Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2606.03512","last_updated":"2026-06-02T11:29:00Z","snapshot_observed_at":"2026-07-06T23:43:46.458792Z","submitted_at":"2026-06-02T11:29:00Z","title":"SPADE: Sketch-guided Path Planning Augmented with Diffusion Experts","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-06-28T10:12:08.238917Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2606.03512"},"observation_digest":"sha256:65127b1102d4769e1992a63ae74318ca08d295fc94dc4994c3ea1dcfcb0ac39e","observation_id":"2f7a7aa6-75c0-4334-828b-b3c5f0f221c9","resolution":{"observed_at":"2026-07-02T03:16:34.465792Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2606.04248","last_updated":"2026-06-02T21:54:30Z","snapshot_observed_at":"2026-08-01T23:01:34.893208Z","submitted_at":"2026-06-02T21:54:30Z","title":"RSC: Decentralized Rigid Formation Flocking for Large-Scale Swarms via Hybrid Predictive Control and Online Reconfiguration","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-06-28T09:29:27.766536Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2606.04248"},"observation_digest":"sha256:042989598d2b761540e91e859df9b38abed7c7e1435186a07ddcbd83c8577480","observation_id":"d101887a-f2d3-4c4c-9a33-2d1409d84aa5","resolution":{"observed_at":"2026-07-02T04:06:34.931212Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2606.06479","last_updated":"2026-07-24T17:49:34Z","snapshot_observed_at":"2026-08-02T12:20:41.406405Z","submitted_at":"2026-06-04T17:57:33Z","title":"Pretraining Recurrent Networks without Recurrence","version":1},"reference_index":104,"source":"pdf_text","source_observed_at":"2026-06-28T02:09:01.018909Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2606.06479"},"observation_digest":"sha256:d35e9c892590b457b986bb4114ffa5dd4a88b120ab0f2a339bf2934ef7894148","observation_id":"a83e2e89-aa70-422b-9997-df8763af5704","resolution":{"observed_at":"2026-07-02T12:26:56.632544Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-02T12:20:55.871350Z","title":"Gordon, and J","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2606.06479","last_updated":"2026-07-24T17:49:34Z","snapshot_observed_at":"2026-08-02T12:20:41.406405Z","submitted_at":"2026-06-04T17:57:33Z","title":"Pretraining Recurrent Networks without Recurrence","version":2},"reference_index":101,"source":"pdf_text","source_observed_at":"2026-08-02T12:20:55.871350Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2606.06479"},"observation_digest":"sha256:dacd7c27fc7146e6e27aa622dde9ec1a1f745f62b52c48449f032ef1146fcaa3","observation_id":"d20eb23d-d09c-455c-bbac-4ddbe1ef385d","resolution":{"observed_at":"2026-08-02T12:20:55.871350Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2606.06660","last_updated":"2026-06-04T19:09:22Z","snapshot_observed_at":"2026-08-03T09:07:15.597536Z","submitted_at":"2026-06-04T19:09:22Z","title":"AEGIS: A Backup Reflex for Physical AI","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-28T00:54:46.723901Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2606.06660"},"observation_digest":"sha256:d39e23b3e89a31c4502c9d46c2f516cd4279d59a474bf0e2c54040b426ec71ad","observation_id":"af715895-7031-4031-9b36-18acbb3bc1e8","resolution":{"observed_at":"2026-07-02T13:47:00.053310Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-07-12T13:49:10.591494Z","title":null,"venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2606.16447","last_updated":"2026-07-09T08:44:02Z","snapshot_observed_at":"2026-08-03T08:58:34.090002Z","submitted_at":"2026-06-15T09:19:34Z","title":"Training and Evaluating Diffusion Policies with Long Context Lengths","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-07-12T13:49:10.591494Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2606.16447"},"observation_digest":"sha256:7fbfa5e9b43a9d47c0a6c242aadb219aecdde1afabac35d15049cb8d78b7b6ce","observation_id":"6670bb56-507c-4448-810e-4fac93f49641","resolution":{"observed_at":"2026-07-12T13:49:10.591494Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2606.20333","last_updated":"2026-06-18T15:04:47Z","snapshot_observed_at":"2026-07-06T23:55:29.991442Z","submitted_at":"2026-06-18T15:04:47Z","title":"SoftSkill: Behavioral Compression for Contextual Adaptation","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-26T17:05:35.902167Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2606.20333"},"observation_digest":"sha256:2de6a31b906001fe48f8d3554cf5c894c0118457a396696d7adcc4ded144617b","observation_id":"850d3fe9-da0b-493b-86bc-b8f528b7d43b","resolution":{"observed_at":"2026-07-04T04:19:34.460700Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2606.27163","last_updated":"2026-07-18T14:45:34Z","snapshot_observed_at":"2026-08-02T10:06:13.054681Z","submitted_at":"2026-06-25T15:31:23Z","title":"Learning to Fold: prizewinning solution at LeHome Challenge 2026 (1st place online, 2nd offline)","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-26T04:58:57.215350Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2606.27163"},"observation_digest":"sha256:2c5f87e4a83dddd25d1416f7a20286bd64055beb8fbfc8576f09942a7669390b","observation_id":"24e005e7-70f4-4cac-a0fa-c51feaeb8255","resolution":{"observed_at":"2026-07-04T13:39:51.456051Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-02T10:06:14.513790Z","title":null,"venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2606.27163","last_updated":"2026-07-18T14:45:34Z","snapshot_observed_at":"2026-08-02T10:06:13.054681Z","submitted_at":"2026-06-25T15:31:23Z","title":"Learning to Fold: prizewinning solution at LeHome Challenge 2026 (1st place online, 2nd offline)","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-02T10:06:14.513790Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2606.27163"},"observation_digest":"sha256:617e6ad713ae11d72031dadf58e2e8c0aa38b9136825699fab3e25da79f755c5","observation_id":"713df568-eb67-4bd0-9895-4dd05d77c8c6","resolution":{"observed_at":"2026-08-02T10:06:14.513790Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2606.30456","last_updated":"2026-06-29T15:23:34Z","snapshot_observed_at":"2026-08-01T15:04:25.603449Z","submitted_at":"2026-06-29T15:23:34Z","title":"Vision-Language-Action Models: Experimental Insights from a Real-World UR5 Platform","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-30T05:25:16.143939Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2606.30456"},"observation_digest":"sha256:ed28c6055027d491c8767f26a00ff97b711b37602d8e7b1b25d4752440bb73be","observation_id":"fa7599b7-7d88-4fba-85cb-d3c81b9ccb40","resolution":{"observed_at":"2026-06-30T14:34:45.915619Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2607.00678","last_updated":"2026-07-06T08:13:44Z","snapshot_observed_at":"2026-07-12T09:22:06.947773Z","submitted_at":"2026-07-01T09:21:20Z","title":"ABot-M0.5: Unified Mobility-and-Manipulation World Action Model","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-07-02T14:24:23.187164Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2607.00678"},"observation_digest":"sha256:4efce7ac24c1e253efbd762342f2435b165f58059e0779ba6af563f4ab7f8a78","observation_id":"d3ccdca5-1ff2-48e8-ac76-02a5f838be54","resolution":{"observed_at":"2026-07-02T14:27:03.201791Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-07-12T09:22:08.000379Z","title":"Gordon, and J","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2607.00678","last_updated":"2026-07-06T08:13:44Z","snapshot_observed_at":"2026-07-12T09:22:06.947773Z","submitted_at":"2026-07-01T09:21:20Z","title":"ABot-M0.5: Unified Mobility-and-Manipulation World Action Model","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-07-12T09:22:08.000379Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2607.00678"},"observation_digest":"sha256:abd9c0f21b343f76dc97bf5ddb5eadd9ea2a52c1f9479f4f92310d725523bf39","observation_id":"947c624b-8755-4da7-8105-64f13b2ca1c7","resolution":{"observed_at":"2026-07-12T09:22:08.000379Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":"1011.0686","doi":"10.48550/arxiv.1011.0686","metadata_source":"pith","pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","venue":"cs.LG","work_id":"73c02b0d-67bb-4f1e-900f-0241ff680534","year":2010},"citing_paper":{"arxiv_id":"2607.01651","last_updated":"2026-07-02T03:23:40Z","snapshot_observed_at":"2026-07-07T00:07:13.555571Z","submitted_at":"2026-07-02T03:23:40Z","title":"One Demonstration Is Enough for Real-World Robotic Reinforcement Learning","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-07-03T12:40:11.283716Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2607.01651"},"observation_digest":"sha256:6bc787fa02d51d119786e279e6c0db04678dbc3b97d8c78857650dcfc5d70aa8","observation_id":"ede12d47-9489-456d-ba66-b585af81b2bc","resolution":{"observed_at":"2026-07-03T12:48:11.091632Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-07-12T00:23:04.763773Z","title":null,"venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2607.03723","last_updated":"2026-07-04T06:11:03Z","snapshot_observed_at":"2026-08-01T06:38:23.685403Z","submitted_at":"2026-07-04T06:11:03Z","title":"OmniTacTune: Policy-Agnostic Real-World RL for Tactile Residual Adaptation of Visual Policies","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-07-12T00:23:04.763773Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2607.03723"},"observation_digest":"sha256:254c0e15a5ab2ea0860865d2f18b278fefcd682807e6fb24a11a0e612f43e2c6","observation_id":"5f9b8222-4e0d-41d8-93d2-30f738776017","resolution":{"observed_at":"2026-07-12T00:23:04.763773Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-01T12:28:59.768254Z","title":"Gordon, and J","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2607.19548","last_updated":"2026-07-21T19:48:21Z","snapshot_observed_at":"2026-08-01T12:28:47.290765Z","submitted_at":"2026-07-21T19:48:21Z","title":"Agent-Centric Animal Pose Forecasting","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-01T12:28:59.768254Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2607.19548"},"observation_digest":"sha256:433326431f66f519bcb9a07ce5af0686323370f73e8a850750a5d3813adae2d0","observation_id":"470722a1-718c-41e8-b870-0f91a287108d","resolution":{"observed_at":"2026-08-01T12:28:59.768254Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-01T08:12:03.719012Z","title":"A reduction of imitation learning and structured prediction to no-regret online learning","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2607.21227","last_updated":"2026-07-23T11:43:31Z","snapshot_observed_at":"2026-08-04T06:11:48.493805Z","submitted_at":"2026-07-23T11:43:31Z","title":"FORGE-plus: Force-Budgeted Recovery for Contact-Rich Assembly with a Frozen LLM Supervisor","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-01T08:12:03.719012Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2607.21227"},"observation_digest":"sha256:878aaa8c62b7652eced1caaf92116db01a3010344f7ea6f9e1fe5b43c9e36c2e","observation_id":"48097ba1-2f7d-41ab-b993-564b9d655d64","resolution":{"observed_at":"2026-08-01T08:12:03.719012Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-01T04:07:45.244964Z","title":"Gordon and J","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.22972","last_updated":"2026-07-28T05:10:29Z","snapshot_observed_at":"2026-08-01T04:07:43.320153Z","submitted_at":"2026-07-25T00:42:04Z","title":"Learned Interventions in Lean 4 grind","version":2},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-01T04:07:45.244964Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2607.22972"},"observation_digest":"sha256:548651252300be523ac4493e0d836d309d786534bea7b5ae107545ba7dae4411","observation_id":"7bfe5a18-9aca-4f44-9584-99cd15d8bac2","resolution":{"observed_at":"2026-08-01T04:07:45.244964Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-07-30T21:05:52.713852Z","title":"J.; and Bagnell, J","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2607.26789","last_updated":"2026-07-29T11:31:33Z","snapshot_observed_at":"2026-08-02T17:59:26.789390Z","submitted_at":"2026-07-29T11:31:33Z","title":"CheckVLA: Execution-Time Verification with Action-Conditioned World Model for Long-Horizon Mobile Manipulation","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-07-30T21:05:52.713852Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2607.26789"},"observation_digest":"sha256:00fc934e03d50876dfa008403c6c9cef0748f97ffba2de6a40727c21c4dc2c5b","observation_id":"35af02fb-0e0b-4c64-a96d-b85025d5b635","resolution":{"observed_at":"2026-07-30T21:05:52.713852Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-08-01T10:23:23.043807Z","title":"Gordon, and J","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2607.27288","last_updated":"2026-07-29T14:46:04Z","snapshot_observed_at":"2026-08-03T19:05:39.789384Z","submitted_at":"2026-07-29T14:46:04Z","title":"Open Security Benchmark: Towards Autonomous Enterprise Cyber Defense","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-01T10:23:23.043807Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2607.27288"},"observation_digest":"sha256:e4159f485dce15848927172cd87212650151a77950b8f7d94debf8eaa0f9570d","observation_id":"6aab12ec-e19c-4b7b-8dba-0c04339806ad","resolution":{"observed_at":"2026-08-01T10:23:23.043807Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1011.0686","snapshot_observed_at":"2026-07-31T23:37:14.917770Z","title":null,"venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2607.27890","last_updated":"2026-07-30T09:05:04Z","snapshot_observed_at":"2026-08-03T02:07:52.914193Z","submitted_at":"2026-07-30T09:05:04Z","title":"Static In, Dynamic Out: Counterfactual Action Augmentation for Moving Object Manipulation","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-07-31T23:37:14.917770Z"},"links":{"cited_paper":"/paper/1011.0686","citing_paper":"/paper/2607.27890"},"observation_digest":"sha256:9056716e2804fb7543cbe5d81a4a275107147796ea712a8f455209946cc822d0","observation_id":"81ebb962-f3cf-453e-bafb-733a638a618b","resolution":{"observed_at":"2026-07-31T23:37:14.917770Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/1011.0686/citation-record","integrity":"/paper/1011.0686/integrity","json":"/paper/1011.0686/citation-record.json","paper":"/paper/1011.0686"},"outbound":[],"paper":{"arxiv_id":"1011.0686","last_updated":"2011-03-16T18:51:21Z","latest_version":3,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-03T22:17:51.826967Z","submitted_at":"2010-11-02T17:55:55Z","title":"A Reduction of Imitation Learning and Structured Prediction to No-Regret Online Learning"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 34 inbound Pith citation observations for arXiv:1011.0686."}