{"as_of":"2026-08-07T10:14:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:58f555d6321e92b7d6d99c37cabde7f8997dc81804e3fc8bee36dd21b12b3539","coverage":[{"denominator":56,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":56,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T17:57:52.560008Z","state":"measured"},{"denominator":66,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":66,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":10,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":10,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T08:55:58.681765Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2509.14257","snapshot_observed_at":"2026-08-04T08:55:58.681765Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.18383","last_updated":"2026-06-18T07:59:59Z","snapshot_observed_at":"2026-08-06T22:36:35.328841Z","submitted_at":"2025-10-21T08:03:14Z","title":"MENTOR: Reinforcement Learning via Flexible Teacher-Optimized Rewards for Tool-Use Distillation","version":3},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-04T08:55:58.681765Z"},"links":{"cited_paper":"/paper/2509.14257","citing_paper":"/paper/2510.18383"},"observation_digest":"sha256:22f1142b6f6d3775f9fbf482f6a89b1e51fb0d8be9990b140011e9296b5fd177","observation_id":"ffdab585-71ee-4ecd-a1da-ae45c1f93a2f","resolution":{"observed_at":"2026-08-04T08:55:58.681765Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"cited_work":{"arxiv_id":"2509.14257","doi":"10.48550/arxiv.2509.14257","metadata_source":"arxiv_reference","pith_arxiv_id":"2509.14257","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Ethan Mendes, Jungsoo Park, and Alan Ritter","venue":"ArXiv.org","work_id":"1172118c-c70c-4d53-922d-4308b0d0e45d","year":2025},"citing_paper":{"arxiv_id":"2604.00626","last_updated":"2026-06-18T17:00:51Z","snapshot_observed_at":"2026-07-13T14:57:31.836214Z","submitted_at":"2026-04-01T08:32:34Z","title":"A Survey of On-Policy Distillation for Large Language Models","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-13T23:03:25.985040Z"},"links":{"cited_paper":"/paper/2509.14257","citing_paper":"/paper/2604.00626"},"observation_digest":"sha256:a57d7ce84861750a73dcb36c9d13bc97cc3e7cb64c711e8456c30be80f8ee671","observation_id":"987b9a5c-e292-430d-b09c-49790ad22543","resolution":{"observed_at":"2026-07-24T02:23:49.832269Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"cited_work":{"arxiv_id":"2509.14257","doi":"10.48550/arxiv.2509.14257","metadata_source":"arxiv_reference","pith_arxiv_id":"2509.14257","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Ethan Mendes, Jungsoo Park, and Alan Ritter","venue":"ArXiv.org","work_id":"1172118c-c70c-4d53-922d-4308b0d0e45d","year":2025},"citing_paper":{"arxiv_id":"2604.00626","last_updated":"2026-06-18T17:00:51Z","snapshot_observed_at":"2026-07-13T14:57:31.836214Z","submitted_at":"2026-04-01T08:32:34Z","title":"A Survey of On-Policy Distillation for Large Language Models","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-21T09:38:56.224611Z"},"links":{"cited_paper":"/paper/2509.14257","citing_paper":"/paper/2604.00626"},"observation_digest":"sha256:8026355cc520d1fdad1ed3fdd6683ed96fe7e7d0bc5ce125062ef8e2bb112293","observation_id":"b1fb5a89-5623-4695-9534-a39089d2948d","resolution":{"observed_at":"2026-07-24T02:23:49.832269Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"cited_work":{"arxiv_id":"2509.14257","doi":"10.48550/arxiv.2509.14257","metadata_source":"arxiv_reference","pith_arxiv_id":"2509.14257","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Ethan Mendes, Jungsoo Park, and Alan Ritter","venue":"ArXiv.org","work_id":"1172118c-c70c-4d53-922d-4308b0d0e45d","year":2025},"citing_paper":{"arxiv_id":"2604.21590","last_updated":"2026-04-23T12:14:52Z","snapshot_observed_at":"2026-07-06T23:08:10.140279Z","submitted_at":"2026-04-23T12:14:52Z","title":"AgenticQwen: Training Small Agentic Language Models with Dual Data Flywheels for Industrial-Scale Tool Use","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-09T21:44:18.943531Z"},"links":{"cited_paper":"/paper/2509.14257","citing_paper":"/paper/2604.21590"},"observation_digest":"sha256:21e51385d2f7320fc2ce5a00cd6fc9eb9d79de270547f716d0d0e770e4d64700","observation_id":"d6cc47e8-8d26-4dcd-afa0-bda9dba8250f","resolution":{"observed_at":"2026-07-24T02:23:49.832269Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"cited_work":{"arxiv_id":"2509.14257","doi":"10.48550/arxiv.2509.14257","metadata_source":"arxiv_reference","pith_arxiv_id":"2509.14257","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Ethan Mendes, Jungsoo Park, and Alan Ritter","venue":"ArXiv.org","work_id":"1172118c-c70c-4d53-922d-4308b0d0e45d","year":2025},"citing_paper":{"arxiv_id":"2605.07276","last_updated":"2026-05-08T05:41:25Z","snapshot_observed_at":"2026-07-06T23:19:41.053744Z","submitted_at":"2026-05-08T05:41:25Z","title":"Signal Reshaping for GRPO in Weak-Feedback Agentic Code Repair","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-11T01:21:43.166304Z"},"links":{"cited_paper":"/paper/2509.14257","citing_paper":"/paper/2605.07276"},"observation_digest":"sha256:feabbebd66ab4550584a3140df7b7c412bdb7cf6a461f37a9653e1a7f0f4a39f","observation_id":"36cb99ce-cc63-4aac-9c49-6454a8c6a77a","resolution":{"observed_at":"2026-07-24T02:23:49.832269Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"cited_work":{"arxiv_id":"2509.14257","doi":"10.48550/arxiv.2509.14257","metadata_source":"arxiv_reference","pith_arxiv_id":"2509.14257","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Ethan Mendes, Jungsoo Park, and Alan Ritter","venue":"ArXiv.org","work_id":"1172118c-c70c-4d53-922d-4308b0d0e45d","year":2025},"citing_paper":{"arxiv_id":"2605.12652","last_updated":"2026-06-01T07:33:52Z","snapshot_observed_at":"2026-08-04T16:10:21.909941Z","submitted_at":"2026-05-12T18:57:44Z","title":"Multi-Rollout On-Policy Distillation via Peer Successes and Failures","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-05-14T21:29:36.803832Z"},"links":{"cited_paper":"/paper/2509.14257","citing_paper":"/paper/2605.12652"},"observation_digest":"sha256:45cac350bb0899a9706b0bd70e7355f784416e7c9cc7c4ca690adeb1d3209c97","observation_id":"94d5937a-bd63-46a5-bda4-9d1dd4795149","resolution":{"observed_at":"2026-07-24T02:23:49.832269Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"cited_work":{"arxiv_id":"2509.14257","doi":"10.48550/arxiv.2509.14257","metadata_source":"arxiv_reference","pith_arxiv_id":"2509.14257","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Ethan Mendes, Jungsoo Park, and Alan Ritter","venue":"ArXiv.org","work_id":"1172118c-c70c-4d53-922d-4308b0d0e45d","year":2025},"citing_paper":{"arxiv_id":"2605.28775","last_updated":"2026-05-27T17:37:00Z","snapshot_observed_at":"2026-07-06T23:38:19.096349Z","submitted_at":"2026-05-27T17:37:00Z","title":"Learn from Weaknesses: Automated Domain Specialization for Small Computer-Use Agents","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-29T14:15:55.180284Z"},"links":{"cited_paper":"/paper/2509.14257","citing_paper":"/paper/2605.28775"},"observation_digest":"sha256:d3045fe68d09fce55f19c3bc65cd13f5476ef05b3474715b693bf77329127a95","observation_id":"a77d67c0-3513-46fb-9fc7-bce953f564bf","resolution":{"observed_at":"2026-07-24T02:23:49.832269Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"cited_work":{"arxiv_id":"2509.14257","doi":"10.48550/arxiv.2509.14257","metadata_source":"arxiv_reference","pith_arxiv_id":"2509.14257","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Ethan Mendes, Jungsoo Park, and Alan Ritter","venue":"ArXiv.org","work_id":"1172118c-c70c-4d53-922d-4308b0d0e45d","year":2025},"citing_paper":{"arxiv_id":"2606.06840","last_updated":"2026-06-05T02:32:24Z","snapshot_observed_at":"2026-08-05T15:03:33.373071Z","submitted_at":"2026-06-05T02:32:24Z","title":"Characterize Then Distill: Mechanistic Reasoning in Large Output Spaces","version":1},"reference_index":162,"source":"arxiv_source","source_observed_at":"2026-06-27T22:22:52.690010Z"},"links":{"cited_paper":"/paper/2509.14257","citing_paper":"/paper/2606.06840"},"observation_digest":"sha256:bfafb4b25af5da2bf5755b17d954fb42e301dca05ac2f7d6db3e941cd64fdc7a","observation_id":"7f578b70-a8b1-4b60-b38d-d3b8a6f3d3c2","resolution":{"observed_at":"2026-07-24T02:23:49.832269Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"cited_work":{"arxiv_id":"2509.14257","doi":"10.48550/arxiv.2509.14257","metadata_source":"arxiv_reference","pith_arxiv_id":"2509.14257","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Ethan Mendes, Jungsoo Park, and Alan Ritter","venue":"ArXiv.org","work_id":"1172118c-c70c-4d53-922d-4308b0d0e45d","year":2025},"citing_paper":{"arxiv_id":"2606.23104","last_updated":"2026-06-22T09:46:35Z","snapshot_observed_at":"2026-07-06T23:57:53.101659Z","submitted_at":"2026-06-22T09:46:35Z","title":"ReNIO: Reweighting Negative Trajectory Importance for LLM On-Policy Distillation","version":1},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-06-26T09:05:36.937295Z"},"links":{"cited_paper":"/paper/2509.14257","citing_paper":"/paper/2606.23104"},"observation_digest":"sha256:e09ffc3c77ba134a7ac2f4a65d25ce709711ba5fd1616f47e5a26debc5bc94d2","observation_id":"c1ea4989-cf15-4d00-b4e3-aa38b984cd2e","resolution":{"observed_at":"2026-07-24T02:23:49.832269Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2509.14257","snapshot_observed_at":"2026-08-02T14:45:23.260495Z","title":"arXiv preprint arXiv:2509.14257 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.16205","last_updated":"2026-05-07T13:16:09Z","snapshot_observed_at":"2026-08-05T18:20:48.142001Z","submitted_at":"2026-05-07T13:16:09Z","title":"It Takes 8 Tokens: Weak-to-Strong Off-Policy RL via Auxiliary Branches","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-02T14:45:23.260495Z"},"links":{"cited_paper":"/paper/2509.14257","citing_paper":"/paper/2607.16205"},"observation_digest":"sha256:adb50e9b779751477f4f627efb607649f3a23941b3830003c17eb26cf94b8bbf","observation_id":"bf9f0abf-d4de-4a97-8516-c5ccdca81ff3","resolution":{"observed_at":"2026-08-02T14:45:23.260495Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2509.14257/citation-record","integrity":"/paper/2509.14257/integrity","json":"/paper/2509.14257/citation-record.json","paper":"/paper/2509.14257"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-04T17:57:52.308605Z","title":"Gpt-4 technical report","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.308605Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:b78f80a65b4857699bda7eda871882c549f592243076157e45648036e73a05de","observation_id":"bf934b61-98de-482e-8dc9-3620a4c356ab","resolution":{"observed_at":"2026-08-04T17:57:52.308605Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.314451Z","title":"Hindsight experience replay","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.314451Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:d3fef7d7047ce9e841538fe16b6a4782d206e202b6fd97b6ecd852ca2cc18d4b","observation_id":"0f2e19eb-bd03-4ab7-9a0e-bb1247a4e9ea","resolution":{"observed_at":"2026-08-04T17:57:52.314451Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2002.06038","last_updated":"2020-02-14T13:57:22Z","snapshot_observed_at":"2026-07-06T08:57:20.804732Z","submitted_at":"2020-02-14T13:57:22Z","title":"Never Give Up: Learning Directed Exploration Strategies","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2002.06038","snapshot_observed_at":"2026-08-04T17:57:52.319352Z","title":"Never give up: Learning directed exploration strategies","venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.319352Z"},"links":{"cited_paper":"/paper/2002.06038","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:b0d9b94563ad1aaf1d49bc21d1d344e1ac1b3823b043e37afe0748347b80bf14","observation_id":"57244e40-7330-4372-87c8-f93f8b3b48f2","resolution":{"observed_at":"2026-08-04T17:57:52.319352Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.13651","last_updated":"2025-06-16T16:16:14Z","snapshot_observed_at":"2026-08-07T00:25:54.557015Z","submitted_at":"2025-06-16T16:16:14Z","title":"xbench: Tracking Agents Productivity Scaling with Profession-Aligned Real-World Evaluations","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.13651","snapshot_observed_at":"2026-08-04T17:57:52.324482Z","title":"xbench: Tracking agents productivity scaling with profession-aligned real-world evaluations","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.324482Z"},"links":{"cited_paper":"/paper/2506.13651","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:33d98aca5cdff0a545701379fcca0f2bac851a4fde68fe94d6da40bbca8ea829","observation_id":"04703e8a-6d7a-439f-bb91-86c287316e75","resolution":{"observed_at":"2026-08-04T17:57:52.324482Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.08691","last_updated":"2023-07-17T17:50:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-07-17T17:50:36Z","title":"FlashAttention-2: Faster Attention with Better Parallelism and Work Partitioning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.08691","snapshot_observed_at":"2026-08-04T17:57:52.329389Z","title":"Flashattention-2: Faster attention with better parallelism and work partitioning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.329389Z"},"links":{"cited_paper":"/paper/2307.08691","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:7b0d7595199a50e481cfb8b088760714857d1200113056510b8ce838a5740411","observation_id":"4f786fea-c0cc-4c9c-93a5-87c2ed7bcba2","resolution":{"observed_at":"2026-08-04T17:57:52.329389Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.16410","last_updated":"2025-05-22T09:00:19Z","snapshot_observed_at":"2026-07-06T21:28:24.644692Z","submitted_at":"2025-05-22T09:00:19Z","title":"Tool-Star: Empowering LLM-Brained Multi-Tool Reasoner via Reinforcement Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.16410","snapshot_observed_at":"2026-08-04T17:57:52.334039Z","title":"Tool-star: Empowering llm-brained multi-tool reasoner via reinforcement learning","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.334039Z"},"links":{"cited_paper":"/paper/2505.16410","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:44f87c84f708b10240fca3d7bbd62ee1865125494ffe93860776a32ca19364f6","observation_id":"4cf73df6-8d84-4a8c-9504-70b53a7eb141","resolution":{"observed_at":"2026-08-04T17:57:52.334039Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.19849","last_updated":"2025-07-26T07:53:11Z","snapshot_observed_at":"2026-07-06T22:03:15.296567Z","submitted_at":"2025-07-26T07:53:11Z","title":"Agentic Reinforced Policy Optimization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.19849","snapshot_observed_at":"2026-08-04T17:57:52.339239Z","title":"Agentic reinforced policy optimization","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.339239Z"},"links":{"cited_paper":"/paper/2507.19849","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:8b278d737a8f9c6e31865688e4ff7c53e42be98714ebb53db1b413c70486d7a8","observation_id":"807572e7-0dd9-4866-b438-53c0622dffce","resolution":{"observed_at":"2026-08-04T17:57:52.339239Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.343851Z","title":"The llama 3 herd of models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.343851Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:35da40993e65950c08856fc20457bd8ce3d29edbb128d4915e4b6303edb4be1d","observation_id":"98eddb0d-1d45-46d3-8e6e-875beae31c78","resolution":{"observed_at":"2026-08-04T17:57:52.343851Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.07985","last_updated":"2024-12-24T04:04:30Z","snapshot_observed_at":"2026-08-05T09:35:42.754650Z","submitted_at":"2024-10-10T14:39:33Z","title":"Omni-MATH: A Universal Olympiad Level Mathematic Benchmark For Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.07985","snapshot_observed_at":"2026-08-04T17:57:52.348384Z","title":"Omni-math: A universal olympiad level mathematic benchmark for large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.348384Z"},"links":{"cited_paper":"/paper/2410.07985","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:70d1d3762deee126ca5e51a37fc5620cabd2b7fe02e44f9c00dda9bbadb4c95d","observation_id":"4c49559d-194a-47fc-bfaf-fecdf23671ac","resolution":{"observed_at":"2026-08-04T17:57:52.348384Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.353164Z","title":"Pal: Program-aided language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.353164Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:466310805e095e7cb2126ab95e8bc389c2cfc9c96999317f53b0ba45ed0cb687","observation_id":"1406dc90-60c5-4cc4-876e-a1602fe7f0c1","resolution":{"observed_at":"2026-08-04T17:57:52.353164Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2103.03874","last_updated":"2021-11-08T21:30:18Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-03-05T18:59:39Z","title":"Measuring Mathematical Problem Solving With the MATH Dataset","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2103.03874","snapshot_observed_at":"2026-08-04T17:57:52.360060Z","title":"Measuring mathematical problem solving with the math dataset","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.360060Z"},"links":{"cited_paper":"/paper/2103.03874","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:845209d2462ac9031d082a5d219bfa78cb96d83bffb1c6291fe30c957854719a","observation_id":"ca7baf44-8589-4a52-aa0a-bd624e934b29","resolution":{"observed_at":"2026-08-04T17:57:52.360060Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2011.01060","last_updated":"2020-11-12T07:47:48Z","snapshot_observed_at":"2026-07-06T10:10:56.466018Z","submitted_at":"2020-11-02T15:42:40Z","title":"Constructing A Multi-hop QA Dataset for Comprehensive Evaluation of Reasoning Steps","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2011.01060","snapshot_observed_at":"2026-08-04T17:57:52.365586Z","title":"Constructing a multi-hop qa dataset for comprehensive evaluation of reasoning steps","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.365586Z"},"links":{"cited_paper":"/paper/2011.01060","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:98481b783da04c0b6402d77570751aa99242a3e34d6778d5badb4637c324233a","observation_id":"2f9bd869-8670-4d95-b032-6ec5bd7a3aa5","resolution":{"observed_at":"2026-08-04T17:57:52.365586Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.370476Z","title":"Distilling llm agent into small models with retrieval and code tools","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.370476Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:40f522ca20fe7632a6d156f88d35e2ae286d7dfee6ac1ac9843c6ba8b10fd770","observation_id":"69e30052-1d97-4d60-91f0-ae1a29c32d63","resolution":{"observed_at":"2026-08-04T17:57:52.370476Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.374888Z","title":"Hg-dagger: Interactive imitation learning with human experts","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.374888Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:18342cfe486b0ea4292125bb1694acc57de727ce79a5dfbb78d88cbc20a723c3","observation_id":"83eadfe2-c2c6-4391-ba8b-148d3fd60ed2","resolution":{"observed_at":"2026-08-04T17:57:52.374888Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.379152Z","title":"Numinamath: The largest public dataset in ai4maths with 860k pairs of competition math problems and solutions","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.379152Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:dbe5241da4039153db39eb6647e6892d688b2d93f9301ad37bd8108977c185cd","observation_id":"1f512c19-b1a2-439d-a731-67425fd2c724","resolution":{"observed_at":"2026-08-04T17:57:52.379152Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.13167","last_updated":"2025-08-06T17:01:02Z","snapshot_observed_at":"2026-08-05T23:57:29.481263Z","submitted_at":"2025-08-06T17:01:02Z","title":"Chain-of-Agents: End-to-End Agent Foundation Models via Multi-Agent Distillation and Agentic RL","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.13167","snapshot_observed_at":"2026-08-04T17:57:52.383322Z","title":"Chain-of-agents: End-to-end agent foundation models via multi-agent distillation and agentic rl","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.383322Z"},"links":{"cited_paper":"/paper/2508.13167","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:f9c8c55073d7dbaa561838efba9484512fc6054ab64e01e81205d931aede5a1c","observation_id":"f01437bd-012f-4fca-8e25-4f829e0cd832","resolution":{"observed_at":"2026-08-04T17:57:52.383322Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.05366","last_updated":"2025-01-09T16:48:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-09T16:48:17Z","title":"Search-o1: Agentic Search-Enhanced Large Reasoning Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.05366","snapshot_observed_at":"2026-08-04T17:57:52.387772Z","title":"Search-o1: Agentic search-enhanced large reasoning models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.387772Z"},"links":{"cited_paper":"/paper/2501.05366","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:a583832a24ee9df06411f3fab87ea454e8eeb13c097f7e516829f36bae8ee117","observation_id":"a51a0c04-09a2-44d8-8b12-5c5bb94f7b2c","resolution":{"observed_at":"2026-08-04T17:57:52.387772Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.21776","last_updated":"2025-10-13T12:40:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-04-30T16:25:25Z","title":"WebThinker: Empowering Large Reasoning Models with Deep Research Capability","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.21776","snapshot_observed_at":"2026-08-04T17:57:52.392186Z","title":"Webthinker: Empowering large reasoning models with deep research capability","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.392186Z"},"links":{"cited_paper":"/paper/2504.21776","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:81b489672dfbd591b1a27d702b866c7815eb1b3b771308b1dd12e70e716c61f3","observation_id":"06ae42c6-e0bf-465c-9288-50a61c21db0a","resolution":{"observed_at":"2026-08-04T17:57:52.392186Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.396713Z","title":"Let's verify step by step","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.396713Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:322aa3ac6003b684f658f79e3881e58ea52f834af9ff6034017b4ed0b1abb8ae","observation_id":"6a0f7ff5-e9e5-4532-9fb5-ae1ef0fd9bb1","resolution":{"observed_at":"2026-08-04T17:57:52.396713Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.03688","last_updated":"2025-10-04T03:54:18Z","snapshot_observed_at":"2026-08-06T20:36:41.418114Z","submitted_at":"2023-08-07T16:08:11Z","title":"AgentBench: Evaluating LLMs as Agents","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.03688","snapshot_observed_at":"2026-08-04T17:57:52.400701Z","title":"Agentbench: Evaluating llms as agents","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.400701Z"},"links":{"cited_paper":"/paper/2308.03688","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:af1cf5821405ac01b526bd8fba8730145302cb56de423613ea1bee2328d73e0c","observation_id":"7b9527c7-e2f1-4979-888d-b49b4f3a3463","resolution":{"observed_at":"2026-08-04T17:57:52.400701Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2212.08410","last_updated":"2023-06-01T12:17:01Z","snapshot_observed_at":"2026-07-06T14:31:38.001143Z","submitted_at":"2022-12-16T11:24:42Z","title":"Teaching Small Language Models to Reason","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2212.08410","snapshot_observed_at":"2026-08-04T17:57:52.405123Z","title":"Teaching small language models to reason","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.405123Z"},"links":{"cited_paper":"/paper/2212.08410","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:7ca33040fe778896d4119880dc6b84b012411cbec7c7df0dd43893214ef6397c","observation_id":"57fad5b4-b2a1-446b-9f23-337d06031d0e","resolution":{"observed_at":"2026-08-04T17:57:52.405123Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.409574Z","title":"Gaia: a benchmark for general ai assistants","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.409574Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:41469ada971a507362722c258be1d6d4642c878d26a63d240c6edc918ad56e00","observation_id":"5cba5e94-ac39-4c7f-8ff6-ada03d3b1e3a","resolution":{"observed_at":"2026-08-04T17:57:52.409574Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.413813Z","title":"Human-level control through deep reinforcement learning","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.413813Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:04466f91d11d7d4dcac9bf2f4f8755d151cc74476a80aea01bbee07c51c270e5","observation_id":"55287cad-5aba-4e9c-a815-0350c9228ed5","resolution":{"observed_at":"2026-08-04T17:57:52.413813Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.418083Z","title":"Reinforcement learning with verifiable rewards: Grpo's effective loss, dynamics, and success amplification","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.418083Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:b1c9f044e74b140896377e04c4c85224a79efd79a18c5f5704052e945aece0ae","observation_id":"d4e68983-782a-4bf9-a99d-1d3c1c1739be","resolution":{"observed_at":"2026-08-04T17:57:52.418083Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1910.00177","last_updated":"2019-10-07T20:23:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2019-10-01T02:23:38Z","title":"Advantage-Weighted Regression: Simple and Scalable Off-Policy Reinforcement Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1910.00177","snapshot_observed_at":"2026-08-04T17:57:52.422210Z","title":"Advantage-weighted regression: Simple and scalable off-policy reinforcement learning","venue":null,"work_id":null,"year":1910},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.422210Z"},"links":{"cited_paper":"/paper/1910.00177","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:e23105980fefa01de481c3ae5107f7757735a4e994256e93c85175b014e1c983","observation_id":"62bf2a10-2a36-4425-9470-ae9a0d2a3da1","resolution":{"observed_at":"2026-08-04T17:57:52.422210Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.14249","last_updated":"2026-02-20T04:23:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-24T05:27:46Z","title":"Humanity's Last Exam","version":10},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.14249","snapshot_observed_at":"2026-08-04T17:57:52.426946Z","title":"Humanity's last exam","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.426946Z"},"links":{"cited_paper":"/paper/2501.14249","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:da1d786f5d5a09827924e281e3c1e0d92717486614541079a496314b63617e7a","observation_id":"4601687e-6a9e-4c99-a0d8-9644cd2d05a7","resolution":{"observed_at":"2026-08-04T17:57:52.426946Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2210.03350","last_updated":"2023-10-17T18:57:17Z","snapshot_observed_at":"2026-07-31T15:52:06.089691Z","submitted_at":"2022-10-07T06:50:23Z","title":"Measuring and Narrowing the Compositionality Gap in Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.03350","snapshot_observed_at":"2026-08-04T17:57:52.431712Z","title":"Measuring and narrowing the compositionality gap in language models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.431712Z"},"links":{"cited_paper":"/paper/2210.03350","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:99e992b57c52203d9d8e35d28f283d5da2bd088f740bf22e6808062797c812ca","observation_id":"9a3d88dd-2261-4a24-bb87-afbe195bc313","resolution":{"observed_at":"2026-08-04T17:57:52.431712Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.13958","last_updated":"2025-04-16T21:45:32Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-04-16T21:45:32Z","title":"ToolRL: Reward is All Tool Learning Needs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.13958","snapshot_observed_at":"2026-08-04T17:57:52.436079Z","title":"Toolrl: Reward is all tool learning needs","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.436079Z"},"links":{"cited_paper":"/paper/2504.13958","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:c1f99dab2d15b3c9d36fee3f7295ae6ef6cf737911904495ced7d5ce5c2a4f76","observation_id":"e3312e06-c97b-4173-8c06-9e9d86958a64","resolution":{"observed_at":"2026-08-04T17:57:52.436079Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.440507Z","title":"Direct preference optimization: Your language model is secretly a reward model","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.440507Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:8f80acc0df255714f8344fdbfb652585c57b5f40fdf2388884f57f4504ba1c4b","observation_id":"33d39b98-445d-4493-8b9a-9aaabce26d36","resolution":{"observed_at":"2026-08-04T17:57:52.440507Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.447756Z","title":"Deepspeed: System optimizations enable training deep learning models with over 100 billion parameters","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.447756Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:80c20094e0e11168e9cfdd334ebf2407c2642aaa76200d2daff1b3161fdf33c5","observation_id":"79c66ba8-9bad-4789-a219-54bbf0ca558a","resolution":{"observed_at":"2026-08-04T17:57:52.447756Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.452032Z","title":"A reduction of imitation learning and structured prediction to no-regret online learning","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.452032Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:f56126a7093ab211d71e98d150f496003ade13919f3bbe4767c77fab115e367f","observation_id":"164402f2-3299-43c1-b303-8cf8ddab24e9","resolution":{"observed_at":"2026-08-04T17:57:52.452032Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.456062Z","title":"Toolformer: Language models can teach themselves to use tools","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.456062Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:a5e0d3ad86dddb253dfd216c3410bf8df7d2dc049a35c3e14bf2049dfe714714","observation_id":"8eca317d-13f0-4273-a6a5-b044e6693dd4","resolution":{"observed_at":"2026-08-04T17:57:52.456062Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-04T17:57:52.460072Z","title":"Proximal policy optimization algorithms","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.460072Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:a65cd43e3a7d2ee866169f8813e30b0f146eb67f31707f55ba062cfe553b9d2a","observation_id":"f8381847-50c3-4d6f-913e-48daffb402aa","resolution":{"observed_at":"2026-08-04T17:57:52.460072Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-04T17:57:52.464074Z","title":"Deepseekmath: Pushing the limits of mathematical reasoning in open language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.464074Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:4efe81e77f6bb23ee759746c44f454330e1dd051b60ce8cdbb49880a7b5c0273","observation_id":"e44d88c6-73b5-4356-a733-9ed466d442a8","resolution":{"observed_at":"2026-08-04T17:57:52.464074Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.19256","last_updated":"2024-10-02T04:01:47Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-09-28T06:20:03Z","title":"HybridFlow: A Flexible and Efficient RLHF Framework","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.19256","snapshot_observed_at":"2026-08-04T17:57:52.468486Z","title":"Hybridflow: A flexible and efficient rlhf framework","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.468486Z"},"links":{"cited_paper":"/paper/2409.19256","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:e3e48ca699cb718666a37ed57385c354f523acbe7b9756e58b0c32b63cc7cac8","observation_id":"f725c070-bc1b-45c4-aef5-e5d86e2420cc","resolution":{"observed_at":"2026-08-04T17:57:52.468486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.03768","last_updated":"2021-03-14T22:44:38Z","snapshot_observed_at":"2026-07-06T10:02:33.297722Z","submitted_at":"2020-10-08T05:13:36Z","title":"ALFWorld: Aligning Text and Embodied Environments for Interactive Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.03768","snapshot_observed_at":"2026-08-04T17:57:52.472836Z","title":"Alfworld: Aligning text and embodied environments for interactive learning","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.472836Z"},"links":{"cited_paper":"/paper/2010.03768","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:bee7855adf8c37ce5d6b4bb1e190ea665a61e63bdb6e145376d1ee3e1d1ca2cb","observation_id":"2c9574c4-9b24-4c40-81cd-23185e7968c7","resolution":{"observed_at":"2026-08-04T17:57:52.472836Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1712.01815","last_updated":"2017-12-05T18:45:38Z","snapshot_observed_at":"2026-08-02T00:39:04.960144Z","submitted_at":"2017-12-05T18:45:38Z","title":"Mastering Chess and Shogi by Self-Play with a General Reinforcement Learning Algorithm","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1712.01815","snapshot_observed_at":"2026-08-04T17:57:52.477057Z","title":"Mastering chess and shogi by self-play with a general reinforcement learning algorithm","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.477057Z"},"links":{"cited_paper":"/paper/1712.01815","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:c02c87d6ae275323152d3f98a4610544561690f3751e0368f18d6420bc592109","observation_id":"525fa1de-ba25-4a75-9403-e531e4b1fd88","resolution":{"observed_at":"2026-08-04T17:57:52.477057Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.23829","last_updated":"2025-04-01T14:48:02Z","snapshot_observed_at":"2026-07-06T21:01:17.931818Z","submitted_at":"2025-03-31T08:22:49Z","title":"Crossing the Reward Bridge: Expanding RL with Verifiable Rewards Across Diverse Domains","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.23829","snapshot_observed_at":"2026-08-04T17:57:52.481344Z","title":"Crossing the reward bridge: Expanding rl with verifiable rewards across diverse domains","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.481344Z"},"links":{"cited_paper":"/paper/2503.23829","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:6b6a7a27fab8deef1429db2c7eeb688e54aa92304016e7a65274e2847f5cb141","observation_id":"500c22dd-a511-42f3-b0b0-879be3b0da6b","resolution":{"observed_at":"2026-08-04T17:57:52.481344Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.21380","last_updated":"2026-04-12T10:37:12Z","snapshot_observed_at":"2026-08-02T16:04:50.458535Z","submitted_at":"2025-03-27T11:20:17Z","title":"Challenging the Boundaries of Reasoning: An Olympiad-Level Math Benchmark for Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.21380","snapshot_observed_at":"2026-08-04T17:57:52.485589Z","title":"Challenging the boundaries of reasoning: An olympiad-level math benchmark for large language models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.485589Z"},"links":{"cited_paper":"/paper/2503.21380","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:c534a43706873d86f179572aa4a6bb36f7c0745059c99d4930841a88db08e154","observation_id":"1655ced4-9cb7-440e-9fb4-69c8622c4622","resolution":{"observed_at":"2026-08-04T17:57:52.485589Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.11805","last_updated":"2025-05-09T21:04:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-19T02:39:27Z","title":"Gemini: A Family of Highly Capable Multimodal Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.11805","snapshot_observed_at":"2026-08-04T17:57:52.489807Z","title":"Gemini: a family of highly capable multimodal models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.489807Z"},"links":{"cited_paper":"/paper/2312.11805","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:f41eb7629af6f99fbac67e8dd19bcd1634394f3a2b79ae0116ce03c573304c55","observation_id":"2064dbd7-f8f9-41b5-a5f6-49ac2eeeb1fe","resolution":{"observed_at":"2026-08-04T17:57:52.489807Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1805.01954","last_updated":"2018-05-11T21:48:52Z","snapshot_observed_at":"2026-08-06T03:56:45.024415Z","submitted_at":"2018-05-04T22:36:58Z","title":"Behavioral Cloning from Observation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1805.01954","snapshot_observed_at":"2026-08-04T17:57:52.494499Z","title":"Behavioral cloning from observation","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.494499Z"},"links":{"cited_paper":"/paper/1805.01954","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:e31a39ca6a5b750392c7f48ffa92fe7ba45ea299a2c20dce48c7fac9d2677230","observation_id":"662435d7-a9ed-4709-9f24-865ebe026354","resolution":{"observed_at":"2026-08-04T17:57:52.494499Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.498831Z","title":"Musique: Multihop questions via single-hop question composition","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.498831Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:a7ebde268e3987acfc3028f76baf634cf39e516b6851afbcc2fb2671c436e88b","observation_id":"304a4e7b-ce8b-4715-a06f-7b67feae6d46","resolution":{"observed_at":"2026-08-04T17:57:52.498831Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.503334Z","title":"Executable code actions elicit better llm agents","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.503334Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:949455df004d931526f1bbd6aaae5e6e1d6309efca5f543eca7e4a2e43a91966","observation_id":"41fc463c-a67f-4a20-8bd0-844b196a04dd","resolution":{"observed_at":"2026-08-04T17:57:52.503334Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.07682","last_updated":"2022-10-26T05:06:24Z","snapshot_observed_at":"2026-08-02T15:56:35.249569Z","submitted_at":"2022-06-15T17:32:01Z","title":"Emergent Abilities of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.07682","snapshot_observed_at":"2026-08-04T17:57:52.507655Z","title":"Emergent abilities of large language models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.507655Z"},"links":{"cited_paper":"/paper/2206.07682","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:04882288cc8708fbf113816f6601fab392501403b16093b87d3ba8d6abe5be9c","observation_id":"4fecf760-912d-4157-8b5f-77a225f22400","resolution":{"observed_at":"2026-08-04T17:57:52.507655Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-04T17:57:52.512404Z","title":"Webwalker: Benchmarking llms in web traversal","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.512404Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:aaa34650a2f5a32a4405576b16a3b3af8fc4c22b6cfbf083bc0289dde80983d9","observation_id":"5dad2497-2441-41ef-86a2-ce2b2e9472e2","resolution":{"observed_at":"2026-08-04T17:57:52.512404Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.516462Z","title":"The rise and potential of large language model based agents: A survey","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.516462Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:64de2cd6b064d0c93180ed7d0b6d4450d34242e310fee36616f1f7b449d283e9","observation_id":"5c599acf-6fbf-4376-9f25-9cca50d56763","resolution":{"observed_at":"2026-08-04T17:57:52.516462Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.12594","last_updated":"2025-06-14T18:19:05Z","snapshot_observed_at":"2026-08-07T00:43:03.992583Z","submitted_at":"2025-06-14T18:19:05Z","title":"A Comprehensive Survey of Deep Research: Systems, Methodologies, and Applications","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.12594","snapshot_observed_at":"2026-08-04T17:57:52.520421Z","title":"A comprehensive survey of deep research: Systems, methodologies, and applications","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.520421Z"},"links":{"cited_paper":"/paper/2506.12594","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:faa93f525862b4326510da93fec33b282b76dd1487a74a3e8dffdf6c4e0934d4","observation_id":"58435922-40f8-4a75-8f4c-0ea8d4098e04","resolution":{"observed_at":"2026-08-04T17:57:52.520421Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-08-04T17:57:52.524678Z","title":"Qwen3 technical report","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.524678Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:4add5170a0a02e6d4fde0a8b940f6eec51a47bd45e446351a52dd8932e935adf","observation_id":"bbb116b9-ce58-4b16-8f9c-5773091e34a0","resolution":{"observed_at":"2026-08-04T17:57:52.524678Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1809.09600","last_updated":"2018-09-25T17:28:20Z","snapshot_observed_at":"2026-07-31T11:19:06.421362Z","submitted_at":"2018-09-25T17:28:20Z","title":"HotpotQA: A Dataset for Diverse, Explainable Multi-hop Question Answering","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1809.09600","snapshot_observed_at":"2026-08-04T17:57:52.529020Z","title":"Hotpotqa: A dataset for diverse, explainable multi-hop question answering","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.529020Z"},"links":{"cited_paper":"/paper/1809.09600","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:7b2c314bde4b4de7d29d345c48d95ed4c8dfeffa5be4572f7c91ddfab4e099b7","observation_id":"ba758956-0719-4da0-b670-da3f0136e338","resolution":{"observed_at":"2026-08-04T17:57:52.529020Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.533363Z","title":"React: Synergizing reasoning and acting in language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.533363Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:5942f63c07f62746ad2db5ace81e7e11a9fc57a9d2270036bece5c84e1c42f1d","observation_id":"eea63485-c17e-4b42-8eb9-db0c30be623c","resolution":{"observed_at":"2026-08-04T17:57:52.533363Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.537759Z","title":"Judging llm-as-a-judge with mt-bench and chatbot arena","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.537759Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:6965c5ad387ff4386afe86cc69b2de722c0113121b6ed1f3d8933c274c4622a5","observation_id":"437d8352-5c08-410a-9f1f-19c85dcfe686","resolution":{"observed_at":"2026-08-04T17:57:52.537759Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13372","last_updated":"2024-06-27T22:44:48Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-03-20T08:08:54Z","title":"LlamaFactory: Unified Efficient Fine-Tuning of 100+ Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.13372","snapshot_observed_at":"2026-08-04T17:57:52.541921Z","title":"Llamafactory: Unified efficient fine-tuning of 100+ language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.541921Z"},"links":{"cited_paper":"/paper/2403.13372","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:09f6900e2a0bac660277c6d79fdd7b5f313da13fb64252e5592c1d6c8ca3ffd9","observation_id":"7e02193e-4937-4a12-a2d2-473644009517","resolution":{"observed_at":"2026-08-04T17:57:52.541921Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.546193Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.546193Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:e3ffbf7b05d6883b5491993303d6aaddb7aa895476f4240ee4b464b7b096186e","observation_id":"ff85f450-61cb-48e0-bb67-3a7ed96e188a","resolution":{"observed_at":"2026-08-04T17:57:52.546193Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.551106Z","title":"@esa (Ref","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.551106Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:9ddca20917797fe33890d3d34dcbaa5f5e9713bc567cbc87d8c042693c05e295","observation_id":"8012ab0c-1af4-4a2d-b580-2c4654dae211","resolution":{"observed_at":"2026-08-04T17:57:52.551106Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.555572Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.555572Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:b67e07210932bf1ca968e28b1ee7b43d61249db9cd5b29722ab5c23ea7c1c2d5","observation_id":"52bbca65-e4ad-47ef-998d-97fa1ee52c1d","resolution":{"observed_at":"2026-08-04T17:57:52.555572Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T17:57:52.560008Z","title":"_7 tlPzKTþr ;]샦3Ӯ쀥ԤN #,¤HE ! , ۹ ]XO< iVzK L m;YcT !Z :4Uy","venue":null,"work_id":null,"year":1970},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.560008Z"},"links":{"citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:fbd131fc67a9b5aa8ff3142696b2b7dcb05436523d363311a09e0180cc0b37b4","observation_id":"b8e01014-a455-4ba4-b90a-e298a7e38715","resolution":{"observed_at":"2026-08-04T17:57:52.560008Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","latest_version":3,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs"},"reference_resolution":{"displayed":56,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":55,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":56},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 56 of 56 outbound references and 10 inbound Pith citation observations for arXiv:2509.14257."}