{"as_of":"2026-08-06T14:43:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:3d165867a4fac3dee1ed1da66d4f559d68837c0e30fdf50fe2faf8ec9be11b6e","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":8,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":8,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-06T06:34:29.942622+00:00","state":"measured"},{"denominator":8,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":8,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-03T23:33:52.937067Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":4,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2412.14135","last_updated":"2024-12-18T18:24:47Z","snapshot_observed_at":"2026-08-03T23:00:22.560900Z","submitted_at":"2024-12-18T18:24:47Z","title":"Scaling of Search and Learning: A Roadmap to Reproduce o1 from Reinforcement Learning Perspective","version":1},"cited_work":{"arxiv_id":"2412.14135","doi":"10.48550/arxiv.2412.14135","metadata_source":"arxiv_reference","pith_arxiv_id":"2412.14135","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling of search and learning: A roadmap to reproduce o1 from reinforcement learning perspective","venue":"arXiv (Cornell University)","work_id":"6a0e5e21-c5d9-4094-9ee7-58c92f267e2b","year":2024},"citing_paper":{"arxiv_id":"2412.18925","last_updated":"2024-12-25T15:12:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-25T15:12:34Z","title":"HuatuoGPT-o1, Towards Medical Complex Reasoning with LLMs","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-15T12:36:50.060335Z"},"links":{"cited_paper":"/paper/2412.14135","citing_paper":"/paper/2412.18925"},"observation_digest":"sha256:32173c8df04d0c6ff63d6c715da848252a1e00bc5da17652dba9a30022b8b494","observation_id":"647261e7-faf9-46f4-8233-cd650ecf5e9e","resolution":{"observed_at":"2026-05-15T12:36:50.283750Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.14135","last_updated":"2024-12-18T18:24:47Z","snapshot_observed_at":"2026-08-03T23:00:22.560900Z","submitted_at":"2024-12-18T18:24:47Z","title":"Scaling of Search and Learning: A Roadmap to Reproduce o1 from Reinforcement Learning Perspective","version":1},"cited_work":{"arxiv_id":"2412.14135","doi":"10.48550/arxiv.2412.14135","metadata_source":"arxiv_reference","pith_arxiv_id":"2412.14135","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling of search and learning: A roadmap to reproduce o1 from reinforcement learning perspective","venue":"arXiv (Cornell University)","work_id":"6a0e5e21-c5d9-4094-9ee7-58c92f267e2b","year":2024},"citing_paper":{"arxiv_id":"2501.05366","last_updated":"2025-01-09T16:48:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-09T16:48:17Z","title":"Search-o1: Agentic Search-Enhanced Large Reasoning Models","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-05-13T17:36:27.515468Z"},"links":{"cited_paper":"/paper/2412.14135","citing_paper":"/paper/2501.05366"},"observation_digest":"sha256:bdd713721cf12a0036201616e07ef2afe609535ba8313ea96de089a3cf764ad4","observation_id":"ff67168c-269a-401b-ae22-67cffee03c28","resolution":{"observed_at":"2026-05-13T17:36:27.640538Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.14135","last_updated":"2024-12-18T18:24:47Z","snapshot_observed_at":"2026-08-03T23:00:22.560900Z","submitted_at":"2024-12-18T18:24:47Z","title":"Scaling of Search and Learning: A Roadmap to Reproduce o1 from Reinforcement Learning Perspective","version":1},"cited_work":{"arxiv_id":"2412.14135","doi":"10.48550/arxiv.2412.14135","metadata_source":"arxiv_reference","pith_arxiv_id":"2412.14135","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling of search and learning: A roadmap to reproduce o1 from reinforcement learning perspective","venue":"arXiv (Cornell University)","work_id":"6a0e5e21-c5d9-4094-9ee7-58c92f267e2b","year":2024},"citing_paper":{"arxiv_id":"2502.17419","last_updated":"2025-06-25T02:24:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-24T18:50:52Z","title":"From System 1 to System 2: A Survey of Reasoning Large Language Models","version":6},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-13T01:36:23.845366Z"},"links":{"cited_paper":"/paper/2412.14135","citing_paper":"/paper/2502.17419"},"observation_digest":"sha256:aec8b05a03b3ddf5517d5e50e691b9edff12921cb92b4e7df34a429bd2edc9b9","observation_id":"454dedb2-e3f9-4e64-9a26-2b809b0efec5","resolution":{"observed_at":"2026-05-13T01:36:24.628273Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.14135","last_updated":"2024-12-18T18:24:47Z","snapshot_observed_at":"2026-08-03T23:00:22.560900Z","submitted_at":"2024-12-18T18:24:47Z","title":"Scaling of Search and Learning: A Roadmap to Reproduce o1 from Reinforcement Learning Perspective","version":1},"cited_work":{"arxiv_id":"2412.14135","doi":"10.48550/arxiv.2412.14135","metadata_source":"arxiv_reference","pith_arxiv_id":"2412.14135","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling of search and learning: A roadmap to reproduce o1 from reinforcement learning perspective","venue":"arXiv (Cornell University)","work_id":"6a0e5e21-c5d9-4094-9ee7-58c92f267e2b","year":2024},"citing_paper":{"arxiv_id":"2506.13351","last_updated":"2026-05-07T20:19:13Z","snapshot_observed_at":"2026-07-31T17:48:52.945760Z","submitted_at":"2025-06-16T10:43:38Z","title":"Direct Reasoning Optimization: Token-Level Reasoning Reflectivity Meets Rubric Gates for Unverifiable Tasks","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-19T09:48:56.990745Z"},"links":{"cited_paper":"/paper/2412.14135","citing_paper":"/paper/2506.13351"},"observation_digest":"sha256:759786233daf15fbce607368b727563c1dfbcbcb6bd8bb2d2cd5bfb1f24a9d6b","observation_id":"e9b1f88b-5d58-4037-921a-64db6211d813","resolution":{"observed_at":"2026-05-19T09:52:14.111450Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.14135","last_updated":"2024-12-18T18:24:47Z","snapshot_observed_at":"2026-08-03T23:00:22.560900Z","submitted_at":"2024-12-18T18:24:47Z","title":"Scaling of Search and Learning: A Roadmap to Reproduce o1 from Reinforcement Learning Perspective","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.14135","snapshot_observed_at":"2026-08-03T23:33:52.937067Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.05385","last_updated":"2026-07-23T06:36:18Z","snapshot_observed_at":"2026-08-04T04:34:49.292424Z","submitted_at":"2025-11-07T16:08:34Z","title":"TeaRAG: A Token-Efficient Agentic Retrieval-Augmented Generation Framework","version":2},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-03T23:33:52.937067Z"},"links":{"cited_paper":"/paper/2412.14135","citing_paper":"/paper/2511.05385"},"observation_digest":"sha256:c402247c8b49abfc4a5be14568a66396ea12ff640e8b7c986c1101299693c532","observation_id":"3372b021-128e-463a-9bf3-b4508077d489","resolution":{"observed_at":"2026-08-03T23:33:52.937067Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.14135","last_updated":"2024-12-18T18:24:47Z","snapshot_observed_at":"2026-08-03T23:00:22.560900Z","submitted_at":"2024-12-18T18:24:47Z","title":"Scaling of Search and Learning: A Roadmap to Reproduce o1 from Reinforcement Learning Perspective","version":1},"cited_work":{"arxiv_id":"2412.14135","doi":"10.48550/arxiv.2412.14135","metadata_source":"arxiv_reference","pith_arxiv_id":"2412.14135","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling of search and learning: A roadmap to reproduce o1 from reinforcement learning perspective","venue":"arXiv (Cornell University)","work_id":"6a0e5e21-c5d9-4094-9ee7-58c92f267e2b","year":2024},"citing_paper":{"arxiv_id":"2605.04065","last_updated":"2026-05-07T04:49:30Z","snapshot_observed_at":"2026-07-06T23:16:54.673178Z","submitted_at":"2026-04-11T07:26:04Z","title":"Free Energy-Driven Reinforcement Learning with Adaptive Advantage Shaping for Unsupervised Reasoning in LLMs","version":2},"reference_index":266,"source":"arxiv_source","source_observed_at":"2026-05-10T16:58:10.013475Z"},"links":{"cited_paper":"/paper/2412.14135","citing_paper":"/paper/2605.04065"},"observation_digest":"sha256:cd4fbfebfb6610820f45605943baea373d7a12ab312249a05675ab765823b8bd","observation_id":"a1b73dc4-bc02-4a79-bf8c-24778fe13642","resolution":{"observed_at":"2026-05-11T07:45:59.486005Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.14135","last_updated":"2024-12-18T18:24:47Z","snapshot_observed_at":"2026-08-03T23:00:22.560900Z","submitted_at":"2024-12-18T18:24:47Z","title":"Scaling of Search and Learning: A Roadmap to Reproduce o1 from Reinforcement Learning Perspective","version":1},"cited_work":{"arxiv_id":"2412.14135","doi":"10.48550/arxiv.2412.14135","metadata_source":"arxiv_reference","pith_arxiv_id":"2412.14135","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling of search and learning: A roadmap to reproduce o1 from reinforcement learning perspective","venue":"arXiv (Cornell University)","work_id":"6a0e5e21-c5d9-4094-9ee7-58c92f267e2b","year":2024},"citing_paper":{"arxiv_id":"2605.04066","last_updated":"2026-05-07T04:57:40Z","snapshot_observed_at":"2026-08-02T15:49:26.057284Z","submitted_at":"2026-04-11T07:34:59Z","title":"Adapt to Thrive! Adaptive Power-Mean Policy Optimization for Improved LLM Reasoning","version":2},"reference_index":251,"source":"arxiv_source","source_observed_at":"2026-05-10T16:51:19.555272Z"},"links":{"cited_paper":"/paper/2412.14135","citing_paper":"/paper/2605.04066"},"observation_digest":"sha256:4ec50405f138432e7df0026d73c07c4f5d7d3d640347479ce7ba0ac50d4c0c1c","observation_id":"9ffcf2c6-bd94-470f-b32f-550aa083dcb3","resolution":{"observed_at":"2026-05-11T08:01:00.754150Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.14135","last_updated":"2024-12-18T18:24:47Z","snapshot_observed_at":"2026-08-03T23:00:22.560900Z","submitted_at":"2024-12-18T18:24:47Z","title":"Scaling of Search and Learning: A Roadmap to Reproduce o1 from Reinforcement Learning Perspective","version":1},"cited_work":{"arxiv_id":"2412.14135","doi":"10.48550/arxiv.2412.14135","metadata_source":"arxiv_reference","pith_arxiv_id":"2412.14135","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling of search and learning: A roadmap to reproduce o1 from reinforcement learning perspective","venue":"arXiv (Cornell University)","work_id":"6a0e5e21-c5d9-4094-9ee7-58c92f267e2b","year":2024},"citing_paper":{"arxiv_id":"2607.00862","last_updated":"2026-07-01T12:27:14Z","snapshot_observed_at":"2026-08-02T04:50:53.247552Z","submitted_at":"2026-07-01T12:27:14Z","title":"CAT: Confidence-Adaptive Thinking for Efficient Reasoning of Large Reasoning Models","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-07-02T13:22:27.566432Z"},"links":{"cited_paper":"/paper/2412.14135","citing_paper":"/paper/2607.00862"},"observation_digest":"sha256:5907692768a38a94106b31620e3bfbc6252f75d539783becbcd6181cf03669c2","observation_id":"58f25752-65fa-4ac8-bbaa-43dd3fcbd1fa","resolution":{"observed_at":"2026-07-02T13:26:58.101193Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2412.14135/citation-record","integrity":"/paper/2412.14135/integrity","json":"/paper/2412.14135/citation-record.json","paper":"/paper/2412.14135"},"outbound":[],"paper":{"arxiv_id":"2412.14135","last_updated":"2024-12-18T18:24:47Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-03T23:00:22.560900Z","submitted_at":"2024-12-18T18:24:47Z","title":"Scaling of Search and Learning: A Roadmap to Reproduce o1 from Reinforcement Learning Perspective"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"thesis":"As of 6 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 8 inbound Pith citation observations for arXiv:2412.14135."}