{"as_of":"2026-08-09T00:52:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:d069019e045bbcea5330565318a80accf92de2982ca466e0a1ae7e6637234d5a","coverage":[{"denominator":87,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":87,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T14:25:53.635679Z","state":"measured"},{"denominator":92,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":92,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":5,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":5,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T10:28:41.579717Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-03T05:57:41.592060Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"cited_work":{"arxiv_id":"2502.06773","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.06773","snapshot_observed_at":"2026-07-03T05:57:41.592060Z","title":"On the emergence of thinking in llms i: Searching for the right intuition","venue":null,"work_id":"79ab0859-4c94-4e77-8733-792f09188125","year":2025},"citing_paper":{"arxiv_id":"2504.21318","last_updated":"2025-04-30T05:05:09Z","snapshot_observed_at":"2026-08-08T08:43:53.657438Z","submitted_at":"2025-04-30T05:05:09Z","title":"Phi-4-reasoning Technical Report","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-05-17T03:40:25.706499Z"},"links":{"cited_paper":"/paper/2502.06773","citing_paper":"/paper/2504.21318"},"observation_digest":"sha256:63ee6966e8deb1538d1c8b4bcf4a6cd43c2b57d8337febf81fb13bc46c16180f","observation_id":"9c08c4c8-1eff-47cd-8605-016f003dce04","resolution":{"observed_at":"2026-05-17T03:40:25.854460Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.06773","snapshot_observed_at":"2026-08-07T10:28:41.579717Z","title":"On the emergence of thinking in llms i: Searching for the right intuition","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.05213","last_updated":"2025-06-05T16:27:49Z","snapshot_observed_at":"2026-08-08T16:04:48.106610Z","submitted_at":"2025-06-05T16:27:49Z","title":"LLM-First Search: Self-Guided Exploration of the Solution Space","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T10:28:41.579717Z"},"links":{"cited_paper":"/paper/2502.06773","citing_paper":"/paper/2506.05213"},"observation_digest":"sha256:e291f978704b7869bd265760f2bc899aa1abe45688dadf57807c0271a5ead2d0","observation_id":"a41acd6a-a1c2-4b54-8f11-fec289847cce","resolution":{"observed_at":"2026-08-07T10:28:41.579717Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.06773","snapshot_observed_at":"2026-08-06T16:50:00.873671Z","title":"D., Zhang, X., Gopi, S., Peng, B., Li, B., Kulkarni, J., and Inan, H","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.12638","last_updated":"2025-07-16T21:21:03Z","snapshot_observed_at":"2026-08-08T14:46:42.003088Z","submitted_at":"2025-07-16T21:21:03Z","title":"Reasoning-Finetuning Repurposes Latent Representations in Base Models","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-06T16:50:00.873671Z"},"links":{"cited_paper":"/paper/2502.06773","citing_paper":"/paper/2507.12638"},"observation_digest":"sha256:3c524ad8c9ab4b65d45ea2de617979e39610495eff5e2ea07a82a2a252ee32cb","observation_id":"0ed703d5-4526-4335-9935-6d8a042fb6da","resolution":{"observed_at":"2026-08-06T16:50:00.873671Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"cited_work":{"arxiv_id":"2502.06773","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.06773","snapshot_observed_at":"2026-07-03T05:57:41.592060Z","title":"On the emergence of thinking in llms i: Searching for the right intuition","venue":null,"work_id":"79ab0859-4c94-4e77-8733-792f09188125","year":2025},"citing_paper":{"arxiv_id":"2606.01532","last_updated":"2026-06-02T03:44:54Z","snapshot_observed_at":"2026-08-05T19:46:58.020708Z","submitted_at":"2026-06-01T01:28:42Z","title":"Rethinking the Role of Positional Encoding: Sliding-Window Transformers without PE Remain Turing Complete","version":2},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-06-28T15:48:48.046003Z"},"links":{"cited_paper":"/paper/2502.06773","citing_paper":"/paper/2606.01532"},"observation_digest":"sha256:9e272a3e22c882354d745ff36e775bfe3d99edef0a681c996cefa3059249f884","observation_id":"83d99c6c-72e6-4eff-8a06-18fbb0536bef","resolution":{"observed_at":"2026-07-01T22:06:16.178705Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"cited_work":{"arxiv_id":"2502.06773","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.06773","snapshot_observed_at":"2026-07-03T05:57:41.592060Z","title":"On the emergence of thinking in llms i: Searching for the right intuition","venue":null,"work_id":"79ab0859-4c94-4e77-8733-792f09188125","year":2025},"citing_paper":{"arxiv_id":"2606.11470","last_updated":"2026-06-09T21:59:37Z","snapshot_observed_at":"2026-07-06T23:50:35.052764Z","submitted_at":"2026-06-09T21:59:37Z","title":"The Periodic Table of LLM Reasoning: A Structured Survey of Reasoning Paradigms, Methods, and Failure Modes","version":1},"reference_index":290,"source":"arxiv_source","source_observed_at":"2026-06-27T12:59:51.091008Z"},"links":{"cited_paper":"/paper/2502.06773","citing_paper":"/paper/2606.11470"},"observation_digest":"sha256:e7584e931aea3d1ed2991b12b459d861326f425b4144df8a49aa601349d25136","observation_id":"7b5fd31f-e1af-469f-91b7-bb688a19f061","resolution":{"observed_at":"2026-07-03T05:57:41.593548Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2502.06773/citation-record","integrity":"/paper/2502.06773/integrity","json":"/paper/2502.06773/citation-record.json","paper":"/paper/2502.06773"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2412.08905","last_updated":"2024-12-12T03:37:41Z","snapshot_observed_at":"2026-08-05T04:04:21.846023Z","submitted_at":"2024-12-12T03:37:41Z","title":"Phi-4 Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.08905","snapshot_observed_at":"2026-08-08T14:25:53.261354Z","title":"Phi-4 technical report","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.261354Z"},"links":{"cited_paper":"/paper/2412.08905","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:a24966165fb095805aa14c586a20579c9b427ba8dfad17207df984b508a80e69","observation_id":"ecc36f18-1e9b-40fe-93ff-1e108ea1af3a","resolution":{"observed_at":"2026-08-08T14:25:53.261354Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:53.267178Z","title":"Amc 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.267178Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:054772e1908d43b772d83445cb9a0b518ecafa4541e8736bc1835f40447468ab","observation_id":"b582b0e9-a95e-4d3e-9f94-0fb09582e460","resolution":{"observed_at":"2026-08-08T14:25:53.267178Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:53.271469Z","title":"Aime 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.271469Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:d199e061aca0b3ba59d50b039967a9b033f4b1e22ed1f618eae19b487f1cbcde","observation_id":"eece90db-4388-4ddf-9c31-7642bde6225a","resolution":{"observed_at":"2026-08-08T14:25:53.271469Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:53.275821Z","title":"Numinamath-cot","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.275821Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:bf42594653f39bb591dbad0dc7d4ba92ce42e09e2ba7f99f313efac6e0eb8a99","observation_id":"b5f48994-4ec5-4212-9c7f-2cefbe971e23","resolution":{"observed_at":"2026-08-08T14:25:53.275821Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21787","last_updated":"2024-12-30T19:03:24Z","snapshot_observed_at":"2026-07-06T18:55:11.576666Z","submitted_at":"2024-07-31T17:57:25Z","title":"Large Language Monkeys: Scaling Inference Compute with Repeated Sampling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21787","snapshot_observed_at":"2026-08-08T14:25:53.280245Z","title":"Large language monkeys: Scaling inference compute with repeated sampling","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.280245Z"},"links":{"cited_paper":"/paper/2407.21787","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:851e09acaa7fd8fbd362ba04919498698086aae0f3d679a18745c3188be84bd3","observation_id":"6a10645d-8bfd-4e75-b884-8f82f92de4b7","resolution":{"observed_at":"2026-08-08T14:25:53.280245Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2212.08073","last_updated":"2022-12-15T06:19:23Z","snapshot_observed_at":"2026-08-02T04:53:58.766070Z","submitted_at":"2022-12-15T06:19:23Z","title":"Constitutional AI: Harmlessness from AI Feedback","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2212.08073","snapshot_observed_at":"2026-08-08T14:25:53.284895Z","title":"Constitutional ai: Harmlessness from ai feedback","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.284895Z"},"links":{"cited_paper":"/paper/2212.08073","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:73bef441d66d34cee3df95f5178335bbd8326f87d3fd9f7af6ec297949321bcd","observation_id":"c281d1d5-4ccc-459e-985e-eb611e13f7f9","resolution":{"observed_at":"2026-08-08T14:25:53.284895Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:53.290041Z","title":"Scaling test-time compute with open models, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.290041Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:00d155528e6a751539c482bf9af5d359ed60874bb364c907f4f57127cda7cda6","observation_id":"1082ab03-ef0f-4ca1-97a8-9f0e7f0cbfcb","resolution":{"observed_at":"2026-08-08T14:25:53.290041Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:53.294520Z","title":"Open-r1: a fully open reproduction of deepseek-r1","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.294520Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:491a63ca7e45f45e7f346ed365268bd11c1c4712933e5d4b58f99b3f28ee74b2","observation_id":"0b217de4-d4a0-4bb4-ad06-04e7fdc930c6","resolution":{"observed_at":"2026-08-08T14:25:53.294520Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.07055","last_updated":"2024-08-13T17:46:12Z","snapshot_observed_at":"2026-07-06T19:00:17.992934Z","submitted_at":"2024-08-13T17:46:12Z","title":"LongWriter: Unleashing 10,000+ Word Generation from Long Context LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.07055","snapshot_observed_at":"2026-08-08T14:25:53.299016Z","title":"Longwriter: Unleashing 10,000+ word generation from long context llms","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.299016Z"},"links":{"cited_paper":"/paper/2408.07055","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:3f44181525632d289425c5f074f5e4d3183104aea824c3cd8dbed820c7727f2e","observation_id":"82bdc4bf-793f-4cd1-ad2b-fd0c52762730","resolution":{"observed_at":"2026-08-08T14:25:53.299016Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.01660","last_updated":"2025-04-11T23:24:37Z","snapshot_observed_at":"2026-07-06T18:24:38.958815Z","submitted_at":"2024-06-03T17:53:25Z","title":"Self-Improving Robust Preference Optimization","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.01660","snapshot_observed_at":"2026-08-08T14:25:53.303921Z","title":"Self-improving robust preference optimization","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.303921Z"},"links":{"cited_paper":"/paper/2406.01660","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:b4dc0513eccac553c14f5e4ab63e98921873dc2ab6b01b88424dfa0545074a76","observation_id":"3357ea57-02a9-44c2-937e-6a19461ad863","resolution":{"observed_at":"2026-08-08T14:25:53.303921Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.03553","last_updated":"2024-09-27T08:16:28Z","snapshot_observed_at":"2026-08-08T14:49:36.757543Z","submitted_at":"2024-05-06T15:20:30Z","title":"AlphaMath Almost Zero: Process Supervision without Process","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.03553","snapshot_observed_at":"2026-08-08T14:25:53.308231Z","title":"Alphamath almost zero: process supervision without process","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.308231Z"},"links":{"cited_paper":"/paper/2405.03553","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:d4158bcce979b2745ecf221deb894ddcf4761f987cb61d64e6b443105e855d41","observation_id":"2402dfe1-17ca-45d9-8aa6-e55610864943","resolution":{"observed_at":"2026-08-08T14:25:53.308231Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:53.312356Z","title":"Codeforces dataset","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.312356Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:b4ebb6e9128e44867b8753d5fc76ab006763f7a2fbeba0b3ac832c15b067f279","observation_id":"ce1aee1b-0de3-4e4b-9af5-2daeb8862248","resolution":{"observed_at":"2026-08-08T14:25:53.312356Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:53.316251Z","title":"Process reinforcement through implicit rewards","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.316251Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:e749af57db8215e8342af273154bd1aba36241fcd9ce6535a77c707b7b37e707","observation_id":"1e9f82f3-63a8-4862-bab1-3583f27d8217","resolution":{"observed_at":"2026-08-08T14:25:53.316251Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.886464Z","title":"Deepseek-r1: Incentivizing reasoning capability in llms via reinforcement learning","venue":null,"work_id":"b838722e-0a87-463a-afa3-255b1f37ad41","year":2025},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.320253Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:974554271604bccfd399895c89048b92940c4d305e003db18b6d3a9c3caeefa9","observation_id":"b2b9e052-b993-479d-b1e8-f235fbb6df7b","resolution":{"observed_at":"2026-08-08T14:25:54.891480Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.870771Z","title":"Flash A ttention-2: Faster attention with better parallelism and work partitioning","venue":null,"work_id":"a168a182-149c-4f98-b616-34935023043d","year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.324290Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:dea7b034c788db80ce9b99a2b4977e4c6371085cb4bb8d00ff325f4a1733dc2b","observation_id":"51617f6d-02d7-4ade-b1ae-f7ee7d57f972","resolution":{"observed_at":"2026-08-08T14:25:54.875355Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.856325Z","title":"Ai achieves silver-medal standard solving international mathematical olympiad problems, 2024","venue":null,"work_id":"1e6773cb-653d-497a-9a98-de221db9756d","year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.328130Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:bcee7edea972ecd523a01b9241c8c6ba25a0aaa807041462064d1f6fc585e824","observation_id":"c4041b13-e1a1-492a-9fae-53fd468c599b","resolution":{"observed_at":"2026-08-08T14:25:54.860648Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.17179","last_updated":"2024-02-09T00:13:46Z","snapshot_observed_at":"2026-07-06T16:25:25.534843Z","submitted_at":"2023-09-29T12:20:19Z","title":"Alphazero-like Tree-Search can Guide Large Language Model Decoding and Training","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.17179","snapshot_observed_at":"2026-08-08T14:25:53.332278Z","title":"Alphazero-like tree-search can guide large language model decoding and training","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.332278Z"},"links":{"cited_paper":"/paper/2309.17179","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:fec3692da575786cb3e486ea49fc68f44c7c37100e94fef7d0d71f22920a5bf1","observation_id":"bdd185cc-189c-4585-a351-a8aeff4a8531","resolution":{"observed_at":"2026-08-08T14:25:53.332278Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.842423Z","title":"Introducing gemini 2.0: our new ai model for the agentic era","venue":null,"work_id":"a6820b27-e001-45bf-8787-759e9906684b","year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.336397Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:0b45f1201913daa4c1605a4d0d91952a66b8e79bdd13991e08cd44b024fa735b","observation_id":"56005d59-b58d-4a7d-a27f-8341ea7a820f","resolution":{"observed_at":"2026-08-08T14:25:54.846667Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-04T20:57:22.329262Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-08T14:25:53.340082Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.340082Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:ea1d21025ef3e55b4bd776ccff051caf5dc52386ede7eb4fcdd09d74b2b682bf","observation_id":"c7e1636c-8b03-4263-91b1-ae5239e0bc6a","resolution":{"observed_at":"2026-08-08T14:25:53.340082Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2009.03300","last_updated":"2021-01-12T18:57:11Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-09-07T17:59:25Z","title":"Measuring Massive Multitask Language Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.03300","snapshot_observed_at":"2026-08-08T14:25:53.344239Z","title":"Measuring massive multitask language understanding","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.344239Z"},"links":{"cited_paper":"/paper/2009.03300","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:b1f5f953c6a23c3113f9e0ee77b1361a5f333084eab130e2b7f9cfffd0237063","observation_id":"6d640c67-b92c-4e0d-9e83-c349fecca3a1","resolution":{"observed_at":"2026-08-08T14:25:53.344239Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2103.03874","last_updated":"2021-11-08T21:30:18Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-03-05T18:59:39Z","title":"Measuring Mathematical Problem Solving With the MATH Dataset","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2103.03874","snapshot_observed_at":"2026-08-08T14:25:53.348541Z","title":"Measuring mathematical problem solving with the math dataset","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.348541Z"},"links":{"cited_paper":"/paper/2103.03874","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:c980fda3e3bfcd0ea0cd20dec1d2fd74f133e814a31f2e59fcb052e2f24a478c","observation_id":"388a53cd-7a7e-43c0-90f1-bcc3a536fd76","resolution":{"observed_at":"2026-08-08T14:25:53.348541Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.01798","last_updated":"2024-03-14T04:27:52Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-03T04:56:12Z","title":"Large Language Models Cannot Self-Correct Reasoning Yet","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.01798","snapshot_observed_at":"2026-08-08T14:25:53.352790Z","title":"Large language models cannot self-correct reasoning yet","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.352790Z"},"links":{"cited_paper":"/paper/2310.01798","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:5fad971ddc971a9c41249d84a1895111e95c7b2100d54c781aec4f257fdbbfc3","observation_id":"c2702d0c-f576-4c21-8bb6-35cca6a3d8e0","resolution":{"observed_at":"2026-08-08T14:25:53.352790Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.04642","last_updated":"2024-03-07T16:36:29Z","snapshot_observed_at":"2026-08-07T01:30:54.327974Z","submitted_at":"2024-03-07T16:36:29Z","title":"Teaching Large Language Models to Reason with Reinforcement Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.04642","snapshot_observed_at":"2026-08-08T14:25:53.357179Z","title":"Teaching large language models to reason with reinforcement learning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.357179Z"},"links":{"cited_paper":"/paper/2403.04642","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:1c28efb2457799941118efe1bc52bf5822b537a2972d6ea18401719116bc2c60","observation_id":"943e4545-6612-4952-a7c6-53c23682a60a","resolution":{"observed_at":"2026-08-08T14:25:53.357179Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.06458","last_updated":"2025-01-11T07:10:23Z","snapshot_observed_at":"2026-08-07T08:07:31.851127Z","submitted_at":"2025-01-11T07:10:23Z","title":"O1 Replication Journey -- Part 3: Inference-time Scaling for Medical Reasoning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.06458","snapshot_observed_at":"2026-08-08T14:25:53.361801Z","title":"O1 replication journey -- part 3: Inference-time scaling for medical reasoning","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.361801Z"},"links":{"cited_paper":"/paper/2501.06458","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:a768fd002678433ff14517a014cadcbde4217e28dd55783ae497a79d9208a37c","observation_id":"a28b3f64-5a69-4c93-831e-8e7869e8ff38","resolution":{"observed_at":"2026-08-08T14:25:53.361801Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14992","last_updated":"2023-10-23T07:24:28Z","snapshot_observed_at":"2026-07-06T15:32:25.931739Z","submitted_at":"2023-05-24T10:28:28Z","title":"Reasoning with Language Model is Planning with World Model","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14992","snapshot_observed_at":"2026-08-08T14:25:53.366490Z","title":"Reasoning with language model is planning with world model","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.366490Z"},"links":{"cited_paper":"/paper/2305.14992","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:3c488cb08e854040b107753b66f277b96ce78049ae5b69a54df313aa1bacce13","observation_id":"b1e892ab-f379-458d-9888-707cb8c2c758","resolution":{"observed_at":"2026-08-08T14:25:53.366490Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21276","last_updated":"2024-10-25T17:43:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-10-25T17:43:01Z","title":"GPT-4o System Card","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21276","snapshot_observed_at":"2026-08-08T14:25:53.370738Z","title":"Gpt-4o system card","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.370738Z"},"links":{"cited_paper":"/paper/2410.21276","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:baa464612efae261feb274cadc6e4d14a2f9d855e1768e36b31aad384eb9bf65","observation_id":"5c92a674-1136-4135-92ae-a8c9385107c5","resolution":{"observed_at":"2026-08-08T14:25:53.370738Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.11651","last_updated":"2025-06-13T16:15:45Z","snapshot_observed_at":"2026-08-08T22:57:04.484148Z","submitted_at":"2025-01-20T18:33:33Z","title":"T1: Advancing Language Model Reasoning through Reinforcement Learning and Inference Scaling","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.11651","snapshot_observed_at":"2026-08-08T14:25:53.374934Z","title":"Advancing language model reasoning through reinforcement learning and inference scaling","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.374934Z"},"links":{"cited_paper":"/paper/2501.11651","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:e6b864ac7bf6f3ff0f9a3dcdc67ba02a5f6c595820d1f241e543061ec598c5e8","observation_id":"89738edf-e17b-4fb8-aa32-0e92980a4dfa","resolution":{"observed_at":"2026-08-08T14:25:53.374934Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.11143","last_updated":"2025-10-09T12:22:46Z","snapshot_observed_at":"2026-07-31T12:28:37.704994Z","submitted_at":"2024-05-20T01:04:40Z","title":"OpenRLHF: An Easy-to-use, Scalable and High-performance RLHF Framework","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.11143","snapshot_observed_at":"2026-08-08T14:25:53.379420Z","title":"Openrlhf: An easy-to-use, scalable and high-performance rlhf framework","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.379420Z"},"links":{"cited_paper":"/paper/2405.11143","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:861e0b26db4381d61c446df7b079c68beb3ef885477051b2df2461aae9dc94d8","observation_id":"71c91751-5438-45e2-bc14-697b74db75eb","resolution":{"observed_at":"2026-08-08T14:25:53.379420Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.16489","last_updated":"2024-11-25T15:31:27Z","snapshot_observed_at":"2026-07-06T19:56:39.115630Z","submitted_at":"2024-11-25T15:31:27Z","title":"O1 Replication Journey -- Part 2: Surpassing O1-preview through Simple Distillation, Big Progress or Bitter Lesson?","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.16489","snapshot_observed_at":"2026-08-08T14:25:53.383766Z","title":"O1 replication journey--part 2: Surpassing o1-preview through simple distillation, big progress or bitter lesson? arXiv preprint arXiv:2411.16489 , 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.383766Z"},"links":{"cited_paper":"/paper/2411.16489","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:7ab9372079b793a74eaec6c3b29c63a13ea9df2f75f14cfef555f550a6a1734d","observation_id":"6b593b83-c156-43af-984f-0629a2a6616e","resolution":{"observed_at":"2026-08-08T14:25:53.383766Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.11694","last_updated":"2024-12-31T01:38:12Z","snapshot_observed_at":"2026-07-06T19:52:03.121993Z","submitted_at":"2024-11-18T16:15:17Z","title":"Enhancing LLM Reasoning with Reward-guided Tree Search","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.11694","snapshot_observed_at":"2026-08-08T14:25:53.387950Z","title":"Technical report: Enhancing llm reasoning with reward-guided tree search","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.387950Z"},"links":{"cited_paper":"/paper/2411.11694","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:e052c6c77acdc285d9de4359f07a73717e336ce646e204f43aec8b1a37011486","observation_id":"9a181384-00a6-44af-a6d7-0d6ebdc467e3","resolution":{"observed_at":"2026-08-08T14:25:53.387950Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.07974","last_updated":"2024-06-06T17:41:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-03-12T17:58:04Z","title":"LiveCodeBench: Holistic and Contamination Free Evaluation of Large Language Models for Code","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.07974","snapshot_observed_at":"2026-08-08T14:25:53.392275Z","title":"Livecodebench: Holistic and contamination free evaluation of large language models for code","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.392275Z"},"links":{"cited_paper":"/paper/2403.07974","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:23ae14c3b2e424d4f07dfff47557a05e9d6b58c74c486cbed527f2b1b960a725","observation_id":"03b4dae8-94e7-44d4-935e-32a759b793f1","resolution":{"observed_at":"2026-08-08T14:25:53.392275Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.16720","last_updated":"2026-04-30T02:46:40Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-21T18:04:31Z","title":"OpenAI o1 System Card","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.16720","snapshot_observed_at":"2026-08-08T14:25:53.396401Z","title":"Openai o1 system card","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.396401Z"},"links":{"cited_paper":"/paper/2412.16720","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:7ee76122eccc805d4792085e62b35c33b55734ccb88de083cc43c9420172b584","observation_id":"a9ce09a3-0a28-4bbe-baa8-a3c690682860","resolution":{"observed_at":"2026-08-08T14:25:53.396401Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1611.05397","last_updated":"2016-11-16T18:21:29Z","snapshot_observed_at":"2026-08-08T00:00:13.424932Z","submitted_at":"2016-11-16T18:21:29Z","title":"Reinforcement Learning with Unsupervised Auxiliary Tasks","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1611.05397","snapshot_observed_at":"2026-08-08T14:25:53.401082Z","title":"Reinforcement learning with unsupervised auxiliary tasks","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.401082Z"},"links":{"cited_paper":"/paper/1611.05397","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:2816c41b08f090db8b14457aaae9c736c59db8f205eb006f8e5cf4d4a7957d85","observation_id":"9fd80777-38eb-4813-b345-735b42bb596c","resolution":{"observed_at":"2026-08-08T14:25:53.401082Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.828005Z","title":"Kimi k1.5: Scaling reinforcement learning with llms","venue":null,"work_id":"457791f2-77c5-4970-81ab-417e24a3297a","year":2025},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.405463Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:50c1cd56950f7a90f961777a9bdcec059593728f6385caad7a90d48cd2eefb1f","observation_id":"9101808f-b24d-40c7-8fe0-e762b331cb81","resolution":{"observed_at":"2026-08-08T14:25:54.832565Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16265","last_updated":"2024-06-26T14:01:15Z","snapshot_observed_at":"2026-08-01T16:15:24.164308Z","submitted_at":"2024-05-25T15:07:33Z","title":"MindStar: Enhancing Math Reasoning in Pre-trained LLMs at Inference Time","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16265","snapshot_observed_at":"2026-08-08T14:25:53.409568Z","title":"Mindstar: Enhancing math reasoning in pre-trained llms at inference time","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.409568Z"},"links":{"cited_paper":"/paper/2405.16265","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:94b3d2a900be9ade91b113447b042efb4c37360da8b5df9ea2192185ce2610c8","observation_id":"79363580-8c3b-4201-904c-95637ea29156","resolution":{"observed_at":"2026-08-08T14:25:53.409568Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12917","last_updated":"2024-10-04T17:28:45Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-09-19T17:16:21Z","title":"Training Language Models to Self-Correct via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12917","snapshot_observed_at":"2026-08-08T14:25:53.413796Z","title":"Training language models to self-correct via reinforcement learning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.413796Z"},"links":{"cited_paper":"/paper/2409.12917","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:23fcb13da44ae4764b0a3de0cc30f13f5a2c61b1fca5df821f6f6a3d1ab6bde5","observation_id":"ab2f45c7-cd1e-4198-985a-51258f2f78d0","resolution":{"observed_at":"2026-08-08T14:25:53.413796Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.813764Z","title":"Numinamath: The largest public dataset in ai4maths with 860k pairs of competition math problems and solutions","venue":null,"work_id":"a218ac7c-01a5-479b-ba95-cb47c4afed90","year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.418178Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:ff2d6beb34aa6b256161d52ae3885995f98bab0e65b98a07c54f0876962c5c44","observation_id":"30c427f1-2c02-460d-8113-97c8523d761f","resolution":{"observed_at":"2026-08-08T14:25:54.818246Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.20050","last_updated":"2023-05-31T17:24:00Z","snapshot_observed_at":"2026-08-05T13:11:04.104454Z","submitted_at":"2023-05-31T17:24:00Z","title":"Let's Verify Step by Step","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.20050","snapshot_observed_at":"2026-08-08T14:25:53.422225Z","title":"Let's verify step by step","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.422225Z"},"links":{"cited_paper":"/paper/2305.20050","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:9265dbacbf0d261aead000fdf648f630628d57c71dcd78fff95106922ffb443c","observation_id":"dd0171c7-9298-42b6-be98-39bf2838d347","resolution":{"observed_at":"2026-08-08T14:25:53.422225Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12875","last_updated":"2024-09-21T06:48:45Z","snapshot_observed_at":"2026-08-07T02:31:13.044154Z","submitted_at":"2024-02-20T10:11:03Z","title":"Chain of Thought Empowers Transformers to Solve Inherently Serial Problems","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12875","snapshot_observed_at":"2026-08-08T14:25:53.426639Z","title":"Chain of thought empowers transformers to solve inherently serial problems","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.426639Z"},"links":{"cited_paper":"/paper/2402.12875","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:04bfa411012b3b4a9aa5bd72423dff4d352b63d5ca6d71674e30bd0561ea0dc1","observation_id":"a58e83aa-0365-45a6-87bb-477e0bcc5451","resolution":{"observed_at":"2026-08-08T14:25:53.426639Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.09583","last_updated":"2025-06-04T08:58:56Z","snapshot_observed_at":"2026-08-02T06:48:43.121988Z","submitted_at":"2023-08-18T14:23:21Z","title":"WizardMath: Empowering Mathematical Reasoning for Large Language Models via Reinforced Evol-Instruct","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.09583","snapshot_observed_at":"2026-08-08T14:25:53.430999Z","title":"Wizardmath: Empowering mathematical reasoning for large language models via reinforced evol-instruct","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.430999Z"},"links":{"cited_paper":"/paper/2308.09583","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:15072eff1a440f2923473edf63a1f9edc66262436450acf510f1c0476156a145","observation_id":"fe092817-e977-4101-87d0-9c9ab5f6a8fd","resolution":{"observed_at":"2026-08-08T14:25:53.430999Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:53.435670Z","title":"Pre-train, prompt, and predict: A systematic survey of prompting methods in natural language processing","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.435670Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:315e7bfe292937a808a7c38f6d7584c5a44f40873e9365e406ac2e430a883de5","observation_id":"85937927-b0c4-406b-bc68-afea3a066a6a","resolution":{"observed_at":"2026-08-08T14:25:53.435670Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.09413","last_updated":"2024-12-22T10:44:13Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-12T16:20:36Z","title":"Imitate, Explore, and Self-Improve: A Reproduction Report on Slow-thinking Reasoning Systems","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.09413","snapshot_observed_at":"2026-08-08T14:25:53.439518Z","title":"Imitate, explore, and self-improve: A reproduction report on slow-thinking reasoning systems","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.439518Z"},"links":{"cited_paper":"/paper/2412.09413","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:055ffb05d76d3be971db5a4d7bed15720581372b3dec9e90b35c6e73ef358804","observation_id":"6ad11a03-cdce-412c-842b-a4a40d04aaa1","resolution":{"observed_at":"2026-08-08T14:25:53.439518Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.789140Z","title":"Llama-3.1-8b","venue":null,"work_id":"84748077-6761-440a-8532-e333d0c052ca","year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.443750Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:0e063c59303f9f11c83661c04c7525b85193b990111626bd72bdad86be1d177c","observation_id":"15b8f081-e7f6-412a-8cb6-8426317c24a6","resolution":{"observed_at":"2026-08-08T14:25:54.793295Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.774754Z","title":"Ray: A distributed framework for emerging \\ AI \\ applications","venue":null,"work_id":"f6dd26af-f783-4e8a-b637-3bbae0c27e16","year":2018},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.447529Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:74b52be4a7118d3d6227816ecc5217fdfdb2ed1545f9c11bd1eda68d3aafa9a8","observation_id":"b208dd9e-f17c-4bc1-a706-1b09b8c00692","resolution":{"observed_at":"2026-08-08T14:25:54.779383Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.07923","last_updated":"2024-04-11T18:03:53Z","snapshot_observed_at":"2026-08-02T23:51:22.170619Z","submitted_at":"2023-10-11T22:35:18Z","title":"The Expressive Power of Transformers with Chain of Thought","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.07923","snapshot_observed_at":"2026-08-08T14:25:53.451451Z","title":"The expresssive power of transformers with chain of thought","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.451451Z"},"links":{"cited_paper":"/paper/2310.07923","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:88e05e93231307c634076e34e0bb53a156f347eb8160682a8d044460d226c63f","observation_id":"ac323491-468e-4ade-8b7c-7ae2a9fea4f6","resolution":{"observed_at":"2026-08-08T14:25:53.451451Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.19393","last_updated":"2025-03-01T06:07:39Z","snapshot_observed_at":"2026-07-06T20:29:11.710285Z","submitted_at":"2025-01-31T18:48:08Z","title":"s1: Simple test-time scaling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.19393","snapshot_observed_at":"2026-08-08T14:25:53.455635Z","title":"s1: Simple test-time scaling","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.455635Z"},"links":{"cited_paper":"/paper/2501.19393","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:1920510406a8cd4fc9abb935c1bf195a8bc2e8c939dd7ef25814b83421f9e8f2","observation_id":"b2b5d512-62b9-456e-893a-26c3ce31dee5","resolution":{"observed_at":"2026-08-08T14:25:53.455635Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.759299Z","title":"Sky-t1: Train your own o1 preview model within \\ 450","venue":null,"work_id":"0de5ed8f-9139-44a1-9cda-3036b0448ada","year":2025},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.459617Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:2b43313a7a17c87b52da16cd81449996c71f6a234c52785280ea5b4366f73772","observation_id":"e3e9fde9-2d69-418b-9f3b-12c5c3db937a","resolution":{"observed_at":"2026-08-08T14:25:54.764289Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.743963Z","title":null,"venue":null,"work_id":"30856b71-d392-4346-85d6-92b676f154df","year":2022},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.463706Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:e46f80749a96b56d19e8f0bc0ea8d49f123803883d55028deee8e1310071a719","observation_id":"b7089b6b-670b-4e80-a58d-dc265ff58890","resolution":{"observed_at":"2026-08-08T14:25:54.748210Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.729497Z","title":"Learning to reason with llms","venue":null,"work_id":"ca097393-80ab-4e79-9dad-c1e828c8de39","year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.468071Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:6c02c57c61ec7722fa473d382beb080853835ee1c8c1a8943475c9eccb3aaa6b","observation_id":"72bc56d0-1ef5-4f09-ac0f-67b5db078475","resolution":{"observed_at":"2026-08-08T14:25:54.733808Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.715287Z","title":"Math-500","venue":null,"work_id":"1d8e467c-68be-4718-9298-6a2f9c814464","year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.472120Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:384b3cf159d918be2d6a1245915fa1e41c334316a2be33fc0e9fcb1e09c13c54","observation_id":"a0a00aeb-efcb-44cc-bc5b-b1de23dd640a","resolution":{"observed_at":"2026-08-08T14:25:54.719735Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.701130Z","title":"Openai humaneval","venue":null,"work_id":"d5d4ab95-b5dd-4a30-9990-43640bf415d9","year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.478200Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:46a1c2c3e857183cebab985c6e4c18895b56c9911767a76360cda7f838535abf","observation_id":"dce3a3b1-830d-44f5-a017-d2b828112f30","resolution":{"observed_at":"2026-08-08T14:25:54.705136Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.686763Z","title":"Openai o1-mini advancing cost-efficient reasoning","venue":null,"work_id":"e5dbd1d9-2ac6-4eb5-a41e-3b25103f413e","year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.482502Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:07f48a36d2ee80f366bbe449d49652c0383fd7a572b46856f8c513a93b6f5d74","observation_id":"ecc12a85-1bea-422a-a208-8604b0edfa99","resolution":{"observed_at":"2026-08-08T14:25:54.691110Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:53.486672Z","title":"Training language models to follow instructions with human feedback","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.486672Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:6a8b0b51e22acb5adf94eb7bfd89ebe0ab2daaaf54151cfdbb8f687653c4cbac","observation_id":"eb58fd79-d437-4c38-b627-6b5ec21ecaad","resolution":{"observed_at":"2026-08-08T14:25:53.486672Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.18982","last_updated":"2024-10-08T15:13:01Z","snapshot_observed_at":"2026-07-06T19:39:15.025945Z","submitted_at":"2024-10-08T15:13:01Z","title":"O1 Replication Journey: A Strategic Progress Report -- Part 1","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.18982","snapshot_observed_at":"2026-08-08T14:25:53.490904Z","title":"O1 replication journey: A strategic progress report--part 1","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.490904Z"},"links":{"cited_paper":"/paper/2410.18982","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:13ae0a4dcd8930b70853a44b46543951bb0bae3d5d7ddbf48c159d58a258c447","observation_id":"aa5ce6f6-2465-4fbc-9d8a-af09777caf75","resolution":{"observed_at":"2026-08-08T14:25:53.490904Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.06195","last_updated":"2024-08-12T14:42:13Z","snapshot_observed_at":"2026-07-31T21:52:59.633187Z","submitted_at":"2024-08-12T14:42:13Z","title":"Mutual Reasoning Makes Smaller LLMs Stronger Problem-Solvers","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.06195","snapshot_observed_at":"2026-08-08T14:25:53.496098Z","title":"Mutual reasoning makes smaller llms stronger problem-solvers","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.496098Z"},"links":{"cited_paper":"/paper/2408.06195","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:74d7dacd3f300c2496c5e879017c70a61c8a70a643cfcab2e2e895431d8c77dd","observation_id":"8b4943d0-f8c2-4144-99f0-d7676e972f86","resolution":{"observed_at":"2026-08-08T14:25:53.496098Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.661649Z","title":"Qwen-2.5-32b","venue":null,"work_id":"88b3af6e-6402-42f7-91f7-bd3bc8e34a1e","year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.500841Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:ac1c2abc6b711b3d2382795867f3bd598bae0ae677d8d95f46117c520db0eed0","observation_id":"c1a0fb65-ef3a-426b-8805-2aef8b147292","resolution":{"observed_at":"2026-08-08T14:25:54.665820Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.646653Z","title":"Qwq-32b-preview","venue":null,"work_id":"a5ae5bf7-b371-45e5-96ab-ea67b5e667b2","year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.506146Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:079c554620b01317309d47f770dd34a76ed72ea321fed6391c4f3538c7ad7abb","observation_id":"1283d6a6-60a1-46c0-ad8d-ccc8592df38b","resolution":{"observed_at":"2026-08-08T14:25:54.651382Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.630965Z","title":"Qwq-longcot-130k-cleaned","venue":null,"work_id":"2fe4bbe5-c276-488c-8bdf-4c234a8118b5","year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.510751Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:f184989951c211dab1bf4e623294857e94e38af7599637cf1f746684be426ff3","observation_id":"8ade35d2-357e-4fcc-906c-02673a1a5888","resolution":{"observed_at":"2026-08-08T14:25:54.635311Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.616484Z","title":"Qwq: Reflect deeply on the boundaries of the unknown","venue":null,"work_id":"da896a62-59c0-4eb9-a7dd-8357c894ec85","year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.515039Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:338cb1e23642ffc6f65b6ec5a7c2effb6a552a8296b17820d743c1c458b5648b","observation_id":"b67dcc04-5003-4918-a1a2-cff852db7ea0","resolution":{"observed_at":"2026-08-08T14:25:54.620823Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.18219","last_updated":"2024-07-26T17:50:27Z","snapshot_observed_at":"2026-08-07T12:10:53.303603Z","submitted_at":"2024-07-25T17:35:59Z","title":"Recursive Introspection: Teaching Language Model Agents How to Self-Improve","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.18219","snapshot_observed_at":"2026-08-08T14:25:53.519496Z","title":"Recursive introspection: Teaching language model agents how to self-improve","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.519496Z"},"links":{"cited_paper":"/paper/2407.18219","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:c9290cef4851eb6f4de93d64782e0d01ddbd2a4ed7487945cf3305244766c9bc","observation_id":"41f5cea6-2626-484a-90ac-1707ef95cbc1","resolution":{"observed_at":"2026-08-08T14:25:53.519496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.12022","last_updated":"2023-11-20T18:57:34Z","snapshot_observed_at":"2026-08-04T22:55:15.345443Z","submitted_at":"2023-11-20T18:57:34Z","title":"GPQA: A Graduate-Level Google-Proof Q&A Benchmark","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.12022","snapshot_observed_at":"2026-08-08T14:25:53.524212Z","title":"Gpqa: A graduate-level google-proof q&a benchmark","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.524212Z"},"links":{"cited_paper":"/paper/2311.12022","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:0a680d7681ae68f33f3cce95d814209c678074da6eab90ffcbc43593b6d46c2c","observation_id":"73ac72b7-cd88-409a-8e9e-983c7e3826ee","resolution":{"observed_at":"2026-08-08T14:25:53.524212Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.601559Z","title":"A program for the machine translation of natural languages","venue":null,"work_id":"116dc45c-cb0f-4ee7-8bf8-f8f160a2f5c6","year":1961},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.528679Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:f817a8361581f9cff755bb6fed6d6118e9a6dc56d01e9e777db5036088ece5fe","observation_id":"118b5f30-a0b3-47a3-8ce1-de26512b8c3e","resolution":{"observed_at":"2026-08-08T14:25:54.606496Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.03314","last_updated":"2024-08-06T17:35:05Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-08-06T17:35:05Z","title":"Scaling LLM Test-Time Compute Optimally can be More Effective than Scaling Model Parameters","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.03314","snapshot_observed_at":"2026-08-08T14:25:53.532654Z","title":"Scaling llm test-time compute optimally can be more effective than scaling model parameters","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.532654Z"},"links":{"cited_paper":"/paper/2408.03314","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:b151b68798fdaf759d45920604063936942ce1c164f2667895d7d39e588402f4","observation_id":"1a275c71-3f7d-4a64-b783-e466d9d6137f","resolution":{"observed_at":"2026-08-08T14:25:53.532654Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.08146","last_updated":"2024-10-10T17:31:23Z","snapshot_observed_at":"2026-08-02T06:08:09.777151Z","submitted_at":"2024-10-10T17:31:23Z","title":"Rewarding Progress: Scaling Automated Process Verifiers for LLM Reasoning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.08146","snapshot_observed_at":"2026-08-08T14:25:53.536838Z","title":"Rewarding progress: Scaling automated process verifiers for llm reasoning","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.536838Z"},"links":{"cited_paper":"/paper/2410.08146","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:e453902213a827364f8f719e9c59fdefd2629da8d93bfcebeac1631b3659596b","observation_id":"46209a82-f91d-419f-94a7-731aeaee85d3","resolution":{"observed_at":"2026-08-08T14:25:53.536838Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.09918","last_updated":"2025-07-11T03:52:42Z","snapshot_observed_at":"2026-08-05T13:59:17.022236Z","submitted_at":"2024-10-13T16:53:02Z","title":"Dualformer: Controllable Fast and Slow Thinking by Learning with Randomized Reasoning Traces","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.09918","snapshot_observed_at":"2026-08-08T14:25:53.541676Z","title":"Dualformer: Controllable fast and slow thinking by learning with randomized reasoning traces","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.541676Z"},"links":{"cited_paper":"/paper/2410.09918","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:1ae35f3e67f2990bd70d0d9aefa234c60e4dc95b61182e34c8e381d16b3c3973","observation_id":"b8806b55-4f4f-45fc-819c-638cf7f0127a","resolution":{"observed_at":"2026-08-08T14:25:53.541676Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-08T14:25:53.545928Z","title":"Proximal policy optimization algorithms","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.545928Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:cf7710b1fec2ad8a6f318cd16d343165ef9db24477f7e76cc9fafad1ba59714e","observation_id":"568c3934-2dd7-4746-81f8-88f9b644773d","resolution":{"observed_at":"2026-08-08T14:25:53.545928Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.12253","last_updated":"2024-12-10T18:19:29Z","snapshot_observed_at":"2026-08-05T22:35:55.296192Z","submitted_at":"2024-04-18T15:21:34Z","title":"Toward Self-Improvement of LLMs via Imagination, Searching, and Criticizing","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.12253","snapshot_observed_at":"2026-08-08T14:25:53.550735Z","title":"Toward self-improvement of llms via imagination, searching, and criticizing","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.550735Z"},"links":{"cited_paper":"/paper/2404.12253","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:acdaab94f5db2feacdd745d3e0ba0d193ac5eba6899e8c8495d38dad6e5fcda9","observation_id":"400e6bbb-f577-4f4a-91de-f675b930a8bf","resolution":{"observed_at":"2026-08-08T14:25:53.550735Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.587155Z","title":"Solving olympiad geometry without human demonstrations","venue":null,"work_id":"2193e429-336f-46b1-95ca-f80e5b6d916d","year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.554997Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:3bac650f7a6fb6b640c19f42466a56928c6b3e01b345e1acc2d9db6a37eee501","observation_id":"26cbd230-74a1-45e0-90b6-77d1584b1b29","resolution":{"observed_at":"2026-08-08T14:25:54.591672Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.14275","last_updated":"2022-11-25T18:19:44Z","snapshot_observed_at":"2026-08-01T02:16:43.109337Z","submitted_at":"2022-11-25T18:19:44Z","title":"Solving math word problems with process- and outcome-based feedback","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.14275","snapshot_observed_at":"2026-08-08T14:25:53.558980Z","title":"Solving math word problems with process-and outcome-based feedback","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.558980Z"},"links":{"cited_paper":"/paper/2211.14275","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:bc738b258863ae8d46367cdd6befd93c4b93c7740d765d30e20b167419fb499b","observation_id":"c967d194-9948-4029-b6f7-2b1915b48a62","resolution":{"observed_at":"2026-08-08T14:25:53.558980Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.10442","last_updated":"2025-04-07T09:09:39Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-11-15T18:59:27Z","title":"Enhancing the Reasoning Ability of Multimodal Large Language Models via Mixed Preference Optimization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.10442","snapshot_observed_at":"2026-08-08T14:25:53.563368Z","title":"Enhancing the reasoning ability of multimodal large language models via mixed preference optimization","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.563368Z"},"links":{"cited_paper":"/paper/2411.10442","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:f36ddbec35667d57f8cb782c5b36ad013a80ac2e81cd1aa3f60dc184adfb0f4b","observation_id":"078bb60e-14d7-4bdd-8e97-d4c805a86e85","resolution":{"observed_at":"2026-08-08T14:25:53.563368Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.14283","last_updated":"2024-07-22T10:01:49Z","snapshot_observed_at":"2026-07-06T18:34:15.291938Z","submitted_at":"2024-06-20T13:08:09Z","title":"Q*: Improving Multi-step Reasoning for LLMs with Deliberative Planning","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.14283","snapshot_observed_at":"2026-08-08T14:25:53.567513Z","title":"Q*: Improving multi-step reasoning for llms with deliberative planning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.567513Z"},"links":{"cited_paper":"/paper/2406.14283","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:a42a8beaf8b8b69cd02e8f33f9c3761b699849ac26c2319c3a752ca5b8f11143","observation_id":"02e6b3c6-4fce-4229-bd2a-91e10e91dc0b","resolution":{"observed_at":"2026-08-08T14:25:53.567513Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.09671","last_updated":"2024-10-12T23:42:16Z","snapshot_observed_at":"2026-08-08T14:45:36.681932Z","submitted_at":"2024-10-12T23:42:16Z","title":"OpenR: An Open Source Framework for Advanced Reasoning with Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.09671","snapshot_observed_at":"2026-08-08T14:25:53.571619Z","title":"Openr: An open source framework for advanced reasoning with large language models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.571619Z"},"links":{"cited_paper":"/paper/2410.09671","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:a2fd5538d190c78d3cb360d1f839a8a815eca1686c77aa546d127a832fb9376c","observation_id":"8889cdf4-a380-4569-9c9b-92db64076048","resolution":{"observed_at":"2026-08-08T14:25:53.571619Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.571593Z","title":"An empirical analysis of compute-optimal inference for problem-solving with language models","venue":null,"work_id":"fa2af39f-d903-4ce3-804d-bf0a13ff4c6a","year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.576225Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:00ed0e83a8cccc4c0bb0da613c16d16bb1f364aedbc1c03464b0d706b0ed1224","observation_id":"8c7e0202-04a5-4727-a56e-85f8bf1c08e8","resolution":{"observed_at":"2026-08-08T14:25:54.576399Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.06508","last_updated":"2024-10-09T03:20:02Z","snapshot_observed_at":"2026-08-05T03:22:34.062603Z","submitted_at":"2024-10-09T03:20:02Z","title":"Towards Self-Improvement of LLMs via MCTS: Leveraging Stepwise Knowledge with Curriculum Preference Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.06508","snapshot_observed_at":"2026-08-08T14:25:53.580233Z","title":"Towards self-improvement of llms via mcts: Leveraging stepwise knowledge with curriculum preference learning","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.580233Z"},"links":{"cited_paper":"/paper/2410.06508","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:3d97f29bbe6983a818c5db0dbd40899407e98deab68b9eeeb8cdfd450e247baf","observation_id":"490f703c-c4cf-40e5-a257-a31153a9bfa5","resolution":{"observed_at":"2026-08-08T14:25:53.580233Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2203.11171","last_updated":"2023-03-07T17:57:37Z","snapshot_observed_at":"2026-07-06T12:50:22.773056Z","submitted_at":"2022-03-21T17:48:52Z","title":"Self-Consistency Improves Chain of Thought Reasoning in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2203.11171","snapshot_observed_at":"2026-08-08T14:25:53.584409Z","title":"Self-consistency improves chain of thought reasoning in language models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.584409Z"},"links":{"cited_paper":"/paper/2203.11171","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:889528bed9a7e22b82f45ab6af56c2f1f79622e8d854b248f66e95bcad7ec3cf","observation_id":"dc1766c5-126c-43cc-9c11-535b88589061","resolution":{"observed_at":"2026-08-08T14:25:53.584409Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.10440","last_updated":"2025-07-21T03:53:30Z","snapshot_observed_at":"2026-08-07T09:36:29.319006Z","submitted_at":"2024-11-15T18:58:31Z","title":"LLaVA-CoT: Let Vision Language Models Reason Step-by-Step","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.10440","snapshot_observed_at":"2026-08-08T14:25:53.588274Z","title":"Llava-o1: Let vision language models reason step-by-step","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.588274Z"},"links":{"cited_paper":"/paper/2411.10440","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:3dd3ccc5d47f163b647e30a6c913c572480e586d55232ad479b1f1d42f323ff7","observation_id":"1fae56f8-5d36-4880-b537-d2c5179ab3c5","resolution":{"observed_at":"2026-08-08T14:25:53.588274Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-07-06T20:18:21.419068Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-08T14:25:53.592132Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-though","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.592132Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:7709d7a53c5cea0deb9d35dc5f19a7d6f5d234c71f3bd17a973aae0a7d6a0de2","observation_id":"04a629f7-ff2c-4e73-a168-ca0f29f98e55","resolution":{"observed_at":"2026-08-08T14:25:53.592132Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.12284","last_updated":"2024-05-03T17:36:07Z","snapshot_observed_at":"2026-08-02T15:00:50.388422Z","submitted_at":"2023-09-21T17:45:42Z","title":"MetaMath: Bootstrap Your Own Mathematical Questions for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.12284","snapshot_observed_at":"2026-08-08T14:25:53.596276Z","title":"Metamath: Bootstrap your own mathematical questions for large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.596276Z"},"links":{"cited_paper":"/paper/2309.12284","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:3db2d0385a8645297b5f7706fd9960173fe0e7b7f1ceb4ea5bf2736fe209b46c","observation_id":"62e15922-f60b-4ea6-811f-e6b923915a3d","resolution":{"observed_at":"2026-08-08T14:25:53.596276Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.03373","last_updated":"2025-02-05T17:13:32Z","snapshot_observed_at":"2026-07-06T20:31:41.231839Z","submitted_at":"2025-02-05T17:13:32Z","title":"Demystifying Long Chain-of-Thought Reasoning in LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.03373","snapshot_observed_at":"2026-08-08T14:25:53.600448Z","title":"Demystifying long chain-of-thought reasoning in llms","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.600448Z"},"links":{"cited_paper":"/paper/2502.03373","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:c0def8c15ae33edd910487565fd0ca1772d83f40fd19f5ce075c0d9a9eeeab55","observation_id":"9754a7d2-d866-4e3e-9803-e7d1ff172adf","resolution":{"observed_at":"2026-08-08T14:25:53.600448Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.01825","last_updated":"2023-09-13T03:57:29Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-08-03T15:34:01Z","title":"Scaling Relationship on Learning Mathematical Reasoning with Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.01825","snapshot_observed_at":"2026-08-08T14:25:53.604965Z","title":"Scaling relationship on learning mathematical reasoning with large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":80,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.604965Z"},"links":{"cited_paper":"/paper/2308.01825","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:0f141fac92bec177c50bcf8ddf85363042a03c765f1968d1abb35880684ec79b","observation_id":"6a28b955-a2f4-4089-9f84-33f32dc4962a","resolution":{"observed_at":"2026-08-08T14:25:53.604965Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:53.609543Z","title":"Tree of thoughts: Deliberate problem solving with large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":81,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.609543Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:23b307b08c67b503a8dd38972ca0a7f34a456082b21035948c67f3ffe407cf5e","observation_id":"47af6924-d35a-47f8-b9c0-07fe3f934934","resolution":{"observed_at":"2026-08-08T14:25:53.609543Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12122","last_updated":"2024-09-18T16:45:37Z","snapshot_observed_at":"2026-07-06T19:17:41.512834Z","submitted_at":"2024-09-18T16:45:37Z","title":"Qwen2.5-Math Technical Report: Toward Mathematical Expert Model via Self-Improvement","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12122","snapshot_observed_at":"2026-08-08T14:25:53.613796Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":82,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.613796Z"},"links":{"cited_paper":"/paper/2409.12122","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:6735b0bf44135b1281473e03ed04cbbed80cb948c91ac66d90bb3a8ab69e6c9b","observation_id":"1049b08b-fa90-4bb7-b99c-ebca1098589a","resolution":{"observed_at":"2026-08-08T14:25:53.613796Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T14:25:54.542937Z","title":"7b model and 8k examples: Emerging reasoning with reinforcement learning is both effective and efficient","venue":null,"work_id":"5b770fde-de42-43bb-971c-64d7b19d4f97","year":2025},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":83,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.618399Z"},"links":{"citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:ee0fd0ef1e96ef9924a5fbefb30aba6b4216ddc03ce8bd6240411761428abc1b","observation_id":"a6936156-37c3-4b86-a63d-0d780164f52d","resolution":{"observed_at":"2026-08-08T14:25:54.549628Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.17140","last_updated":"2024-06-06T03:59:24Z","snapshot_observed_at":"2026-07-06T18:05:55.328415Z","submitted_at":"2024-04-26T03:41:28Z","title":"Small Language Models Need Strong Verifiers to Self-Correct Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.17140","snapshot_observed_at":"2026-08-08T14:25:53.622557Z","title":"Small language models need strong verifiers to self-correct reasoning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":84,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.622557Z"},"links":{"cited_paper":"/paper/2404.17140","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:7c4cf87d770c5c35d2ca32d53bc51b97aae62ed00a346921ab058b526a2a4b07","observation_id":"0f9bad28-0108-4b8f-bf7c-123fe2271a93","resolution":{"observed_at":"2026-08-08T14:25:53.622557Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07394","last_updated":"2024-06-13T07:19:06Z","snapshot_observed_at":"2026-07-06T18:29:00.001236Z","submitted_at":"2024-06-11T16:01:07Z","title":"Accessing GPT-4 level Mathematical Olympiad Solutions via Monte Carlo Tree Self-refine with LLaMa-3 8B","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07394","snapshot_observed_at":"2026-08-08T14:25:53.627005Z","title":"Accessing gpt-4 level mathematical olympiad solutions via monte carlo tree self-refine with llama-3 8b","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":85,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.627005Z"},"links":{"cited_paper":"/paper/2406.07394","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:c945874be7202a998d4de19268415858923d5038d8c94e41f81ec565733f3907","observation_id":"388520e0-f5bf-489f-a245-0ef8e55cbf75","resolution":{"observed_at":"2026-08-08T14:25:53.627005Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.02884","last_updated":"2024-11-21T07:07:59Z","snapshot_observed_at":"2026-08-07T12:03:21.367108Z","submitted_at":"2024-10-03T18:12:29Z","title":"LLaMA-Berry: Pairwise Optimization for O1-like Olympiad-Level Mathematical Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.02884","snapshot_observed_at":"2026-08-08T14:25:53.631264Z","title":"Llama-berry: Pairwise optimization for o1-like olympiad-level mathematical reasoning","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.631264Z"},"links":{"cited_paper":"/paper/2410.02884","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:6bcab896dcfb22a60b6205b74319a025c346dd698c55d1a30d8994b6ab224286","observation_id":"407f26e5-679b-458a-860b-e946dfad7d99","resolution":{"observed_at":"2026-08-08T14:25:53.631264Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03816","last_updated":"2024-11-18T05:36:16Z","snapshot_observed_at":"2026-08-04T20:04:21.125115Z","submitted_at":"2024-06-06T07:40:00Z","title":"ReST-MCTS*: LLM Self-Training via Process Reward Guided Tree Search","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.03816","snapshot_observed_at":"2026-08-08T14:25:53.635679Z","title":"Rest-mcts*: Llm self-training via process reward guided tree search","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":87,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.635679Z"},"links":{"cited_paper":"/paper/2406.03816","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:fec0dced225f08db1a880c586c1273d5e41c5051dbb91e49dfa00c6c5b43c6df","observation_id":"e4199f7f-f4b0-424c-8252-8677147d40fd","resolution":{"observed_at":"2026-08-08T14:25:53.635679Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-08T22:57:17.515334Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition"},"reference_resolution":{"displayed":87,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":66,"verified_exact":0,"verified_fuzzy":21},"total_outbound_references":87},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 87 of 87 outbound references and 5 inbound Pith citation observations for arXiv:2502.06773."}