{"as_of":"2026-08-23T17:49:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:10f28223678a32d90368cd7601afa681bafe9b676b084a8480d1f02da97327ae","coverage":[{"denominator":51,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":51,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T21:31:38.710562Z","state":"measured"},{"denominator":90,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":90,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-23T06:30:58.430688+00:00","state":"measured"},{"denominator":39,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":39,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T11:08:27.062988Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"pith","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":2,"observed_at":"2026-08-05T02:28:24.338817Z","source":"pith"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-09T14:49:08.848136Z","title":"2, 4 Xie, Y ., Goyal, A., Zheng, W., Kan, M.-Y ., Lillicrap, T","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.01633","last_updated":"2025-06-25T15:31:17Z","snapshot_observed_at":"2026-08-16T09:09:35.076144Z","submitted_at":"2025-02-03T18:59:01Z","title":"Adversarial Reasoning at Jailbreaking Time","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-09T14:49:08.848136Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2502.01633"},"observation_digest":"sha256:1ef00bef2195c3c921db64793a00b1db76f94b09b63a30b5040b097f42ed8d12","observation_id":"8bfa863d-6e8e-441f-8bf1-1d4aeaf68823","resolution":{"observed_at":"2026-08-09T14:49:08.848136Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-09T17:37:40.799741Z","title":"Towards System 2 reasoning in LLMs: learning how to think with meta chain-of-thought","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.01694","last_updated":"2025-03-01T10:27:24Z","snapshot_observed_at":"2026-08-22T02:01:27.458290Z","submitted_at":"2025-02-02T18:19:14Z","title":"Metastable Dynamics of Chain-of-Thought Reasoning: Provable Benefits of Search, RL and Distillation","version":2},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-08-09T17:37:40.799741Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2502.01694"},"observation_digest":"sha256:bb236d4bc476e8d469229df295399b6ed2df86b26c1f7a0c164d3fd4904866f7","observation_id":"1d7d5871-f5a4-48c6-919a-71dabc5f150f","resolution":{"observed_at":"2026-08-09T17:37:40.799741Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-08T14:25:53.592132Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-though","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-13T11:53:50.444757Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.592132Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:0327315a5160367bf98ad392f688296a8830e12322d16ac7dcce2545ffb8411c","observation_id":"04a629f7-ff2c-4e73-a168-ca0f29f98e55","resolution":{"observed_at":"2026-08-08T14:25:53.592132Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-08T13:11:51.767100Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-though","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.07316","last_updated":"2025-05-21T13:38:27Z","snapshot_observed_at":"2026-08-16T15:36:15.556886Z","submitted_at":"2025-02-11T07:26:50Z","title":"CodeI/O: Condensing Reasoning Patterns via Code Input-Output Prediction","version":4},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-08T13:11:51.767100Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2502.07316"},"observation_digest":"sha256:8bb6d51ed3ecb163fd9d6a30c83f0cb0aa2a71fdd2393726ba0463dfa47311a1","observation_id":"2d42c829-c4b4-42f6-891e-3213d9c493c7","resolution":{"observed_at":"2026-08-08T13:11:51.767100Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2501.04682","doi":"10.48550/arxiv.2501.04682","metadata_source":"pith","pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":"cs.AI","work_id":"9e24c386-134f-4c23-b861-1bf5acc21088","year":2025},"citing_paper":{"arxiv_id":"2502.17419","last_updated":"2025-06-25T02:24:46Z","snapshot_observed_at":"2026-08-20T08:13:05.630967Z","submitted_at":"2025-02-24T18:50:52Z","title":"From System 1 to System 2: A Survey of Reasoning Large Language Models","version":6},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-05-13T01:36:23.845366Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2502.17419"},"observation_digest":"sha256:17a2a5178f6c5c12c28633867c3436c752293dd27d08633f644d38f5798d6fb7","observation_id":"256165f8-9fdd-418e-bf86-fc992171ebd7","resolution":{"observed_at":"2026-05-13T01:36:24.604401Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2501.04682","doi":"10.48550/arxiv.2501.04682","metadata_source":"pith","pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":"cs.AI","work_id":"9e24c386-134f-4c23-b861-1bf5acc21088","year":2025},"citing_paper":{"arxiv_id":"2504.01943","last_updated":"2025-08-07T23:04:55Z","snapshot_observed_at":"2026-08-14T07:26:32.101930Z","submitted_at":"2025-04-02T17:50:31Z","title":"OpenCodeReasoning: Advancing Data Distillation for Competitive Coding","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-17T19:21:42.081762Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2504.01943"},"observation_digest":"sha256:eec12cbbb402f92e9d1031e5afd70ee38f4596febca4f40f7535bf3593858826","observation_id":"247b2d95-bd10-49f6-a353-eb33305b3478","resolution":{"observed_at":"2026-05-17T19:21:42.185923Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-16T11:08:27.062988Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-though","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2504.16414","last_updated":"2025-08-06T11:20:14Z","snapshot_observed_at":"2026-08-16T17:16:29.664799Z","submitted_at":"2025-04-23T04:36:19Z","title":"Evaluating Multi-Hop Reasoning in Large Language Models: A Chemistry-Centric Case Study","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-16T11:08:27.062988Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2504.16414"},"observation_digest":"sha256:f44e40456b528d5f99fc38fb860c47f8279c383675d010d1ac9caed0a90e92ea","observation_id":"13d59eb0-efc0-4df1-b3e2-b9aa560dd2f3","resolution":{"observed_at":"2026-08-16T11:08:27.062988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2501.04682","doi":"10.48550/arxiv.2501.04682","metadata_source":"pith","pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":"cs.AI","work_id":"9e24c386-134f-4c23-b861-1bf5acc21088","year":2025},"citing_paper":{"arxiv_id":"2504.19678","last_updated":"2026-03-06T19:01:27Z","snapshot_observed_at":"2026-08-20T06:16:20.613961Z","submitted_at":"2025-04-28T11:08:22Z","title":"From LLM Reasoning to Autonomous AI Agents: A Comprehensive Review","version":2},"reference_index":233,"source":"pdf_text","source_observed_at":"2026-05-15T02:57:37.873567Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2504.19678"},"observation_digest":"sha256:f8bba3877e3670ac82e9c683b8ce8330723863d943a8b3b2510556868e354f80","observation_id":"4dcc5592-b690-4986-a571-0d356895bd88","resolution":{"observed_at":"2026-05-15T02:57:38.188188Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-15T23:08:00.424521Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-though","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.05408","last_updated":"2025-05-08T16:50:06Z","snapshot_observed_at":"2026-08-17T01:26:11.454863Z","submitted_at":"2025-05-08T16:50:06Z","title":"Crosslingual Reasoning through Test-Time Scaling","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:00.424521Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2505.05408"},"observation_digest":"sha256:030e2114d03e46dfb1c0c4c660498369ac9204b4a70056f68ccafa2e9578f2f9","observation_id":"007cd9e1-c964-4d72-b0c1-324c490d9653","resolution":{"observed_at":"2026-08-15T23:08:00.424521Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-15T20:51:06.668503Z","title":"Towards system 2 reasoning in LLMs: learning how to think with meta chain-of-thought","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.11866","last_updated":"2025-05-17T06:17:57Z","snapshot_observed_at":"2026-08-20T14:46:55.753408Z","submitted_at":"2025-05-17T06:17:57Z","title":"Position Paper: Bounded Alignment: What (Not) To Expect From AGI Agents","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-15T20:51:06.668503Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2505.11866"},"observation_digest":"sha256:7c7b060f5cedd8b5ee25fe2bb116142f0d8e84d67fb45266d2d946b3a2d1ddae","observation_id":"96ff42be-c498-4d2f-88a8-fb9c87f065f1","resolution":{"observed_at":"2026-08-15T20:51:06.668503Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-15T20:50:26.122396Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-though","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.12031","last_updated":"2025-05-17T14:47:36Z","snapshot_observed_at":"2026-08-19T15:38:20.338234Z","submitted_at":"2025-05-17T14:47:36Z","title":"LLM-based Automated Theorem Proving Hinges on Scalable Synthetic Data Generation","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-15T20:50:26.122396Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2505.12031"},"observation_digest":"sha256:c4f3cef9b6e6f8ca95f1b6ddd24a0d46e4f3a42e743480aebf161b858c7d0c69","observation_id":"72dc2695-d719-40a8-b2e9-7ebd9677ce84","resolution":{"observed_at":"2026-08-15T20:50:26.122396Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-07T15:42:41.245910Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.14212","last_updated":"2025-05-20T11:16:29Z","snapshot_observed_at":"2026-08-15T05:44:48.071906Z","submitted_at":"2025-05-20T11:16:29Z","title":"Automatic Dataset Generation for Knowledge Intensive Question Answering Tasks","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-07T15:42:41.245910Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2505.14212"},"observation_digest":"sha256:f8c0e204e2ae0456a20d9f7db4fae8095ee13f06f90ffde7da38cd87a4c11b05","observation_id":"1189bdc3-20cf-4933-abe3-eb27f0a95585","resolution":{"observed_at":"2026-08-07T15:42:41.245910Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-07T14:45:19.472575Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.17829","last_updated":"2025-05-23T12:42:50Z","snapshot_observed_at":"2026-08-20T18:38:54.052870Z","submitted_at":"2025-05-23T12:42:50Z","title":"Stepwise Reasoning Checkpoint Analysis: A Test Time Scaling Method to Enhance LLMs' Reasoning","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-07T14:45:19.472575Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2505.17829"},"observation_digest":"sha256:a42bf90770605b833abf196f0418461bd6eb788ae20f56172fc1db4396797f0f","observation_id":"13e43f31-d822-4b72-a7f6-63725eaf9c21","resolution":{"observed_at":"2026-08-07T14:45:19.472575Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-07T14:07:49.488600Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-though.arXiv preprint arXiv:2501.04682, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.19949","last_updated":"2025-05-26T13:15:26Z","snapshot_observed_at":"2026-08-20T03:49:01.362977Z","submitted_at":"2025-05-26T13:15:26Z","title":"Which Data Attributes Stimulate Math and Code Reasoning? An Investigation via Influence Functions","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T14:07:49.488600Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2505.19949"},"observation_digest":"sha256:5bc7c6573d4f7b888de4ab58f2d4c780f677354d4d107ce48a8203dbd864bca4","observation_id":"04384deb-b68a-49d0-904e-743caf9baf8e","resolution":{"observed_at":"2026-08-07T14:07:49.488600Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-07T14:02:52.066758Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-though,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.20223","last_updated":"2025-05-26T17:06:00Z","snapshot_observed_at":"2026-08-14T01:14:39.568702Z","submitted_at":"2025-05-26T17:06:00Z","title":"Chain-of-Thought for Autonomous Driving: A Comprehensive Survey and Future Prospects","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T14:02:52.066758Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2505.20223"},"observation_digest":"sha256:7d5e869b845bf18e0b6739431536abe81dc8ab8a7f4e0ca4e90f1c20ed8c547e","observation_id":"b8eb9a20-a1c3-4105-b6bf-bd2cdb7a5227","resolution":{"observed_at":"2026-08-07T14:02:52.066758Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-07T13:30:22.966069Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-though","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.21765","last_updated":"2025-05-27T20:59:29Z","snapshot_observed_at":"2026-08-17T14:07:17.104324Z","submitted_at":"2025-05-27T20:59:29Z","title":"Don't Think Longer, Think Wisely: Optimizing Thinking Dynamics for Large Reasoning Models","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T13:30:22.966069Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2505.21765"},"observation_digest":"sha256:d03106cba58bd3d60576e6c708ed04fd0f53897c55afbe47fc32151c011e7765","observation_id":"8f039778-8cbe-4d26-bb60-f9b10116127f","resolution":{"observed_at":"2026-08-07T13:30:22.966069Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-07T12:11:10.909618Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-though,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.00309","last_updated":"2025-06-28T05:54:45Z","snapshot_observed_at":"2026-08-09T10:36:52.271268Z","submitted_at":"2025-05-30T23:37:37Z","title":"Evaluation of LLMs for mathematical problem solving","version":3},"reference_index":107,"source":"pdf_text","source_observed_at":"2026-08-07T12:11:10.909618Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2506.00309"},"observation_digest":"sha256:aec3b6e35414273833c516298c9afc88c259ba275ad713a8bdad112028d0c954","observation_id":"1d84176a-66b9-479c-be9c-215a0d76e0bd","resolution":{"observed_at":"2026-08-07T12:11:10.909618Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-07T10:28:31.296126Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-though.arXiv preprint arXiv:2501.04682,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.05256","last_updated":"2025-06-06T02:38:39Z","snapshot_observed_at":"2026-08-22T04:55:04.355256Z","submitted_at":"2025-06-05T17:17:05Z","title":"Just Enough Thinking: Efficient Reasoning with Adaptive Length Penalties Reinforcement Learning","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T10:28:31.296126Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2506.05256"},"observation_digest":"sha256:a97922a163d3589ed08de3975d466803685af2043d78e7dee8b06f2f57750695","observation_id":"79876e20-bc1d-4a1e-be92-fa16de0f8d84","resolution":{"observed_at":"2026-08-07T10:28:31.296126Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-07T05:51:30.733420Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-though.arXiv preprint arXiv:2501.04682,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.06923","last_updated":"2025-06-07T21:23:00Z","snapshot_observed_at":"2026-08-18T04:51:46.950390Z","submitted_at":"2025-06-07T21:23:00Z","title":"Boosting LLM Reasoning via Spontaneous Self-Correction","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T05:51:30.733420Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2506.06923"},"observation_digest":"sha256:50e4674fcc0e0cb6981a90e1376937b875ca4bbc202d3c9fdea59a36b92c191a","observation_id":"45d72fcc-112d-4e52-9821-6762d0f869d5","resolution":{"observed_at":"2026-08-07T05:51:30.733420Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-07T04:38:21.180225Z","title":"Towards System 2 reasoning in LLMs: Learning how to think with meta chain-of-thought,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.10161","last_updated":"2025-06-11T20:27:08Z","snapshot_observed_at":"2026-08-16T21:31:48.736248Z","submitted_at":"2025-06-11T20:27:08Z","title":"Can LLMs Generate Good Stories? Insights and Challenges from a Narrative Planning Perspective","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T04:38:21.180225Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2506.10161"},"observation_digest":"sha256:e87a299ec248db023382c878f34f33ca1b3383be17abb9dce2ce7db8977c550a","observation_id":"a5f89e6b-bcec-48fd-9ae1-d4c9ec9005eb","resolution":{"observed_at":"2026-08-07T04:38:21.180225Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-15T20:13:19.653697Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-though.arXiv preprint arXiv:2501.04682, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12860","last_updated":"2025-06-15T14:21:28Z","snapshot_observed_at":"2026-08-16T20:27:15.337845Z","submitted_at":"2025-06-15T14:21:28Z","title":"QFFT, Question-Free Fine-Tuning for Adaptive Reasoning","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:19.653697Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2506.12860"},"observation_digest":"sha256:6aa878651ec3414ad0084eb4d61c4aa803d8e66083c5602db91cc6a2737d744d","observation_id":"0e2059b1-2d5c-40fa-9711-59a7e9ad626a","resolution":{"observed_at":"2026-08-15T20:13:19.653697Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-06T21:21:41.978513Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00417","last_updated":"2025-07-01T04:10:15Z","snapshot_observed_at":"2026-08-14T02:32:06.313774Z","submitted_at":"2025-07-01T04:10:15Z","title":"ASTRO: Teaching Language Models to Reason by Reflecting and Backtracking In-Context","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-06T21:21:41.978513Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2507.00417"},"observation_digest":"sha256:11ad76bb15782cd71ba4fdd3ee60475e8e75f0997ca29b3eac0c271921f9c968","observation_id":"1e719be4-ba21-4113-98e0-ea540a68ee6d","resolution":{"observed_at":"2026-08-06T21:21:41.978513Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-06T14:26:03.628070Z","title":", Snell, C","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.06503","last_updated":"2025-07-25T15:56:25Z","snapshot_observed_at":"2026-08-21T15:52:46.821368Z","submitted_at":"2025-07-25T15:56:25Z","title":"Understanding Human Limits in Pattern Recognition: A Computational Model of Sequential Reasoning in Rock, Paper, Scissors","version":1},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-08-06T14:26:03.628070Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2508.06503"},"observation_digest":"sha256:f515427ccd8843279fab7b838150c5ca518368df1f9460121e3e70eb139ef4e8","observation_id":"3a202bf3-a2d3-4ef3-ada5-668443b92403","resolution":{"observed_at":"2026-08-06T14:26:03.628070Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T16:15:27.565145Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.18812","last_updated":"2025-08-26T08:47:58Z","snapshot_observed_at":"2026-08-16T16:20:49.965923Z","submitted_at":"2025-08-26T08:47:58Z","title":"STARec: An Efficient Agent Framework for Recommender Systems via Autonomous Deliberate Reasoning","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-05T16:15:27.565145Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2508.18812"},"observation_digest":"sha256:542acd54508f299798a8c89a569b540c47f624f96dca7c3694327da21db3d605","observation_id":"3024f050-71e9-454f-a12e-4f5c0ca8f254","resolution":{"observed_at":"2026-08-05T16:15:27.565145Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T15:18:57.675232Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.20018","last_updated":"2025-08-27T16:27:19Z","snapshot_observed_at":"2026-08-15T13:03:03.060918Z","submitted_at":"2025-08-27T16:27:19Z","title":"SWIRL: A Staged Workflow for Interleaved Reinforcement Learning in Mobile GUI Control","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-05T15:18:57.675232Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2508.20018"},"observation_digest":"sha256:4ba606a81481609ce5dba45456d898bc37135f18c95d64ab228ecc6fdda28d0c","observation_id":"12716f24-d199-44e6-a838-6115cd11bfe8","resolution":{"observed_at":"2026-08-05T15:18:57.675232Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2501.04682","doi":"10.48550/arxiv.2501.04682","metadata_source":"pith","pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":"cs.AI","work_id":"9e24c386-134f-4c23-b861-1bf5acc21088","year":2025},"citing_paper":{"arxiv_id":"2509.23629","last_updated":"2026-05-07T02:29:03Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-09-28T04:10:37Z","title":"Emergent Slow Thinking in LLMs as Inverse Tree Freezing","version":3},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-05-18T12:43:49.628082Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2509.23629"},"observation_digest":"sha256:bc63d2ac40f298727a1a95a9e176521646f953f8e3df1fefaaf46021889b2e12","observation_id":"ab711e7b-bdbb-47bc-9862-8bbd970643c5","resolution":{"observed_at":"2026-05-18T12:46:24.254006Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-03T05:01:00.667656Z","title":"Towards system 2 reasoning in LLM s: Learning how to think with meta chain-of-though","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.03542","last_updated":"2026-06-03T16:48:57Z","snapshot_observed_at":"2026-08-15T15:14:08.381551Z","submitted_at":"2026-02-03T13:56:54Z","title":"Can Large Language Models Generalize Procedures Across Representations?","version":2},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-08-03T05:01:00.667656Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2602.03542"},"observation_digest":"sha256:3e35a21ce9d7afa74186a370a9d4708e939ee9680cf429fc718bac0f0c190efc","observation_id":"7759d371-465d-42e6-a4e0-5715fa22c1b2","resolution":{"observed_at":"2026-08-03T05:01:00.667656Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2501.04682","doi":"10.48550/arxiv.2501.04682","metadata_source":"pith","pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":"cs.AI","work_id":"9e24c386-134f-4c23-b861-1bf5acc21088","year":2025},"citing_paper":{"arxiv_id":"2602.22508","last_updated":"2026-05-11T05:22:14Z","snapshot_observed_at":"2026-08-09T18:36:08.061095Z","submitted_at":"2026-02-26T00:56:15Z","title":"Metacognitive Behavioral Tuning of Large Language Models for Multi-Hop Question Answering","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-15T19:35:14.853273Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2602.22508"},"observation_digest":"sha256:479951e0566c9046e334269457ddbc3abf30d56fa5d394276486f61d9345b78a","observation_id":"cb91586b-ee73-4db7-b113-bb1c61aaa225","resolution":{"observed_at":"2026-05-15T19:36:32.774090Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2501.04682","doi":"10.48550/arxiv.2501.04682","metadata_source":"pith","pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":"cs.AI","work_id":"9e24c386-134f-4c23-b861-1bf5acc21088","year":2025},"citing_paper":{"arxiv_id":"2604.10449","last_updated":"2026-04-12T04:15:31Z","snapshot_observed_at":"2026-08-11T00:32:22.016000Z","submitted_at":"2026-04-12T04:15:31Z","title":"AdverMCTS: Combating Pseudo-Correctness in Code Generation via Adversarial Monte Carlo Tree Search","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-05-10T16:35:16.056397Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2604.10449"},"observation_digest":"sha256:2df249d66a4e349622d66c452786e01afdb53151b7a085a88fd7e88573514fc2","observation_id":"5b04ef12-ae1c-4829-ae52-2ab6d720433b","resolution":{"observed_at":"2026-05-11T08:40:57.471120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2501.04682","doi":"10.48550/arxiv.2501.04682","metadata_source":"pith","pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":"cs.AI","work_id":"9e24c386-134f-4c23-b861-1bf5acc21088","year":2025},"citing_paper":{"arxiv_id":"2604.12816","last_updated":"2026-04-14T14:43:39Z","snapshot_observed_at":"2026-07-06T23:00:56.449281Z","submitted_at":"2026-04-14T14:43:39Z","title":"The role of System 1 and System 2 semantic memory structure in human and LLM biases","version":1},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-05-10T14:49:45.560437Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2604.12816"},"observation_digest":"sha256:5dee5203f62efcf42eed4207e55bf01e60a2f52bbc3ac730dc4d5c484235ca53","observation_id":"33d398e5-8c71-4645-b9c0-b7f2ee423585","resolution":{"observed_at":"2026-05-11T11:31:01.953584Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2501.04682","doi":"10.48550/arxiv.2501.04682","metadata_source":"pith","pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":"cs.AI","work_id":"9e24c386-134f-4c23-b861-1bf5acc21088","year":2025},"citing_paper":{"arxiv_id":"2604.17928","last_updated":"2026-04-20T08:09:01Z","snapshot_observed_at":"2026-08-14T00:24:03.820656Z","submitted_at":"2026-04-20T08:09:01Z","title":"HEALing Entropy Collapse: Enhancing Exploration in Few-Shot RLVR via Hybrid-Domain Entropy Dynamics Alignment","version":1},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-05-10T05:23:08.478393Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2604.17928"},"observation_digest":"sha256:50126ac8f3b8ab4df98656feb6d4569c0bb16018f03b27d7463d323e6777899f","observation_id":"721e6429-bce1-4b85-b247-02a43b4f540a","resolution":{"observed_at":"2026-05-10T05:25:55.197063Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2501.04682","doi":"10.48550/arxiv.2501.04682","metadata_source":"pith","pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":"cs.AI","work_id":"9e24c386-134f-4c23-b861-1bf5acc21088","year":2025},"citing_paper":{"arxiv_id":"2605.02860","last_updated":"2026-05-04T17:37:16Z","snapshot_observed_at":"2026-08-12T17:50:30.651793Z","submitted_at":"2026-05-04T17:37:16Z","title":"Standing on the Shoulders of Giants: Stabilized Knowledge Distillation for Cross--Language Code Clone Detection","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-05-08T17:59:12.781703Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2605.02860"},"observation_digest":"sha256:3a5744619440c03b2d31725d4e6aead0ccec746ee0ee679ea1d3c8442b541aca","observation_id":"2bc966ce-8d80-4f03-9f9e-fa2c590e86da","resolution":{"observed_at":"2026-05-09T06:55:41.618176Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2501.04682","doi":"10.48550/arxiv.2501.04682","metadata_source":"pith","pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":"cs.AI","work_id":"9e24c386-134f-4c23-b861-1bf5acc21088","year":2025},"citing_paper":{"arxiv_id":"2605.25745","last_updated":"2026-05-25T11:57:09Z","snapshot_observed_at":"2026-08-06T11:29:59.606951Z","submitted_at":"2026-05-25T11:57:09Z","title":"Selective Latent Thinking: Adaptive Compression of LLM Reasoning Chains","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-29T21:42:08.610000Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2605.25745"},"observation_digest":"sha256:4f73939762740378a7bf4cd2803b786c523691f0947b78de9e1ca0e626d3ca7f","observation_id":"393e81c2-649a-4ec1-9997-8d24228a4c01","resolution":{"observed_at":"2026-06-29T21:43:59.096395Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2501.04682","doi":"10.48550/arxiv.2501.04682","metadata_source":"pith","pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":"cs.AI","work_id":"9e24c386-134f-4c23-b861-1bf5acc21088","year":2025},"citing_paper":{"arxiv_id":"2606.00726","last_updated":"2026-07-10T01:15:34Z","snapshot_observed_at":"2026-08-14T07:08:33.006893Z","submitted_at":"2026-05-30T13:38:06Z","title":"Latent Reward Steering: An Adaptive Inference-Time Framework that Implicitly Promotes Cognitive Behaviors in Reasoning LLMs","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-06-28T18:49:56.917505Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2606.00726"},"observation_digest":"sha256:5a86308a87b26adcbfd123b89d7d86e737e08a0d2577cf3b7b75da5f482c6d58","observation_id":"7b945b25-ef91-42fe-93d5-30762f7f5493","resolution":{"observed_at":"2026-06-28T19:52:35.564696Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2501.04682","doi":"10.48550/arxiv.2501.04682","metadata_source":"pith","pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":"cs.AI","work_id":"9e24c386-134f-4c23-b861-1bf5acc21088","year":2025},"citing_paper":{"arxiv_id":"2606.02871","last_updated":"2026-06-01T20:36:06Z","snapshot_observed_at":"2026-07-06T23:43:12.737017Z","submitted_at":"2026-06-01T20:36:06Z","title":"Adaptive Latent Agentic Reasoning","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-06-28T14:24:22.855486Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2606.02871"},"observation_digest":"sha256:f8687d4a5fca127381f0a8cae5a60be9374cad3543f4d26627fee88ff95eb662","observation_id":"005d2321-65cf-4fe6-bfb0-abdf0b108bdb","resolution":{"observed_at":"2026-07-01T23:26:22.287107Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2501.04682","doi":"10.48550/arxiv.2501.04682","metadata_source":"pith","pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":"cs.AI","work_id":"9e24c386-134f-4c23-b861-1bf5acc21088","year":2025},"citing_paper":{"arxiv_id":"2606.05736","last_updated":"2026-06-04T05:55:15Z","snapshot_observed_at":"2026-07-06T23:45:42.379051Z","submitted_at":"2026-06-04T05:55:15Z","title":"VTI-CoT: Visual-Textual Interleaved Chain of Thought for Video Reasoning","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-06-28T01:52:44.785582Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2606.05736"},"observation_digest":"sha256:193c1140058c6f6a86df2234fd2e4d8c8f9a217e176668fa032afc7b61ef70b3","observation_id":"4da3fa61-6b23-4657-b548-f6c13b8f5957","resolution":{"observed_at":"2026-07-02T12:46:56.796999Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2501.04682","doi":"10.48550/arxiv.2501.04682","metadata_source":"pith","pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":"cs.AI","work_id":"9e24c386-134f-4c23-b861-1bf5acc21088","year":2025},"citing_paper":{"arxiv_id":"2606.11470","last_updated":"2026-08-10T19:13:02Z","snapshot_observed_at":"2026-08-14T23:09:31.705348Z","submitted_at":"2026-06-09T21:59:37Z","title":"The Periodic Table of LLM Reasoning: A Structured Survey of Reasoning Paradigms, Methods, and Failure Modes","version":1},"reference_index":270,"source":"arxiv_source","source_observed_at":"2026-06-27T12:59:51.091008Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2606.11470"},"observation_digest":"sha256:003a28060d68b06da69d80924a12f2b2efeada31d808da8a17f0af0af753df01","observation_id":"48f844cf-f839-4a2e-a473-f01a4b005141","resolution":{"observed_at":"2026-06-27T13:00:56.106338Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2501.04682","doi":"10.48550/arxiv.2501.04682","metadata_source":"pith","pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards system 2 reasoning in llms: Learning how to think with meta chain-of-thought","venue":"cs.AI","work_id":"9e24c386-134f-4c23-b861-1bf5acc21088","year":2025},"citing_paper":{"arxiv_id":"2607.06720","last_updated":"2026-07-07T18:36:04Z","snapshot_observed_at":"2026-08-19T23:14:36.284430Z","submitted_at":"2026-07-07T18:36:04Z","title":"When Does In-Context Search Help? A Sampling-Complexity Theory of Reflection-Driven Reasoning","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-07-10T22:46:24.057572Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2607.06720"},"observation_digest":"sha256:1fe320837057a10366d074e5c68768a80305e58af421f70291f3b5d5fdebe6c1","observation_id":"858cd4ae-daa6-4bae-83d1-ee2133bf73a1","resolution":{"observed_at":"2026-07-10T22:47:36.883086Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04682","snapshot_observed_at":"2026-07-30T22:49:43.471344Z","title":"arXiv preprint arXiv:2501.04682 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.23420","last_updated":"2026-07-26T02:35:47Z","snapshot_observed_at":"2026-08-22T01:15:30.699662Z","submitted_at":"2026-07-26T02:35:47Z","title":"LA-RL: Label-Aware Self-Reflection for Reinforcement Learning in Information Extraction","version":1},"reference_index":168,"source":"arxiv_source","source_observed_at":"2026-07-30T22:49:43.471344Z"},"links":{"cited_paper":"/paper/2501.04682","citing_paper":"/paper/2607.23420"},"observation_digest":"sha256:1ea5ce72612e3f0424a2ac2aaee971066ce76d174b97bf5d9e1c57835125c391","observation_id":"7ec64dce-f88e-4901-979b-9c89215a189c","resolution":{"observed_at":"2026-07-30T22:49:43.471344Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2501.04682/citation-record","integrity":"/paper/2501.04682/integrity","json":"/paper/2501.04682/citation-record.json","paper":"/paper/2501.04682"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.479242Z","title":null,"venue":null,"work_id":"33bed27b-7ed7-473d-b2bf-053a7d5501fb","year":2022},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.515566Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:825ec4601060a28aa8102b2f46eab6b0ae29dd889c2c889228ad0422afbca67a","observation_id":"55614dbc-f940-453c-bec9-a89584f90594","resolution":{"observed_at":"2026-08-10T21:31:39.483933Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.462746Z","title":"(23) In this formulation, the model will not learn to generate a solution, requiring the use of a separate step to summarize the search process into a final solution","venue":null,"work_id":"6118986a-c2c4-496a-a498-415827a6c998","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.519697Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:43289c58b9036fceefc0c83b254aa429d3240a26f1409d5aa31f1cdd4cd67d40","observation_id":"f1739aee-9a4a-41c7-a883-8a596e57e653","resolution":{"observed_at":"2026-08-10T21:31:39.468389Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.19165","last_updated":"2023-05-30T16:09:19Z","snapshot_observed_at":"2026-08-19T07:02:36.987726Z","submitted_at":"2023-05-30T16:09:19Z","title":"Strategic Reasoning with Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.19165","snapshot_observed_at":"2026-08-10T21:31:38.494670Z","title":"Kanishk Gandhi, Denise Lee, Gabriel Grand, Muxin Liu, Winson Cheng, Archit Sharma, and Noah D Goodman","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.494670Z"},"links":{"cited_paper":"/paper/2305.19165","citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:eb696741d03b173b4186f7b744e9ab6b98372f5e4b93b6c60e001a5b1f95b160","observation_id":"fd07db51-b5e2-4280-aeca-c6bb4f289934","resolution":{"observed_at":"2026-08-10T21:31:38.494670Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.434008Z","title":"simulation","venue":null,"work_id":"aa9416c8-4c5f-4072-985b-80643a3d68bc","year":2024},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.527750Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:b21e7c8f6a84a86dde5e7d3359973f8a415bc9a29bf232575b9a5d77b3b157a0","observation_id":"6a4b70d3-068d-4f5f-80dd-50c616b3973c","resolution":{"observed_at":"2026-08-10T21:31:39.438602Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14083","last_updated":"2024-04-26T21:05:19Z","snapshot_observed_at":"2026-08-17T06:38:47.622164Z","submitted_at":"2024-02-21T19:17:28Z","title":"Beyond A*: Better Planning with Transformers via Search Dynamics Bootstrapping","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14083","snapshot_observed_at":"2026-08-10T21:31:38.504790Z","title":"Sergey Levine, Aviral Kumar, George Tucker, and Justin Fu","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.504790Z"},"links":{"cited_paper":"/paper/2402.14083","citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:d2a4eba58b489961b61228c915c5c478655d0c85467da073b905c2b2b8a1b28f","observation_id":"9058365a-e13b-41f1-b6d5-d771bf840bb0","resolution":{"observed_at":"2026-08-10T21:31:38.504790Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.448001Z","title":"To handle this, we can mask incorrect steps/branches in the above loss ℒ(𝜃) = min 𝜃 −E(q,Z,S)∼𝒟train ⎡ ⎣ |Z|∑︁ 𝑖=1 𝐼{z𝑖+1∈ S} log 𝜋𝜃(z𝑖+1|Z𝑖, q) ⎤ ⎦","venue":null,"work_id":"8d8154c0-a9fc-477a-bb36-ef412594675f","year":2024},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.523645Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:48189a8c36144a78f813428fbe357dbfd5f794125a28fa2164775b46107f4adf","observation_id":"a157908c-c097-4e61-8803-b759b590d386","resolution":{"observed_at":"2026-08-10T21:31:39.452931Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.419012Z","title":null,"venue":null,"work_id":"face0e2e-add7-4736-928a-19284abe95d2","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.532361Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:35cd8d485cc9703e33d2be5be21b75b7d181cfc77b944c3f6a5c1af49fcd793a","observation_id":"372eec14-2950-48f9-ada0-808ad41c948a","resolution":{"observed_at":"2026-08-10T21:31:39.423944Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.404674Z","title":null,"venue":null,"work_id":"f0a7ce75-b04c-429f-80ce-2c1e4110e67c","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.537004Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:0a40e16d852d479b2e7291cddc081ad96c1d9c3ff5520795a393926cd3b7b0e4","observation_id":"5f3ad01d-136f-47b4-b8b2-1e1df8e956b7","resolution":{"observed_at":"2026-08-10T21:31:39.408944Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.390771Z","title":null,"venue":null,"work_id":"7166b4c4-ad61-4bb3-b5bb-405843c41824","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.541700Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:3536cab7ed218d317a6f5d882330de603ff135e10a1f63bb6c1a4884c3cf500a","observation_id":"0cfdc1eb-392c-4ad9-bd97-aef2f0281601","resolution":{"observed_at":"2026-08-10T21:31:39.395041Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.375865Z","title":"0.0 Alternatively","venue":null,"work_id":"91ba8c48-9d9c-4ac0-966c-cbe4c00bfb4a","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.546254Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:8dceaf7dd6150cb3d7090625cd36caf69d9400b6c8993ef083db68befdf9de09","observation_id":"46dbb544-ce2f-4977-b924-19a625fa92a3","resolution":{"observed_at":"2026-08-10T21:31:39.381115Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.361486Z","title":null,"venue":null,"work_id":"1b9518e3-c8df-4ea5-892e-743226bc0839","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.550908Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:a4d5c0ccd960d2b7d0dff20fa2e5f512bc023ccfe34d6f42950e55e28859caa7","observation_id":"372d4e9b-c6d4-4b25-902b-ee2a9f1f1ee1","resolution":{"observed_at":"2026-08-10T21:31:39.365929Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.347312Z","title":null,"venue":null,"work_id":"3fc10392-bf38-4d05-b5d1-30a33095ee01","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.556593Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:17b953e47ef12585c58ae05201b9ec878dc3f63db8bb9a4b07aa000a56585d93","observation_id":"68e1aced-dd97-435d-a3bb-a9c271376a7b","resolution":{"observed_at":"2026-08-10T21:31:39.351899Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.326709Z","title":null,"venue":null,"work_id":"d759ca07-423f-479d-a7be-66487ae874c6","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.560945Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:22a702c87572cf3383add2f93e845d3b8adaa0fe9a4a5bfaf6678ecb4c0ce8fc","observation_id":"2ff2bd1f-5e4e-4c10-8a93-891f477faf26","resolution":{"observed_at":"2026-08-10T21:31:39.336955Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.311838Z","title":"0.0 Alternatively","venue":null,"work_id":"c428085f-dd95-419d-b865-5adda3133a21","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.565173Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:ceb37bdfb787c2e87713133691494a3628e79fee5ea796412be0b0421e4cd973","observation_id":"53b19ea1-3505-4aa9-828a-351233c962c6","resolution":{"observed_at":"2026-08-10T21:31:39.316704Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.296730Z","title":null,"venue":null,"work_id":"1a787268-8686-4333-80b6-517fe73abd83","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.569735Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:9ad002715df818ad6b414928bac0e4fd4f7771a3bba9dcc66338b19e93818b3b","observation_id":"c5fcc718-0f9b-467b-8900-f561b9e35279","resolution":{"observed_at":"2026-08-10T21:31:39.301040Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.282501Z","title":"This pattern continues until the last member, who has only 1 option left","venue":null,"work_id":"eb91a88a-3743-48cd-abbf-32b00de67ce2","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.574262Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:6020ef7e94cf40d6c766021582a6cc26207a548073405a833e39f13314b062cb","observation_id":"0d67f6a9-ba90-4aa0-8972-30601a2a67a2","resolution":{"observed_at":"2026-08-10T21:31:39.287054Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.269478Z","title":null,"venue":null,"work_id":"2c127bd1-4569-4181-b2e1-df4c57170d3a","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.578776Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:f663d118c2ddbc57db33ac8901af30e3d3f606ae1ac8f6c7a85eed3458c0d404","observation_id":"c86f69eb-53a4-4919-816f-a553bbb88e58","resolution":{"observed_at":"2026-08-10T21:31:39.273510Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.254922Z","title":null,"venue":null,"work_id":"71ebbb69-1f3f-45a4-8e29-7ebf82f64255","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.582967Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:4fffa72f00a22933c3134f6bd3ac62437c88bf07e80c614c9a193173dcfa4a05","observation_id":"4a8c7fa7-c3ed-424c-a9de-0b920e0af3e8","resolution":{"observed_at":"2026-08-10T21:31:39.259657Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.237068Z","title":"This means that the second person has 11 choices","venue":null,"work_id":"bd182bbe-5a9a-4849-a14b-2c567349730e","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.587638Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:4de699156ec5248094661024fd9b203a1ccd6e05bef10db682e404ec7af890d7","observation_id":"8ee345e1-1ed8-457e-b306-19bc7b0be126","resolution":{"observed_at":"2026-08-10T21:31:39.242004Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.219072Z","title":"0.0 65 Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought Alternatively","venue":null,"work_id":"06d56e54-c6bb-40bc-b983-1b8bfb71b4fc","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.592048Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:098813fdf172004dde0fcaa52bf71cd471c69a1196914a1399fbe3d881034fe5","observation_id":"5a9e06f3-d4c5-45e6-b3f7-26b668015c22","resolution":{"observed_at":"2026-08-10T21:31:39.223766Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.204469Z","title":null,"venue":null,"work_id":"f76b9d70-c06a-4bec-8fbe-09d6ac17b49d","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.596203Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:2934ccbea4d390cf270d4bf1054a8176f39d3e05647f62bf08d7c941c5e084da","observation_id":"f8af18b5-386c-4418-8023-e132c72d2782","resolution":{"observed_at":"2026-08-10T21:31:39.208984Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.190060Z","title":"This gives 2 options for the second person","venue":null,"work_id":"a2ad0bf9-41b9-40c0-a666-5a3dd51d9fc0","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.600151Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:69d34c43cd1f462887af2111f8b27e6ea7e319fd15be318fab48721c961babde","observation_id":"1710cfb2-87ee-4d80-a304-724156f987e5","resolution":{"observed_at":"2026-08-10T21:31:39.194991Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.175901Z","title":"This gives 2 options for the third person, but we need to consider the case where the second person moved, so there's only 1 additional option for the third person","venue":null,"work_id":"9d3615e5-8e6e-4801-ba4c-ac655082cc8d","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.604540Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:ab9766fa36ca14cb4839443e6a7f61936d735436951f168da126a36e823c418f","observation_id":"81cb25c4-cf5f-4960-a5d9-04ce547c64f2","resolution":{"observed_at":"2026-08-10T21:31:39.180612Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.160713Z","title":"decision point","venue":null,"work_id":"b27df32b-73b3-4d05-a1a7-54ec1347111e","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.608351Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:ca57db0eb3797b15f5333123948a1879eed35853292045a19b2dc5869b187ef8","observation_id":"22108f74-f6ee-4d8f-bbf8-d7a313f9f3fc","resolution":{"observed_at":"2026-08-10T21:31:39.165235Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.145519Z","title":null,"venue":null,"work_id":"de4c0f13-5d8d-4500-9cf8-ba17ec2f5b10","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.612573Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:b95902acafc9a23a69612834d4f7d587cf5be00e9cebdce63c29078d604a80e2","observation_id":"59c0539f-70be-4a9e-951c-ece3feaa94f0","resolution":{"observed_at":"2026-08-10T21:31:39.150512Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.130472Z","title":null,"venue":null,"work_id":"6f599f47-bd90-44e2-b9eb-eb9097bd2dff","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.616753Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:db90d4e47582584bf5a279e5c5fca3eee8ac9addc2a408b6dd9e3da57497f485","observation_id":"2639d50b-645c-48b5-acd1-ebc65b2db4b2","resolution":{"observed_at":"2026-08-10T21:31:39.135063Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.116013Z","title":"0.046875 The pattern continues until the last person, who always has only 1 option","venue":null,"work_id":"1ed4fc75-9903-4eb4-882b-4233daae00e4","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.620881Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:882cf30ec65f955e9c55b797fc11a540d94504328e81da6b556e013d2daab8fc","observation_id":"ab244d9a-f143-46d1-9777-c861ef48536c","resolution":{"observed_at":"2026-08-10T21:31:39.120590Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.103043Z","title":null,"venue":null,"work_id":"476a4005-c29f-4b18-9983-4588e546eab2","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.625120Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:0bd5f17afce7e24a48e96201648bbdfb19aacdb55237208c117a7994fcb081c6","observation_id":"84172cab-feaf-47d5-843d-d368ce51b74b","resolution":{"observed_at":"2026-08-10T21:31:39.107084Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.088067Z","title":null,"venue":null,"work_id":"2cb6ecb6-dc22-4170-901e-f2ca8d27730f","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.629663Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:a20441ed39bc3070b4dfabca8a0808c21d29d81bce0b762f3c2964b68bb35e9a","observation_id":"964890b6-2971-4cf2-ab7d-76b9652623ff","resolution":{"observed_at":"2026-08-10T21:31:39.092497Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.073637Z","title":"0.015625 To determine the number of ways the jury can be seated, we need to consider the following steps: 0.01171875","venue":null,"work_id":"fa7817b9-fe05-4026-b33b-ff49a07b4e47","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.634031Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:95085e1e5061614f9525472fc8028df063f7ecca9ec169759565d2b7172adba6","observation_id":"54407608-aff4-46b7-8b08-3bb90de9e30c","resolution":{"observed_at":"2026-08-10T21:31:39.078587Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.044093Z","title":null,"venue":null,"work_id":"b3694739-0eae-4c4c-a6c4-83b8555fc7d2","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.642918Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:4a4d55cb850fe425e5aa393826e50243ec813cc0f738f35a585925b4565a0a41","observation_id":"9edd5251-bb20-4176-beed-527befbec73c","resolution":{"observed_at":"2026-08-10T21:31:39.048760Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.029107Z","title":"0.046875 Let's represent this using Python code to calculate the number of different ways the jury can be seated","venue":null,"work_id":"ddc03c59-df86-4b3a-84e0-43f00f9b4e73","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.647032Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:f31d566e157afb89f956efc02d30fc60ea1c258fddb57e80ae4ea4bceed39830","observation_id":"f843ece3-a5f0-476a-8851-300337fea477","resolution":{"observed_at":"2026-08-10T21:31:39.033971Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.014212Z","title":null,"venue":null,"work_id":"604fac4a-2982-4bfa-b0c6-5b5ce5b92fb0","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.651273Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:3d13631ec3e64ddfbc6c0467a75410df204178fc4759ae0d7f057a1aad89d233","observation_id":"99ef2542-0d99-48a2-83ca-899f0bbe0cba","resolution":{"observed_at":"2026-08-10T21:31:39.019360Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:38.999502Z","title":"0.04296875 It seems there was an issue with the initial state of the dynamic programming approach","venue":null,"work_id":"de040f8d-c940-4514-a85b-3b0db3cf5703","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.655600Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:4c3014a19072089817614070fee99d8c608387916772670d628b0e0ecefb2646","observation_id":"847cf571-9eda-42cf-8641-a994965fdbeb","resolution":{"observed_at":"2026-08-10T21:31:39.004840Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:38.984336Z","title":null,"venue":null,"work_id":"6bd517bb-5151-4c1c-b6ab-ddd90a82afac","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.659984Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:6f2fbffb79150ed9d4f73e2475f447f6b55138b450c2687de79e44bf28eda502","observation_id":"d594a3c8-5354-411a-9b9e-0eca6e19a043","resolution":{"observed_at":"2026-08-10T21:31:38.989182Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:38.968988Z","title":"0.046875 Let's re-implement the logic more carefully: 0.046875","venue":null,"work_id":"54d2015f-c36a-4692-9d5a-d1ed634468f1","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.664398Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:fa8faa7ffe470d5fc08faa848f336e72d776750daacbd9be1398ddaf894289df","observation_id":"2e095092-6968-4035-a66c-c49fd7e73224","resolution":{"observed_at":"2026-08-10T21:31:38.974005Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:38.952036Z","title":null,"venue":null,"work_id":"b36e1c6b-845a-40cf-83b0-5816a029f04e","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.668764Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:bd208ac8c01e39905e0eafd92a5032c1ebf373ec5066779bf062643646f905a6","observation_id":"d3cf6c9a-28d6-4732-a081-0d7844f103a8","resolution":{"observed_at":"2026-08-10T21:31:38.957889Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:38.933493Z","title":"0.04296875 Let's reconsider the logic and the initialization of the DP table: 0.0546875 Here's an updated approach: 0.04296875","venue":null,"work_id":"9f876133-3ba0-49b0-a500-768b4ccd7796","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.673034Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:eea3dcebb8e7966aa5284b40b95bbf1cac1d2ca18ef7fc8b7d3d3e0cc6989e75","observation_id":"8d793347-f22e-4238-8ee9-c5ef3c27118e","resolution":{"observed_at":"2026-08-10T21:31:38.940590Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:38.919732Z","title":null,"venue":null,"work_id":"5d746bd5-9bd2-4d54-a876-fb165410da4d","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.677258Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:36b3dc713bd30b274a673321b21e561cf802d6e2d2673198d39c027a1023f82d","observation_id":"3c73b5d9-45fa-419b-b98b-3c2f91504dae","resolution":{"observed_at":"2026-08-10T21:31:38.923628Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:38.906024Z","title":null,"venue":null,"work_id":"eab1eb17-eb5c-4561-9e0c-fd8f58b18c28","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.680917Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:4e1a4aa34e25be464fc97d0e605162ebc986feb660ea7bbf2674437c85e3d4aa","observation_id":"a646f3f9-bfb3-402c-a2bc-18860d48b7e5","resolution":{"observed_at":"2026-08-10T21:31:38.910307Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:38.890043Z","title":null,"venue":null,"work_id":"43cbea95-239f-432c-8c51-50d4323a0043","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.684987Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:9739cc01a048a4e27ea29f46871f37305c7a64ef519b7e5ec8d87f3ab9e012cb","observation_id":"589a7009-0b52-4e60-95ae-98c6cc1c870c","resolution":{"observed_at":"2026-08-10T21:31:38.895776Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.058920Z","title":null,"venue":null,"work_id":"74116a6c-e1db-4797-b33a-e2bf378536fe","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.688955Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:bcaa7fa3914a0ceb9d90965909db6f782acb66d7a9571ee60ad861a36c57c1ee","observation_id":"d47ca26d-d35b-4556-9d0f-0bf7ccbd69c5","resolution":{"observed_at":"2026-08-10T21:31:39.063794Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:38.874483Z","title":null,"venue":null,"work_id":"b86d2ba3-02cf-4d92-a4e8-da57a082c631","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.692787Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:d3aa34a92e748ab82ce845ad4904aa0c6bd19a1daf1d921083b0b6ec22f1ec3b","observation_id":"ce512472-3a06-44bd-ac7f-da76aff519fb","resolution":{"observed_at":"2026-08-10T21:31:38.879037Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:38.859369Z","title":null,"venue":null,"work_id":"72d2d358-7360-43ca-ad18-a4b037a79e16","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.696779Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:cb51f67499cb7ae41cba93c1e37eecee3a0a04daf9cbf9d4e74a32168051e0c2","observation_id":"0168eada-a72a-463f-a843-12970bc78d3e","resolution":{"observed_at":"2026-08-10T21:31:38.863981Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:38.844598Z","title":null,"venue":null,"work_id":"fb35c083-852c-433e-bbbe-fea26f683eab","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.700711Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:4ca25c71c3726ae22ccfd2fe9ca2354ec87d0aaa4adb39a9130018ae888bdc3f","observation_id":"e047fcfc-c841-4a25-88be-759ef3f3c423","resolution":{"observed_at":"2026-08-10T21:31:38.849136Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:38.830041Z","title":null,"venue":null,"work_id":"b1b0a04f-85ce-4a96-b04d-389824eb12b8","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.705854Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:cf6fe4a6bbeb6c46ca3c2fc2350459f1cae3372ff7472085842dcc48537f0de0","observation_id":"85f59b34-7740-40ce-b5f0-b43b191c126c","resolution":{"observed_at":"2026-08-10T21:31:38.834630Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:38.813925Z","title":"other real numbers","venue":null,"work_id":"9a2e27e0-7594-490d-8057-59c983d3ddf4","year":null},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.710562Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:1d76d97583b8d56b58f1191629418ad55a7bc2c3e44ca01bdf56ed82e83fd489","observation_id":"cd65c577-2a5a-492c-9062-06931078bbc2","resolution":{"observed_at":"2026-08-10T21:31:38.819603Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1312.6114","last_updated":"2022-12-10T21:04:00Z","snapshot_observed_at":"2026-08-14T23:50:45.029465Z","submitted_at":"2013-12-20T20:58:10Z","title":"Auto-Encoding Variational Bayes","version":11},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1312.6114","snapshot_observed_at":"2026-08-10T21:31:38.499956Z","title":"Levente Kocsis and Csaba Szepesvári","venue":null,"work_id":null,"year":2006},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":2013,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.499956Z"},"links":{"cited_paper":"/paper/1312.6114","citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:d10cc1e1d859e499963b7d330a99eb5b9d3c4b6bbc7453fc6ae0ca304359e214","observation_id":"6db28aa4-7bde-4002-a0fa-75a77a7a52c8","resolution":{"observed_at":"2026-08-10T21:31:38.499956Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21787","last_updated":"2024-12-30T19:03:24Z","snapshot_observed_at":"2026-08-14T10:00:52.343929Z","submitted_at":"2024-07-31T17:57:25Z","title":"Large Language Monkeys: Scaling Inference Compute with Repeated Sampling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21787","snapshot_observed_at":"2026-08-10T21:31:38.483287Z","title":"Anonymous","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.483287Z"},"links":{"cited_paper":"/paper/2407.21787","citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:e64cf1e40cc66bfebf211709e52c4feafc9af3ad6dcf2e6a768df9e7ca63b3a9","observation_id":"0148582c-79e7-4f2f-b802-7749504fb2bc","resolution":{"observed_at":"2026-08-10T21:31:38.483287Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:31:39.493646Z","title":"Jeffrey L Elman","venue":null,"work_id":"1cc968e3-6a01-4abd-9c12-1a803e647250","year":1990},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.489647Z"},"links":{"citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:6028b6965e7cccbf5e7b42bc16b15c3cbd92e3403d03f8e824cf0835f65802d3","observation_id":"b0f33ea5-0f85-459c-9a59-37c5aef995cb","resolution":{"observed_at":"2026-08-10T21:31:39.498599Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-10T21:31:38.510018Z","title":"Zhihong Shao, Peiyi Wang, Qihao Zhu, Runxin Xu, Junxiao Song, Xiao Bi, Haowei Zhang, Mingchuan Zhang, Y","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-10T21:31:38.510018Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2501.04682"},"observation_digest":"sha256:40fb52fdb31a9b5f312cafffbf2f9e0ad57378fa5c937f05feb72bab1aa598c6","observation_id":"3c2b4f18-2d8a-4f10-a38a-285196c12f6a","resolution":{"observed_at":"2026-08-10T21:31:38.510018Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2501.04682","last_updated":"2025-01-08T18:42:48Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-19T22:21:38.527415Z","submitted_at":"2025-01-08T18:42:48Z","title":"Towards System 2 Reasoning in LLMs: Learning How to Think With Meta Chain-of-Thought"},"reference_resolution":{"displayed":51,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":32,"verified_exact":0,"verified_fuzzy":19},"total_outbound_references":51},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"thesis":"As of 23 August 2026, this Paper Citation Record lists 51 of 51 outbound references and 39 inbound Pith citation observations for arXiv:2501.04682."}