{"as_of":"2026-08-07T23:50:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:e0b48a3ff8925248a26f4ce7a605f894208e76b1be0163d33ea0f0b10ae9ea68","coverage":[{"denominator":39,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":39,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T21:13:13.121454Z","state":"measured"},{"denominator":41,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":41,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T00:24:20.051022Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-22T05:51:08.098335Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"cited_work":{"arxiv_id":"2507.00711","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.00711","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Large reasoning mod- els are not thinking straight: on the unreliability of thinking trajectories","venue":null,"work_id":"57f5e5d9-8c22-45d3-9850-fac7fec6211c","year":2025},"citing_paper":{"arxiv_id":"2605.22211","last_updated":"2026-05-21T09:16:27Z","snapshot_observed_at":"2026-07-06T23:32:35.159280Z","submitted_at":"2026-05-21T09:16:27Z","title":"CLORE: Content-Level Optimization for Reasoning Efficiency","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-22T05:50:23.111591Z"},"links":{"cited_paper":"/paper/2507.00711","citing_paper":"/paper/2605.22211"},"observation_digest":"sha256:ad2b5c9aeb9ffff81a27f8bf9f0f28b9a70504cdf8fffe3c3fce546974bc809d","observation_id":"346cb749-5c01-423c-9f43-4a55b660d881","resolution":{"observed_at":"2026-05-22T05:51:08.101200Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.00711","snapshot_observed_at":"2026-08-06T00:24:20.051022Z","title":"Large reasoning mod- els are not thinking straight: on the unreliability of thinking trajectories.arXiv preprint arXiv:2507.00711, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.01319","last_updated":"2026-08-02T15:41:12Z","snapshot_observed_at":"2026-08-07T12:53:08.016668Z","submitted_at":"2026-08-02T15:41:12Z","title":"Cognitive Demand Steering for Adaptive Meta-Reasoning in Large Language Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T00:24:20.051022Z"},"links":{"cited_paper":"/paper/2507.00711","citing_paper":"/paper/2608.01319"},"observation_digest":"sha256:0658f04a76f533611ac62419fdd4bedc87af2de5d11d5f022a0eb817e2001870","observation_id":"bb111e22-3a5f-45b7-8ec3-b659771fe20d","resolution":{"observed_at":"2026-08-06T00:24:20.051022Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2507.00711/citation-record","integrity":"/paper/2507.00711/integrity","json":"/paper/2507.00711/citation-record.json","paper":"/paper/2507.00711"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:12:45.030972Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-06T21:12:45.030972Z"},"links":{"citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:e4f17d8605b9a6c061abab947ca7722c91845b44059a1ec933cfbef1cfbd834b","observation_id":"0ce28430-1447-4a20-b1cc-6df3890b4f6e","resolution":{"observed_at":"2026-08-06T21:12:45.030972Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.02151","last_updated":"2025-04-17T18:55:45Z","snapshot_observed_at":"2026-07-06T17:54:43.685212Z","submitted_at":"2024-04-02T17:58:27Z","title":"Jailbreaking Leading Safety-Aligned LLMs with Simple Adaptive Attacks","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.02151","snapshot_observed_at":"2026-08-06T21:13:12.722979Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.722979Z"},"links":{"cited_paper":"/paper/2404.02151","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:71daebc52a790646118c13fcdee2d27f571ae4376f3fdd2f7bfa42933a12d3a1","observation_id":"96204faa-602b-40d8-8148-6bd242a9b07a","resolution":{"observed_at":"2026-08-06T21:13:12.722979Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.11791","last_updated":"2024-08-21T17:24:15Z","snapshot_observed_at":"2026-07-06T19:04:09.389583Z","submitted_at":"2024-08-21T17:24:15Z","title":"Critique-out-Loud Reward Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.11791","snapshot_observed_at":"2026-08-06T21:13:12.739806Z","title":"Chang, and Prithviraj Ammanabrolu","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.739806Z"},"links":{"cited_paper":"/paper/2408.11791","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:80b51b0ba06e4b5f38a30c7aadeaa0ed2091b6cefb191b64c9c01ba684c48edd","observation_id":"54cf2e96-04ce-4fd7-abac-42db1f557aa2","resolution":{"observed_at":"2026-08-06T21:13:12.739806Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.15074","last_updated":"2023-10-23T11:55:58Z","snapshot_observed_at":"2026-08-03T17:43:44.363296Z","submitted_at":"2023-05-24T11:55:59Z","title":"Have LLMs Advanced Enough? A Challenging Problem Solving Benchmark For Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.15074","snapshot_observed_at":"2026-08-06T21:13:12.754676Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.754676Z"},"links":{"cited_paper":"/paper/2305.15074","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:9a36bca5f5a4a79227cf10aa64186e7c94a10ca9e5e89448cc0356734cd16b5c","observation_id":"afe36068-81f1-4725-be62-bd3f6346308b","resolution":{"observed_at":"2026-08-06T21:13:12.754676Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.12499","last_updated":"2023-09-21T21:45:17Z","snapshot_observed_at":"2026-07-06T16:22:05.800327Z","submitted_at":"2023-09-21T21:45:17Z","title":"CodePlan: Repository-level Coding using LLMs and Planning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.12499","snapshot_observed_at":"2026-08-06T21:13:12.780233Z","title":"Ashok, and Shashank Shet","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.780233Z"},"links":{"cited_paper":"/paper/2309.12499","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:d5315e78b140851a376f0e1ef80c24984897a952fc66ed68739cb9339ccc1630","observation_id":"b140e735-dc0f-4d79-acd9-c8039dcdf2d8","resolution":{"observed_at":"2026-08-06T21:13:12.780233Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:13:13.981188Z","title":null,"venue":null,"work_id":"7e0ae322-504a-4880-9cd4-04a34f6d369c","year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.797219Z"},"links":{"citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:9df1ed10ac0e0f473718c6fe86533f8d51315b9feac1a6d575333213a546e784","observation_id":"27a800bb-65c4-4fa6-ae2a-48e29170def8","resolution":{"observed_at":"2026-08-06T21:13:14.006360Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.21187","last_updated":"2025-02-01T07:57:37Z","snapshot_observed_at":"2026-08-01T16:43:44.704797Z","submitted_at":"2024-12-30T18:55:12Z","title":"Do NOT Think That Much for 2+3=? On the Overthinking of o1-Like LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.21187","snapshot_observed_at":"2026-08-06T21:13:12.810430Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.810430Z"},"links":{"cited_paper":"/paper/2412.21187","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:997ce887eb70953ff53d85153eae95a38aadc13d5a3dda9e3f89ceba9a7fcd0e","observation_id":"28e052de-9fed-4c6e-8f18-652c624f2982","resolution":{"observed_at":"2026-08-06T21:13:12.810430Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.00114","last_updated":"2024-08-07T00:52:07Z","snapshot_observed_at":"2026-07-06T18:55:15.704075Z","submitted_at":"2024-07-31T18:47:11Z","title":"Inductive or Deductive? Rethinking the Fundamental Reasoning Abilities of LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.00114","snapshot_observed_at":"2026-08-06T21:13:12.820950Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.820950Z"},"links":{"cited_paper":"/paper/2408.00114","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:0db830b070377bd97e8d094b66c71c417913cede5d264056803cdafb0d2cf2bc","observation_id":"33e92626-dde2-4ddb-87fd-6ed17f0d0c79","resolution":{"observed_at":"2026-08-06T21:13:12.820950Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.08600","last_updated":"2023-10-04T13:17:38Z","snapshot_observed_at":"2026-07-06T16:19:05.495349Z","submitted_at":"2023-09-15T17:56:55Z","title":"Sparse Autoencoders Find Highly Interpretable Features in Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.08600","snapshot_observed_at":"2026-08-06T21:13:12.829504Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.829504Z"},"links":{"cited_paper":"/paper/2309.08600","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:de21a2cbb775d1e45e940b230d55b6b35b4359fef41db7879f001c3950a8f1ce","observation_id":"40829b0d-c122-4f7c-90c6-6caa63f73831","resolution":{"observed_at":"2026-08-06T21:13:12.829504Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-06T21:13:12.842189Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.842189Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:eda0319cc5b3fb4213087805163297637db61158b36bf4a577e42f3b288fd4bf","observation_id":"e9a7d6db-5477-4974-9125-92b73594a13a","resolution":{"observed_at":"2026-08-06T21:13:12.842189Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11944","last_updated":"2024-11-06T22:37:30Z","snapshot_observed_at":"2026-08-04T08:40:10.913790Z","submitted_at":"2024-06-17T17:49:00Z","title":"Transcoders Find Interpretable LLM Feature Circuits","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.11944","snapshot_observed_at":"2026-08-06T21:13:12.849806Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.849806Z"},"links":{"cited_paper":"/paper/2406.11944","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:ed8c895008859af97f723b3d6b07327c2e714262d8e15fc32ad1c890237c675b","observation_id":"85e77101-b6ba-4207-bf9b-8c9ccf225721","resolution":{"observed_at":"2026-08-06T21:13:12.849806Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.01707","last_updated":"2024-12-25T13:32:54Z","snapshot_observed_at":"2026-07-06T19:26:22.646942Z","submitted_at":"2024-10-02T16:15:31Z","title":"Interpretable Contrastive Monte Carlo Tree Search Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.01707","snapshot_observed_at":"2026-08-06T21:13:12.856172Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.856172Z"},"links":{"cited_paper":"/paper/2410.01707","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:848a7a2f842066bb229742541152e8af4408361f47a9e3fd3639709176be5c96","observation_id":"286efa0d-5526-4a3c-879d-0ff2fbc83719","resolution":{"observed_at":"2026-08-06T21:13:12.856172Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.11651","last_updated":"2025-06-13T16:15:45Z","snapshot_observed_at":"2026-08-02T09:44:45.484509Z","submitted_at":"2025-01-20T18:33:33Z","title":"T1: Advancing Language Model Reasoning through Reinforcement Learning and Inference Scaling","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.11651","snapshot_observed_at":"2026-08-06T21:13:12.864750Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.864750Z"},"links":{"cited_paper":"/paper/2501.11651","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:244d25a84183ee7984a79b5ae887495169cbe8720d7e7277b09b63119cbd3e5f","observation_id":"c480f1c5-9c68-4f25-9b8e-77c4d3e525d8","resolution":{"observed_at":"2026-08-06T21:13:12.864750Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.01817","last_updated":"2024-06-12T01:13:11Z","snapshot_observed_at":"2026-08-01T23:26:05.428862Z","submitted_at":"2024-02-02T14:43:18Z","title":"LLMs Can't Plan, But Can Help Planning in LLM-Modulo Frameworks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.01817","snapshot_observed_at":"2026-08-06T21:13:12.873393Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.873393Z"},"links":{"cited_paper":"/paper/2402.01817","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:ce4509b4fbe5d528e929751ac0f16531f7f3cd7f640b89673eda48db0110e624","observation_id":"7ceaeb7e-002e-4659-9f89-320e7d1bfcba","resolution":{"observed_at":"2026-08-06T21:13:12.873393Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15124","last_updated":"2025-04-14T22:39:09Z","snapshot_observed_at":"2026-07-06T19:55:37.400185Z","submitted_at":"2024-11-22T18:44:04Z","title":"Tulu 3: Pushing Frontiers in Open Language Model Post-Training","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.15124","snapshot_observed_at":"2026-08-06T21:13:12.896156Z","title":"Miranda, Alisa Liu, Nouha Dziri, Shane Lyu, Yuling Gu, Saumya Malik, Victoria Graf, Jena D","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.896156Z"},"links":{"cited_paper":"/paper/2411.15124","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:45650a317729905bf76a252e6358da2d7ea2a38056ad1796f86366365a69d02d","observation_id":"20b71181-3b55-44b3-9e09-d7e19684da4b","resolution":{"observed_at":"2026-08-06T21:13:12.896156Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.05147","last_updated":"2024-08-19T07:51:05Z","snapshot_observed_at":"2026-08-04T11:44:14.524984Z","submitted_at":"2024-08-09T16:06:42Z","title":"Gemma Scope: Open Sparse Autoencoders Everywhere All At Once on Gemma 2","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.05147","snapshot_observed_at":"2026-08-06T21:13:12.903707Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.903707Z"},"links":{"cited_paper":"/paper/2408.05147","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:d006fa5d43bf77c553ed1a8af2b3a650b32acf77513c67c181b9fcb0c66111f0","observation_id":"2ea5ab97-a966-4569-b43e-1040cb7710ba","resolution":{"observed_at":"2026-08-06T21:13:12.903707Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:13:12.912119Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.912119Z"},"links":{"citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:f8adab39bbfbf24f7f13215edda53146930f9025eecf63889c814b9f28e8b352","observation_id":"da248f8f-3eb0-4973-a573-60992ab917bf","resolution":{"observed_at":"2026-08-06T21:13:12.912119Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:13:13.903934Z","title":"Tang, Manan Roongta, Colin Cai, Jeffrey Luo, Li Erran Li, Raluca Ada Popa, and Ion Stoica","venue":null,"work_id":"35b735c6-20fd-4ea4-911c-4d406045d68f","year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.920953Z"},"links":{"citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:9806204ae614717359cc4d711f0f6f73678a27c79563076366271a8aea2db1e0","observation_id":"07a0e4b0-5561-4412-bd97-e5488e71df49","resolution":{"observed_at":"2026-08-06T21:13:13.934844Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.09858","last_updated":"2025-04-14T04:08:16Z","snapshot_observed_at":"2026-08-07T16:06:02.705385Z","submitted_at":"2025-04-14T04:08:16Z","title":"Reasoning Models Can Be Effective Without Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.09858","snapshot_observed_at":"2026-08-06T21:13:12.934750Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.934750Z"},"links":{"cited_paper":"/paper/2504.09858","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:bba51c45114debcee8eea7d58b8edc1c17aa50ce90a0a5cde448199b62d412f2","observation_id":"8982f4d3-cdaf-442a-8f71-4c6e704d8f9a","resolution":{"observed_at":"2026-08-06T21:13:12.934750Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.19393","last_updated":"2025-03-01T06:07:39Z","snapshot_observed_at":"2026-07-06T20:29:11.710285Z","submitted_at":"2025-01-31T18:48:08Z","title":"s1: Simple test-time scaling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.19393","snapshot_observed_at":"2026-08-06T21:13:12.943940Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.943940Z"},"links":{"cited_paper":"/paper/2501.19393","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:f4694fd6e43a6d9565367f23c5ff8f87eeaed1aa1eacc8eea0cd66f1f414d797","observation_id":"97966a76-53d8-4352-83d2-13ade9108578","resolution":{"observed_at":"2026-08-06T21:13:12.943940Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:13:13.828864Z","title":null,"venue":null,"work_id":"fb8c5a40-a76f-48be-883e-d2744fe8d17f","year":2024},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.956511Z"},"links":{"citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:63ddf4428387e1a5fda77ab1bdc841584b63bb713ece531b0413a934b9f5d118","observation_id":"c3ff820c-9dc1-4636-ba4b-cad3612ee30f","resolution":{"observed_at":"2026-08-06T21:13:13.861666Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:13:13.754651Z","title":null,"venue":null,"work_id":"2338661a-35db-466d-b4a0-5f03b0c97ac4","year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.962582Z"},"links":{"citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:c69d9a530cbc296460e89f5f68be47f01cbd0696720eee9f9a5f5abf3c43dac7","observation_id":"c2a09a0e-f022-4ddd-82f0-b58aeceb7866","resolution":{"observed_at":"2026-08-06T21:13:13.787663Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:13:13.684468Z","title":null,"venue":null,"work_id":"77cd435b-6a32-414f-83a3-be54d36fd39e","year":2024},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.970392Z"},"links":{"citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:19409589949a40d6e615a0e4bb0e3bd616681115378c475abc36594e28ef124d","observation_id":"f8a169c1-74fa-4ee8-b01a-005438ee54b1","resolution":{"observed_at":"2026-08-06T21:13:13.715774Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20304","last_updated":"2024-05-30T17:50:04Z","snapshot_observed_at":"2026-08-04T05:00:56.526394Z","submitted_at":"2024-05-30T17:50:04Z","title":"Group Robust Preference Optimization in Reward-free RLHF","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.20304","snapshot_observed_at":"2026-08-06T21:13:12.979442Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.979442Z"},"links":{"cited_paper":"/paper/2405.20304","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:46c32a8e14b09c998511ac4203ae290330495878263d9f266791a9ff6daf4c48","observation_id":"4ca0d123-1808-4516-9361-94dde4e34aa5","resolution":{"observed_at":"2026-08-06T21:13:12.979442Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-06T21:13:12.991359Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:12.991359Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:ff39f7f6710c32198e9e17e4f976c22e5201636d49981c5f517f159ad7d0f9a0","observation_id":"9a3fc0a6-59c2-4f22-9ec6-c354e781314b","resolution":{"observed_at":"2026-08-06T21:13:12.991359Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.19595","last_updated":"2025-08-17T21:21:08Z","snapshot_observed_at":"2026-08-07T16:38:33.934628Z","submitted_at":"2025-03-25T12:21:26Z","title":"Optimizing Language Models for Inference Time Objectives using Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.19595","snapshot_observed_at":"2026-08-06T21:13:13.002002Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:13.002002Z"},"links":{"cited_paper":"/paper/2503.19595","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:36f29450899420fa29428db583bd4d03dc33fb2436f4caf8a99fa372cdaa3327","observation_id":"a5a03000-ddb6-4db8-a58a-639361ea3c9c","resolution":{"observed_at":"2026-08-06T21:13:13.002002Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.12253","last_updated":"2024-12-10T18:19:29Z","snapshot_observed_at":"2026-08-05T22:35:55.296192Z","submitted_at":"2024-04-18T15:21:34Z","title":"Toward Self-Improvement of LLMs via Imagination, Searching, and Criticizing","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.12253","snapshot_observed_at":"2026-08-06T21:13:13.009844Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:13.009844Z"},"links":{"cited_paper":"/paper/2404.12253","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:f07d6daab5b54114fc2eba55cde3223b5d45b424e29584a4466deea0e7cd6704","observation_id":"775a8845-ebc0-4eff-818c-9e53b93e8c2e","resolution":{"observed_at":"2026-08-06T21:13:13.009844Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.13373","last_updated":"2024-09-20T10:20:46Z","snapshot_observed_at":"2026-07-06T19:18:41.002837Z","submitted_at":"2024-09-20T10:20:46Z","title":"LLMs Still Can't Plan; Can LRMs? A Preliminary Evaluation of OpenAI's o1 on PlanBench","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.13373","snapshot_observed_at":"2026-08-06T21:13:13.017038Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:13.017038Z"},"links":{"cited_paper":"/paper/2409.13373","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:0757abce2d8d6b72e3404da8b73fb29a778e3e89646e0f0283309510055bf39b","observation_id":"0eea2786-3587-46b6-8b5f-be4c58072569","resolution":{"observed_at":"2026-08-06T21:13:13.017038Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.04692","last_updated":"2024-06-07T07:04:10Z","snapshot_observed_at":"2026-08-07T21:51:22.136369Z","submitted_at":"2024-06-07T07:04:10Z","title":"Mixture-of-Agents Enhances Large Language Model Capabilities","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.04692","snapshot_observed_at":"2026-08-06T21:13:13.023536Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:13.023536Z"},"links":{"cited_paper":"/paper/2406.04692","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:fc1b215dcecfe512b0df9c7dc3ff7a6571deb0c556f19914b538f024435ae56e","observation_id":"2e217cc9-f0fb-42d4-a558-6d7fcf82e397","resolution":{"observed_at":"2026-08-06T21:13:13.023536Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.04091","last_updated":"2023-05-26T07:06:48Z","snapshot_observed_at":"2026-07-06T15:24:07.662207Z","submitted_at":"2023-05-06T16:34:37Z","title":"Plan-and-Solve Prompting: Improving Zero-Shot Chain-of-Thought Reasoning by Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.04091","snapshot_observed_at":"2026-08-06T21:13:13.029610Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:13.029610Z"},"links":{"cited_paper":"/paper/2305.04091","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:ddeef0ad28e419d8ea853997be6e3cf925ea68b4c5683afa48fe7d42b869013b","observation_id":"5c404f8c-fe45-48a9-8648-53428f96f589","resolution":{"observed_at":"2026-08-06T21:13:13.029610Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2203.11171","last_updated":"2023-03-07T17:57:37Z","snapshot_observed_at":"2026-07-06T12:50:22.773056Z","submitted_at":"2022-03-21T17:48:52Z","title":"Self-Consistency Improves Chain of Thought Reasoning in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2203.11171","snapshot_observed_at":"2026-08-06T21:13:13.038644Z","title":"Chi, Sharan Narang, Aakanksha Chowdhery, and Denny Zhou","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:13.038644Z"},"links":{"cited_paper":"/paper/2203.11171","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:55354bdcedc589c0af50a2ad93b14d67c311b3d5a9de2aacdc2211eea1a27264","observation_id":"e72b9d8a-2607-441c-8556-fd0ff6952531","resolution":{"observed_at":"2026-08-06T21:13:13.038644Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.18585","last_updated":"2025-02-18T16:51:53Z","snapshot_observed_at":"2026-07-06T20:28:33.378335Z","submitted_at":"2025-01-30T18:58:18Z","title":"Thoughts Are All Over the Place: On the Underthinking of o1-Like LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.18585","snapshot_observed_at":"2026-08-06T21:13:13.046279Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:13.046279Z"},"links":{"cited_paper":"/paper/2501.18585","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:b71dadf504254c5c5b245c476013982ef53b19aa43f643aefd3e673401cdfe8d","observation_id":"0fd5dae3-0522-43a6-ac4b-fe8238e94796","resolution":{"observed_at":"2026-08-06T21:13:13.046279Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:13:13.612297Z","title":"Chi, Quoc V Le, and Denny Zhou","venue":null,"work_id":"a310a092-beda-492a-8898-93e6b00ff27c","year":2022},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:13.051777Z"},"links":{"citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:1768e85d7cf08b6d408da27d03a803dde5f82827bbc4b3cff9c20bb4a91ec8c5","observation_id":"f90ae4c5-9a03-4e7a-9329-81609b711af0","resolution":{"observed_at":"2026-08-06T21:13:13.645210Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:13:13.528581Z","title":null,"venue":null,"work_id":"1a37e705-7ca7-4565-ab28-c5afc4806a85","year":2023},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:13.060486Z"},"links":{"citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:68155ac9e9e916003742605b3d9e7259527e0fa9ee331644c2a101a62a6576d9","observation_id":"bfeba1c9-4df4-465b-9e94-68c2ebfa3adc","resolution":{"observed_at":"2026-08-06T21:13:13.561826Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.03373","last_updated":"2025-02-05T17:13:32Z","snapshot_observed_at":"2026-07-06T20:31:41.231839Z","submitted_at":"2025-02-05T17:13:32Z","title":"Demystifying Long Chain-of-Thought Reasoning in LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.03373","snapshot_observed_at":"2026-08-06T21:13:13.066825Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:13.066825Z"},"links":{"cited_paper":"/paper/2502.03373","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:722b00ff07087877cc00e2f1c28ef0e2f2147c7e4d760ce3a53247389d0a8b4a","observation_id":"6f9f34ad-ba7c-4b13-aa4e-b1ca80116534","resolution":{"observed_at":"2026-08-06T21:13:13.066825Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.13837","last_updated":"2025-11-24T06:11:04Z","snapshot_observed_at":"2026-07-06T21:11:34.701779Z","submitted_at":"2025-04-18T17:59:56Z","title":"Does Reinforcement Learning Really Incentivize Reasoning Capacity in LLMs Beyond the Base Model?","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.13837","snapshot_observed_at":"2026-08-06T21:13:13.072280Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:13.072280Z"},"links":{"cited_paper":"/paper/2504.13837","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:9cfa5e814f4ede442eda0dc2122026587c530d176b25b6a3bdde91df2bd2afd8","observation_id":"ded393c9-1296-4f9d-90eb-923d45517bf6","resolution":{"observed_at":"2026-08-06T21:13:13.072280Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.08603","last_updated":"2025-01-31T05:28:15Z","snapshot_observed_at":"2026-07-06T20:21:13.914417Z","submitted_at":"2025-01-15T06:00:50Z","title":"Monte Carlo Tree Search for Comprehensive Exploration in LLM-Based Automatic Heuristic Design","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.08603","snapshot_observed_at":"2026-08-06T21:13:13.077881Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:13.077881Z"},"links":{"cited_paper":"/paper/2501.08603","citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:51c33aa03e4b33603c450d917a9f8d6168481d64b4d6d27d34175f7f6a24f071","observation_id":"b7acd3fa-06b0-4497-87a5-1012cee4c031","resolution":{"observed_at":"2026-08-06T21:13:13.077881Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:13:13.092764Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:13.092764Z"},"links":{"citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:ea613447157b427c05fbd65cfa33e91ed371d8fba4237fb1f745c3bf2338e388","observation_id":"8fb864ed-aafa-4f21-baaf-9e5bd23f14f1","resolution":{"observed_at":"2026-08-06T21:13:13.092764Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:13:13.121454Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-06T21:13:13.121454Z"},"links":{"citing_paper":"/paper/2507.00711"},"observation_digest":"sha256:88d1aef9eba7f01a30c904160438852c081bef68af0c0bf5f8532eb3df72f36f","observation_id":"84031adb-6967-4cab-9ef1-9efdd79a6d42","resolution":{"observed_at":"2026-08-06T21:13:13.121454Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2507.00711","last_updated":"2025-07-01T12:14:22Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-07T07:21:03.621965Z","submitted_at":"2025-07-01T12:14:22Z","title":"Large Reasoning Models are not thinking straight: on the unreliability of thinking trajectories"},"reference_resolution":{"displayed":39,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":37,"verified_exact":0,"verified_fuzzy":2},"total_outbound_references":39},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 39 of 39 outbound references and 2 inbound Pith citation observations for arXiv:2507.00711."}