{"as_of":"2026-08-22T00:49:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2ff83fbbde9569a0d32536ab6c7088baeb61cb08d5d449cf4fd21818a051cf15","coverage":[{"denominator":39,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":39,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-30T21:44:25.238171Z","state":"measured"},{"denominator":39,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":39,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-21T06:32:19.484+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.23458/citation-record","integrity":"/paper/2607.23458/integrity","json":"/paper/2607.23458/citation-record.json","paper":"/paper/2607.23458"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.087781Z","title":"2026 , url =","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.087781Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:b196d14bf1a49a5860c25105df80957cd4d953e3f73ceeb876be9bec1e81de01","observation_id":"5c90a8f7-8983-47fd-af4a-21b69c0a0911","resolution":{"observed_at":"2026-07-30T21:44:25.087781Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.092939Z","title":"Premise-Augmented Reasoning Chains Improve Error Identification in Math Reasoning with","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.092939Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:96b0f00fed9da5f6a1728a3cc21493457de40a3c9f3ca54a049413fa53446d54","observation_id":"45c354a4-4f56-402d-bc6a-db3d07063996","resolution":{"observed_at":"2026-07-30T21:44:25.092939Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.103776Z","title":"Graph of Verification: Structured Verification of","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.103776Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:244aa2095c4cb38e0068a2c0073f2a3dddd187cf17b70b3ecd87450f3e326763","observation_id":"b1e97867-3c4e-42c9-a4af-ecf71bf40d53","resolution":{"observed_at":"2026-07-30T21:44:25.103776Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.107284Z","title":"2025 , url =","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.107284Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:a6c51de4956e0abe4050f86a4a21616a40b482d573dfa584857e77359f141370","observation_id":"63cd86c1-567f-4118-bb42-e50157cb2def","resolution":{"observed_at":"2026-07-30T21:44:25.107284Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.110637Z","title":"2024 , url =","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.110637Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:9efde33bf282e8071ad1e7701efe17366b8419b4cfb7e7cbefec9c4316c2d9dc","observation_id":"54445dbf-8bab-4579-87d6-4705c8e2e69b","resolution":{"observed_at":"2026-07-30T21:44:25.110637Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.113805Z","title":"Proceedings of the 62nd Annual Meeting of the Association for Computational Linguistics (ACL) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.113805Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:170097061e4d6e4a12621c95c834ed05c3fcf626f752406ec6871adbadda4293","observation_id":"663a7cfe-f43f-4d7f-adbe-0dcfee81707a","resolution":{"observed_at":"2026-07-30T21:44:25.113805Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.117046Z","title":"2024 , url =","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.117046Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:df17bcf08ce0db2ed7ad36c01eebc9fa076b6c7a86bc60eedccc0986a8e7714e","observation_id":"b4f52f9d-0bc0-407d-9728-7da8179566a3","resolution":{"observed_at":"2026-07-30T21:44:25.117046Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.120520Z","title":"2024 , url =","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.120520Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:efc816cdcc3b1289232ecb16230732482b48e25a572d22aeaf83c5374e681cc8","observation_id":"51350e97-d834-4bcd-ac8f-e1ba7d675582","resolution":{"observed_at":"2026-07-30T21:44:25.120520Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.123447Z","title":"2026 , url =","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.123447Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:e8ba6c07bbfbffcf06beaeeeb20388014a794b0318d5ded8937bf7b843c76984","observation_id":"a77be292-0f70-4c61-8f16-db713e73029b","resolution":{"observed_at":"2026-07-30T21:44:25.123447Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.126712Z","title":"Advances in Neural Information Processing Systems (NeurIPS) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.126712Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:bdf28f836b66dbebf8bc649d012a1be2f03b5f5e69cf0242a7db3b141882ae7d","observation_id":"d37cb9d2-23d2-4ea8-a28b-f365ad0f5051","resolution":{"observed_at":"2026-07-30T21:44:25.126712Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.133088Z","title":"Transactions on Machine Learning Research , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.133088Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:77ad6ad7370e2d50aae7e89c15ad643916f365b9a50212dedf7512f30336c53d","observation_id":"d079c9ba-5d78-4cbe-af49-30f920f45aba","resolution":{"observed_at":"2026-07-30T21:44:25.133088Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.136303Z","title":"Computational Linguistics , volume =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.136303Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:5e6e6d34783d47b6ac631e47459084347eaecb81fed42daf609c53bd65e14d87","observation_id":"ab858bdf-66ff-4807-8d8a-d6f9b8fa57bc","resolution":{"observed_at":"2026-07-30T21:44:25.136303Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.139165Z","title":"Human Brain Mapping , volume =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.139165Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:67109297ae5fe5a109ca83ab1a4be91240270c6a62d46a508b6df0f2200a401e","observation_id":"53932c48-2ec8-4f48-8b87-dcfdb8b2ec5d","resolution":{"observed_at":"2026-07-30T21:44:25.139165Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.141891Z","title":"Transactions of the Association for Computational Linguistics , volume =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.141891Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:02352adce07076b9c7232bc70f67032ce64b16c5c12a7aa325d24f78c7523cf3","observation_id":"6ff198eb-19d3-413e-8f75-89b0dba71691","resolution":{"observed_at":"2026-07-30T21:44:25.141891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.144953Z","title":"Advances in Neural Information Processing Systems (NeurIPS) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.144953Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:3a228d2783126398045bdc68cfb57c3cd9a6ea256dd6227d0a9d11f38faaa7eb","observation_id":"177f441b-a874-4651-b1a6-ec2343634603","resolution":{"observed_at":"2026-07-30T21:44:25.144953Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.168173Z","title":"ICLR 2026 Workshop , year =","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.168173Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:5f1194baf80fc514ee29edde8c0e89eb307076b1b7567e29aeb090cadd00705e","observation_id":"3cb644b2-8d2b-4cf7-85d7-3581afffe700","resolution":{"observed_at":"2026-07-30T21:44:25.168173Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.08679","last_updated":"2026-06-16T17:36:22Z","snapshot_observed_at":"2026-08-16T12:51:00.810519Z","submitted_at":"2025-03-11T17:56:30Z","title":"Chain-of-Thought Reasoning In The Wild Is Not Always Faithful","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.08679","snapshot_observed_at":"2026-07-30T21:44:25.171383Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.171383Z"},"links":{"cited_paper":"/paper/2503.08679","citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:4ba0d59abd073a6349ff6a2044cd26c2b2b38b24d58deb180f3dc36d3d1eb372","observation_id":"656cdbbf-656c-4531-827b-b2f40dc799bf","resolution":{"observed_at":"2026-07-30T21:44:25.171383Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.174410Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.174410Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:f225fa9142da30d8471139b86b3c4106f9038c119b0f0ca922b4df9f180656ac","observation_id":"fa43d79d-f3d7-4bd7-a3b4-eb45291fee1f","resolution":{"observed_at":"2026-07-30T21:44:25.174410Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14897","last_updated":"2024-06-21T13:39:14Z","snapshot_observed_at":"2026-08-16T14:16:10.089897Z","submitted_at":"2024-02-22T17:23:53Z","title":"Chain-of-Thought Unfaithfulness as Disguised Accuracy","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14897","snapshot_observed_at":"2026-07-30T21:44:25.177613Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.177613Z"},"links":{"cited_paper":"/paper/2402.14897","citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:0c491835ec9edc2558496a63498b8c91c98ecf6fea787b00e775b343f8817e91","observation_id":"d9f4320e-7cb0-493e-a4e6-78958db79175","resolution":{"observed_at":"2026-07-30T21:44:25.177613Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.05410","last_updated":"2025-05-08T16:51:43Z","snapshot_observed_at":"2026-08-15T05:29:13.739206Z","submitted_at":"2025-05-08T16:51:43Z","title":"Reasoning Models Don't Always Say What They Think","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.05410","snapshot_observed_at":"2026-07-30T21:44:25.180976Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.180976Z"},"links":{"cited_paper":"/paper/2505.05410","citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:b970dd351934e96cf8510746423ca4d9a2a22a175ec4938d2cc084c854c240d1","observation_id":"f6c2a3ca-6a69-4867-be91-3ad2dfbec8bd","resolution":{"observed_at":"2026-07-30T21:44:25.180976Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2603.01437","last_updated":"2026-07-23T06:30:11Z","snapshot_observed_at":"2026-08-20T05:10:20.587943Z","submitted_at":"2026-03-02T04:33:55Z","title":"Post-Hoc Reasoning in Chain of Thought: Decoding and Steering Pre-Committed Answers","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2603.01437","snapshot_observed_at":"2026-07-30T21:44:25.183808Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.183808Z"},"links":{"cited_paper":"/paper/2603.01437","citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:9b62cc199edecfc9c5139bbfd5fea6aabd6b9db8b13ea71f8233434fbe2ee604","observation_id":"4ea66b1d-bee3-425d-beb3-8da6555f0ee9","resolution":{"observed_at":"2026-07-30T21:44:25.183808Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.186637Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.186637Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:bc0333a847b79237c8e468cba9f1aed05275cd8ba4c795069392f15778c5361e","observation_id":"167ab9f8-bb50-4e86-a8d3-8ab5447bec22","resolution":{"observed_at":"2026-07-30T21:44:25.186637Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.189486Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.189486Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:0f8ddc38e34ff5596ed2aea3df5a4d4f20ef35b8e322b32ef1f2dc88ae971795","observation_id":"4452d943-fddf-4feb-82f9-673f53d06f80","resolution":{"observed_at":"2026-07-30T21:44:25.189486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.192488Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.192488Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:830a7cc4f1b922df5165398e4d71448f88da452474b27a45f34b76ec0f918a09","observation_id":"55eb3452-66a2-4e9d-8d41-78268878e668","resolution":{"observed_at":"2026-07-30T21:44:25.192488Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.25052","last_updated":"2026-05-24T12:57:01Z","snapshot_observed_at":"2026-08-14T13:05:55.119852Z","submitted_at":"2026-05-24T12:57:01Z","title":"Faithfulness Metrics Don't Measure Faithfulness: A Meta-Evaluation with Ground Truth","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.25052","snapshot_observed_at":"2026-07-30T21:44:25.195610Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.195610Z"},"links":{"cited_paper":"/paper/2605.25052","citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:d4f72b609d8891a6bac5bfd730f9c15f8a8d4c9e884804b9985b10e4edae82aa","observation_id":"d73f7de8-d029-4222-863a-2631efb2447d","resolution":{"observed_at":"2026-07-30T21:44:25.195610Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.198641Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.198641Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:f2f33c0a10fa35f4f2ac540d7d44c20586c36a953213a1b8ef686eb93076bbf1","observation_id":"8e6d1c9e-6a5b-4728-92d8-a1e100e05ec9","resolution":{"observed_at":"2026-07-30T21:44:25.198641Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.17408","last_updated":"2026-04-19T21:58:15Z","snapshot_observed_at":"2026-08-10T23:27:19.301813Z","submitted_at":"2025-11-21T17:08:48Z","title":"The Impact of Off-Policy Training Data on Probe Generalisation","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.17408","snapshot_observed_at":"2026-07-30T21:44:25.201514Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.201514Z"},"links":{"cited_paper":"/paper/2511.17408","citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:5d141d0d2b7f78d43392a6514dace3cdc05332e422e1a395754fee976a9d6d01","observation_id":"f6d7ce6b-f995-4274-8788-22a11b04f80c","resolution":{"observed_at":"2026-07-30T21:44:25.201514Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.13702","last_updated":"2023-07-17T01:08:39Z","snapshot_observed_at":"2026-08-13T18:02:27.155379Z","submitted_at":"2023-07-17T01:08:39Z","title":"Measuring Faithfulness in Chain-of-Thought Reasoning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.13702","snapshot_observed_at":"2026-07-30T21:44:25.204356Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.204356Z"},"links":{"cited_paper":"/paper/2307.13702","citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:8cb1585814fa28a8a398904987efd68e17005cbaf02d7767ae25ec0ebf571be0","observation_id":"9d428722-5969-4338-9261-21b267d0f34e","resolution":{"observed_at":"2026-07-30T21:44:25.204356Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.207228Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.207228Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:f2e2dd8d5964b4d826396d105e11fa8829febd9ac530b019ffa76cc2d1178489","observation_id":"8530bae0-d2a5-49d9-8c83-32598ef46996","resolution":{"observed_at":"2026-07-30T21:44:25.207228Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.209862Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.209862Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:b647ea773fd32daf894523cbbfafb78af031211bb20be73d35abc97956723ada","observation_id":"a936a45c-9cb7-4647-8874-c3c9f21a5930","resolution":{"observed_at":"2026-07-30T21:44:25.209862Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.212786Z","title":"Nichols and Andrew P","venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.212786Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:4ebaa393aaf64c79c5f797f8710ce3f4c901e619f67ddc03eba3e9b421d408a7","observation_id":"e28dad47-8a30-464a-a7c9-9633bf262f26","resolution":{"observed_at":"2026-07-30T21:44:25.212786Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.215788Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.215788Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:790f673c13f32a808e3870ee19ed4cd8e11f25f8e1a3051c16ba72ed110b54a8","observation_id":"92653279-6142-412e-abe1-66b16d248ce2","resolution":{"observed_at":"2026-07-30T21:44:25.215788Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.219552Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.219552Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:ca7c983d55250fe5faecf4e37381504c756a7e017dc721fd9d7bfeb4cfc98cd8","observation_id":"da77993a-be31-43a7-974c-ef48f790ad7f","resolution":{"observed_at":"2026-07-30T21:44:25.219552Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.25603","last_updated":"2026-05-25T08:54:55Z","snapshot_observed_at":"2026-08-10T09:24:29.678831Z","submitted_at":"2026-05-25T08:54:55Z","title":"Detecting Unfaithful Chain-of-Thought via Circuit-Guided Internal-External Discrepancy","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.25603","snapshot_observed_at":"2026-07-30T21:44:25.222539Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.222539Z"},"links":{"cited_paper":"/paper/2605.25603","citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:f51f8953b2f0242080877e2b4779256949e4915f5b3f11699ec2ee294ee72443","observation_id":"a041051a-837f-4e86-92af-67c7ceb8a7b3","resolution":{"observed_at":"2026-07-30T21:44:25.222539Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-30T21:44:25.225351Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.225351Z"},"links":{"citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:2bfa74df62b2dba0a018d914e11a540824ca2bc0274e0d5148f2ed9384206908","observation_id":"f3ba1e7c-6578-46e8-918b-b3118d3253b0","resolution":{"observed_at":"2026-07-30T21:44:25.225351Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.12404","last_updated":"2025-05-04T23:07:56Z","snapshot_observed_at":"2026-08-18T06:51:37.465225Z","submitted_at":"2024-07-17T08:32:03Z","title":"Analyzing the Generalization and Reliability of Steering Vectors","version":8},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.12404","snapshot_observed_at":"2026-07-30T21:44:25.228410Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.228410Z"},"links":{"cited_paper":"/paper/2407.12404","citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:3be2340f31e9dce1f82d0cee19a12b2022c2194f37fd3a60f973f96ffb67b267","observation_id":"14af5936-ed9a-44d3-a843-2273bb7bd155","resolution":{"observed_at":"2026-07-30T21:44:25.228410Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.04388","last_updated":"2023-12-09T21:25:02Z","snapshot_observed_at":"2026-08-12T21:03:44.597326Z","submitted_at":"2023-05-07T22:44:25Z","title":"Language Models Don't Always Say What They Think: Unfaithful Explanations in Chain-of-Thought Prompting","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.04388","snapshot_observed_at":"2026-07-30T21:44:25.231592Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.231592Z"},"links":{"cited_paper":"/paper/2305.04388","citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:eefd274773b6ca5cc1a57421c44c9dc9ed57a8362e702afb38627859d28e9160","observation_id":"bceb4ed0-d5e8-471d-ad53-464515de4a19","resolution":{"observed_at":"2026-07-30T21:44:25.231592Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2512.23032","last_updated":"2026-05-08T14:16:13Z","snapshot_observed_at":"2026-08-19T13:20:05.756139Z","submitted_at":"2025-12-28T18:18:02Z","title":"Is Chain-of-Thought Really Not Explainability? Chain-of-Thought Can Be Faithful without Hint Verbalization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2512.23032","snapshot_observed_at":"2026-07-30T21:44:25.235255Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.235255Z"},"links":{"cited_paper":"/paper/2512.23032","citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:c6bd9a33d3b9ed95b818e81c9fb4ae6f87e2574318876f4a8219fbbaa3ad5386","observation_id":"bc03e437-6f18-4ee7-8339-75f662a4b0c2","resolution":{"observed_at":"2026-07-30T21:44:25.235255Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.06559","last_updated":"2025-05-26T14:03:32Z","snapshot_observed_at":"2026-08-14T07:02:14.518172Z","submitted_at":"2024-12-09T15:11:40Z","title":"ProcessBench: Identifying Process Errors in Mathematical Reasoning","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.06559","snapshot_observed_at":"2026-07-30T21:44:25.238171Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-07-30T21:44:25.238171Z"},"links":{"cited_paper":"/paper/2412.06559","citing_paper":"/paper/2607.23458"},"observation_digest":"sha256:e79638485b8690eb0efb2dfeee74fff949378897899e7fbd2712047fc19af55a","observation_id":"ba25b37c-5141-4c06-82f7-a91bcdc0d0a9","resolution":{"observed_at":"2026-07-30T21:44:25.238171Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2607.23458","last_updated":"2026-07-26T04:43:53Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-17T18:54:37.760375Z","submitted_at":"2026-07-26T04:43:53Z","title":"Two Regimes of Chain-of-Thought Unfaithfulness: Behavioral Detection Fails Where Models Are Wrong"},"reference_resolution":{"displayed":39,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":39,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":39},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"thesis":"As of 22 August 2026, this Paper Citation Record lists 39 of 39 outbound references and 0 inbound Pith citation observations for arXiv:2607.23458."}