{"as_of":"2026-08-11T07:27:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:c66b7bf0300ecb643f67611c0a55f8cf1b78194f725ca473db1ca801910bd5de","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":17,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":17,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-11T06:34:44.6726+00:00","state":"measured"},{"denominator":17,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":17,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T16:41:22.141140Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":2,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":"2310.18512","doi":"10.48550/arxiv.2310.18512","metadata_source":"arxiv_reference","pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2310.18512 , year=","venue":"arXiv (Cornell University)","work_id":"eb34c557-1225-4ed5-bdd4-0781473b6cf1","year":2023},"citing_paper":{"arxiv_id":"2412.14093","last_updated":"2024-12-20T02:22:19Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-18T17:41:24Z","title":"Alignment faking in large language models","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-12T22:50:11.846863Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2412.14093"},"observation_digest":"sha256:97f891e3e707acc5b58fda25aa26a705729535bd66b6b70e046c3d74969e42ed","observation_id":"06a3266a-2cee-4b34-8223-3dcf28092312","resolution":{"observed_at":"2026-05-12T22:50:11.946185Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-10T16:41:22.141140Z","title":"different","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2501.13011","last_updated":"2025-04-10T16:25:31Z","snapshot_observed_at":"2026-08-10T16:29:47.737564Z","submitted_at":"2025-01-22T16:53:08Z","title":"MONA: Myopic Optimization with Non-myopic Approval Can Mitigate Multi-step Reward Hacking","version":2},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-10T16:41:22.141140Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2501.13011"},"observation_digest":"sha256:1d6653f858c82d5e3ff2a6f4af7b4850416f94c7356b112b3dece43217eb12ea","observation_id":"42aceb86-a0bb-481f-b238-dfaa9b454102","resolution":{"observed_at":"2026-08-10T16:41:22.141140Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-06T22:03:54.383501Z","title":"and Greenblatt, R","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.22777","last_updated":"2025-07-13T15:36:35Z","snapshot_observed_at":"2026-08-10T05:25:44.170625Z","submitted_at":"2025-06-28T06:37:10Z","title":"Teaching Models to Verbalize Reward Hacking in Chain-of-Thought Reasoning","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T22:03:54.383501Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2506.22777"},"observation_digest":"sha256:ccdb291e57845b08c86e9bc86249233475e37979181f3741d8b05088b10af527","observation_id":"67be5f1f-89f9-49de-b9de-97969602ee1a","resolution":{"observed_at":"2026-08-06T22:03:54.383501Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-06T19:34:55.843136Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.05246","last_updated":"2025-07-07T17:54:52Z","snapshot_observed_at":"2026-08-09T22:19:12.076180Z","submitted_at":"2025-07-07T17:54:52Z","title":"When Chain of Thought is Necessary, Language Models Struggle to Evade Monitors","version":1},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-06T19:34:55.843136Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2507.05246"},"observation_digest":"sha256:d9a8219f6ee84328077cba4bc162008c8c8a0e56c7a0601699455ff75c1443c0","observation_id":"31c74176-d014-4c06-b0de-dc539a7e1980","resolution":{"observed_at":"2026-08-06T19:34:55.843136Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":"2310.18512","doi":"10.48550/arxiv.2310.18512","metadata_source":"arxiv_reference","pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2310.18512 , year=","venue":"arXiv (Cornell University)","work_id":"eb34c557-1225-4ed5-bdd4-0781473b6cf1","year":2023},"citing_paper":{"arxiv_id":"2507.11473","last_updated":"2025-12-07T02:14:12Z","snapshot_observed_at":"2026-08-10T09:38:50.886870Z","submitted_at":"2025-07-15T16:43:41Z","title":"Chain of Thought Monitorability: A New and Fragile Opportunity for AI Safety","version":2},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-05-20T14:19:44.695462Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2507.11473"},"observation_digest":"sha256:ace4e872b9e10c3302862d86afb9003d5a624d6ac217de8d1c3d4986a870ce46","observation_id":"1bf448f2-6dcd-4830-905a-704b7d8ece0f","resolution":{"observed_at":"2026-05-20T14:19:44.761298Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-02T23:26:30.628642Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.13904","last_updated":"2026-07-23T02:27:49Z","snapshot_observed_at":"2026-08-10T06:53:01.769687Z","submitted_at":"2026-02-14T21:53:47Z","title":"Diagnosing Pathological Chain-of-Thought in Reasoning Models","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-02T23:26:30.628642Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2602.13904"},"observation_digest":"sha256:06e0f52f5fa1b8620884f4925469d62d9007d03ebde933f9e510504e0cecc83b","observation_id":"1f21787f-b906-4a30-bddd-cd0a03fb9630","resolution":{"observed_at":"2026-08-02T23:26:30.628642Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":"2310.18512","doi":"10.48550/arxiv.2310.18512","metadata_source":"arxiv_reference","pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2310.18512 , year=","venue":"arXiv (Cornell University)","work_id":"eb34c557-1225-4ed5-bdd4-0781473b6cf1","year":2023},"citing_paper":{"arxiv_id":"2604.16242","last_updated":"2026-04-17T17:01:24Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-17T17:01:24Z","title":"Detecting and Suppressing Reward Hacking with Gradient Fingerprints","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T08:18:37.665350Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2604.16242"},"observation_digest":"sha256:1733345c6ca4e63e2f5bf649e612445efcb4c838956140f41fd20c234a88a4aa","observation_id":"3d71b3f4-33e8-4484-be57-83039ffb90f6","resolution":{"observed_at":"2026-05-10T08:22:37.572076Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":"2310.18512","doi":"10.48550/arxiv.2310.18512","metadata_source":"arxiv_reference","pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2310.18512 , year=","venue":"arXiv (Cornell University)","work_id":"eb34c557-1225-4ed5-bdd4-0781473b6cf1","year":2023},"citing_paper":{"arxiv_id":"2605.11746","last_updated":"2026-05-12T08:24:47Z","snapshot_observed_at":"2026-08-03T19:32:09.592136Z","submitted_at":"2026-05-12T08:24:47Z","title":"When Reasoning Traces Become Performative: Step-Level Evidence that Chain-of-Thought Is an Imperfect Oversight Channel","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-05-13T06:30:12.558660Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2605.11746"},"observation_digest":"sha256:d9802c9a2cb73fd9e8cb7ba751e57d7e3b46531b0201f2e50fc71b044742ebff","observation_id":"bc293fa8-4e7f-4c12-b12d-7002267f0c87","resolution":{"observed_at":"2026-05-13T06:32:24.393171Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":"2310.18512","doi":"10.48550/arxiv.2310.18512","metadata_source":"arxiv_reference","pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2310.18512 , year=","venue":"arXiv (Cornell University)","work_id":"eb34c557-1225-4ed5-bdd4-0781473b6cf1","year":2023},"citing_paper":{"arxiv_id":"2605.24396","last_updated":"2026-08-10T01:19:49Z","snapshot_observed_at":"2026-08-11T07:22:11.826661Z","submitted_at":"2026-05-23T04:42:45Z","title":"Understanding and Mitigating Premature Confidence for Better LLM Reasoning","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-06-30T14:03:25.913615Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2605.24396"},"observation_digest":"sha256:068b00712c76048035a983f6bb8d4286347ae82bb0722e22753008cf302e7a44","observation_id":"03f5e166-a971-4774-ad45-9f8e85f5c4a0","resolution":{"observed_at":"2026-06-30T14:04:44.385114Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":"2310.18512","doi":"10.48550/arxiv.2310.18512","metadata_source":"arxiv_reference","pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2310.18512 , year=","venue":"arXiv (Cornell University)","work_id":"eb34c557-1225-4ed5-bdd4-0781473b6cf1","year":2023},"citing_paper":{"arxiv_id":"2605.26537","last_updated":"2026-05-26T04:38:56Z","snapshot_observed_at":"2026-07-06T23:36:20.072253Z","submitted_at":"2026-05-26T04:38:56Z","title":"Conceptual Steganography","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-06-29T18:46:39.761975Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2605.26537"},"observation_digest":"sha256:837b2449a56449fb2570baf5bb4745ecceffa853c4bb8772d52714c835c21e73","observation_id":"f752b928-caa1-4e30-a480-7f9f7a324065","resolution":{"observed_at":"2026-06-29T18:53:51.621461Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":"2310.18512","doi":"10.48550/arxiv.2310.18512","metadata_source":"arxiv_reference","pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2310.18512 , year=","venue":"arXiv (Cornell University)","work_id":"eb34c557-1225-4ed5-bdd4-0781473b6cf1","year":2023},"citing_paper":{"arxiv_id":"2606.09931","last_updated":"2026-06-07T16:36:58Z","snapshot_observed_at":"2026-08-08T10:57:53.795341Z","submitted_at":"2026-06-07T16:36:58Z","title":"A Note on the Strategic Confinement Problem","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-06-27T17:30:12.165074Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2606.09931"},"observation_digest":"sha256:58a33acf2db0dcd3283fc5f0bd93dde98eeabafb8e817abeacebd51249873c71","observation_id":"4ad341bd-0be5-4de4-80fe-cc441e544fac","resolution":{"observed_at":"2026-06-27T17:31:06.842944Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":"2310.18512","doi":"10.48550/arxiv.2310.18512","metadata_source":"arxiv_reference","pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2310.18512 , year=","venue":"arXiv (Cornell University)","work_id":"eb34c557-1225-4ed5-bdd4-0781473b6cf1","year":2023},"citing_paper":{"arxiv_id":"2606.11063","last_updated":"2026-06-09T16:24:16Z","snapshot_observed_at":"2026-08-09T22:11:04.022394Z","submitted_at":"2026-06-09T16:24:16Z","title":"CIAware-Bench: Benchmarking Control Intervention Awareness Across Frontier LLMs","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-06-27T13:23:28.061924Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2606.11063"},"observation_digest":"sha256:737c70124cf42816f435a099a137ecde603b1595c55b96d0688aa6c8d0f339b1","observation_id":"122c54e7-21d8-478a-b498-880e09ae5697","resolution":{"observed_at":"2026-07-03T05:07:39.423665Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":"2310.18512","doi":"10.48550/arxiv.2310.18512","metadata_source":"arxiv_reference","pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2310.18512 , year=","venue":"arXiv (Cornell University)","work_id":"eb34c557-1225-4ed5-bdd4-0781473b6cf1","year":2023},"citing_paper":{"arxiv_id":"2606.19603","last_updated":"2026-06-17T21:15:19Z","snapshot_observed_at":"2026-08-02T22:25:31.054596Z","submitted_at":"2026-06-17T21:15:19Z","title":"Comparing Linear Probes with Mahalanobis Cosine Similarity","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-06-26T20:40:03.169825Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2606.19603"},"observation_digest":"sha256:c100db60829dcd6cba4609c375592cb4d65805c0f85f0d4bc88585c17f189f25","observation_id":"e8a0d71d-82ff-46f5-ac74-7566df5fb004","resolution":{"observed_at":"2026-07-04T01:09:18.923298Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":"2310.18512","doi":"10.48550/arxiv.2310.18512","metadata_source":"arxiv_reference","pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2310.18512 , year=","venue":"arXiv (Cornell University)","work_id":"eb34c557-1225-4ed5-bdd4-0781473b6cf1","year":2023},"citing_paper":{"arxiv_id":"2606.28615","last_updated":"2026-06-26T21:14:56Z","snapshot_observed_at":"2026-08-04T19:40:32.344244Z","submitted_at":"2026-06-26T21:14:56Z","title":"What LLMs explain is not what they believe: Evaluating explanation sufficiency under models' own input beliefs","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-06-30T00:38:21.949283Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2606.28615"},"observation_digest":"sha256:41ea86e4cd9df6d5ee0eab1b0bede61c3e554c53f41b1cc6f880207b566b2e02","observation_id":"0b3af246-dbce-44ae-b803-3e0a0600999d","resolution":{"observed_at":"2026-07-01T16:25:50.102094Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-01T23:29:50.791477Z","title":"Preventing language models from hiding their reasoning, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.15434","last_updated":"2026-08-09T21:27:54Z","snapshot_observed_at":"2026-08-11T07:22:21.526508Z","submitted_at":"2026-07-16T20:07:47Z","title":"Coercion and Deception in AI-to-AI Management: An Agentic Benchmark of Unprompted Escalation","version":4},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-01T23:29:50.791477Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2607.15434"},"observation_digest":"sha256:f55b5367edfccc6cadf081f48425b4d7d94e8968539fe1f33be445cf342058f5","observation_id":"3fd1c6e4-9836-428f-8115-7be3bbfda935","resolution":{"observed_at":"2026-08-01T23:29:50.791477Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-01T10:02:59.970123Z","title":"arXiv preprint arXiv:2310.18512 , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.20379","last_updated":"2026-07-22T17:10:23Z","snapshot_observed_at":"2026-08-07T23:15:05.523185Z","submitted_at":"2026-07-22T17:10:23Z","title":"Train the Model, Not the Reader: Decodability Supervision for Verifiable Activation Explanations","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-01T10:02:59.970123Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2607.20379"},"observation_digest":"sha256:4ae1abe9d76e8c84493593c90159aed613540f043cd7a8f7b9936413b6484a76","observation_id":"9df42a77-34b4-467c-8d69-28d57c845de5","resolution":{"observed_at":"2026-08-01T10:02:59.970123Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.18512","snapshot_observed_at":"2026-08-01T04:14:45.406064Z","title":"Preventing language models from hiding their reasoning, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.22925","last_updated":"2026-07-24T21:32:48Z","snapshot_observed_at":"2026-08-08T05:49:21.548828Z","submitted_at":"2026-07-24T21:32:48Z","title":"Not All LLM Reasoning is Visible in the Chain-of-Thought","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-01T04:14:45.406064Z"},"links":{"cited_paper":"/paper/2310.18512","citing_paper":"/paper/2607.22925"},"observation_digest":"sha256:7fb2967d33701ece7561457ffe7a99804e13d5f3f48d1975d131761f84c797b1","observation_id":"c7dbb06f-cfdd-481b-a3dc-6ebe56957df8","resolution":{"observed_at":"2026-08-01T04:14:45.406064Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2310.18512/citation-record","integrity":"/paper/2310.18512/integrity","json":"/paper/2310.18512/citation-record.json","paper":"/paper/2310.18512"},"outbound":[],"paper":{"arxiv_id":"2310.18512","last_updated":"2023-10-31T19:13:43Z","latest_version":2,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-06T07:03:02.030103Z","submitted_at":"2023-10-27T22:02:29Z","title":"Preventing Language Models From Hiding Their Reasoning"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"thesis":"As of 11 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 17 inbound Pith citation observations for arXiv:2310.18512."}