{"as_of":"2026-08-07T16:45:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2396432c41d0c28e42e363919f560c9e879b827bd270df634f61a345f20c90dc","coverage":[{"denominator":37,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":37,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T11:01:15.053772Z","state":"measured"},{"denominator":39,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":39,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-21T21:44:36.351517Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-21T21:45:40.587851Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"cited_work":{"arxiv_id":"2507.23221","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.23221","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Long Ouyang, Jeff Wu, Xu Jiang, Diogo Almeida, Carroll L","venue":null,"work_id":"15637379-4a59-4ade-bdcb-e970149f0155","year":null},"citing_paper":{"arxiv_id":"2509.22739","last_updated":"2026-05-15T15:16:23Z","snapshot_observed_at":"2026-07-06T22:30:55.313733Z","submitted_at":"2025-09-25T23:25:47Z","title":"Painless Activation Steering: An Automated, Lightweight Approach for Post-Training Large Language Models","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-21T21:44:36.351517Z"},"links":{"cited_paper":"/paper/2507.23221","citing_paper":"/paper/2509.22739"},"observation_digest":"sha256:965a73356ebcb621b5f4be054329cb0f570e626162eb429a3ac0858d7f75006d","observation_id":"8d0a9c56-e1aa-4ea9-9116-4297c04a8041","resolution":{"observed_at":"2026-05-21T21:45:40.590400Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"cited_work":{"arxiv_id":"2507.23221","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.23221","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Long Ouyang, Jeff Wu, Xu Jiang, Diogo Almeida, Carroll L","venue":null,"work_id":"15637379-4a59-4ade-bdcb-e970149f0155","year":null},"citing_paper":{"arxiv_id":"2605.17028","last_updated":"2026-05-16T14:57:15Z","snapshot_observed_at":"2026-07-06T23:28:06.971959Z","submitted_at":"2026-05-16T14:57:15Z","title":"PARALLAX: Separating Genuine Hallucination Detection from Benchmark Construction Artifacts","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-05-19T20:09:47.750043Z"},"links":{"cited_paper":"/paper/2507.23221","citing_paper":"/paper/2605.17028"},"observation_digest":"sha256:e44913beaf2ffa254b178d446a1b2bdf6b304ff91d803ca6b70d5dea2037d563","observation_id":"4097f8f0-7c55-4ee8-97d6-95ccbfba6437","resolution":{"observed_at":"2026-05-19T20:12:44.923519Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2507.23221/citation-record","integrity":"/paper/2507.23221/integrity","json":"/paper/2507.23221/citation-record.json","paper":"/paper/2507.23221"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"1610.01644","last_updated":"2018-11-22T23:40:00Z","snapshot_observed_at":"2026-07-06T05:13:30.860932Z","submitted_at":"2016-10-05T20:59:01Z","title":"Understanding intermediate layers using linear classifier probes","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1610.01644","snapshot_observed_at":"2026-08-06T11:01:14.932935Z","title":"and Bengio, Y","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.932935Z"},"links":{"cited_paper":"/paper/1610.01644","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:b44445aeb3a3ad8fa7dcce38e7fe5ab88387ce82d38c4a7399c02b644bd8e746","observation_id":"25dac620-b13b-4dcb-80fe-451c71a158d6","resolution":{"observed_at":"2026-08-06T11:01:14.932935Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T11:01:14.936933Z","title":"and Mitchell, T","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.936933Z"},"links":{"citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:2735b8d4b955655cb9dcdfc9ee3e3d328f9c64b5945289d6416ae0325ff4a587","observation_id":"b3d95ff1-40eb-4a7a-905e-475f22780c93","resolution":{"observed_at":"2026-08-06T11:01:14.936933Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T11:01:15.580612Z","title":"Discovering Latent Knowledge in Language Models Without Supervision","venue":null,"work_id":"fec569c0-301b-4973-8662-585393ba4d27","year":2023},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.940142Z"},"links":{"citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:fece29823f755da6700964f77aad565d0e572dfa3f44afa9c644f4aa76fa66c6","observation_id":"f213f5fe-6f27-4bc3-beb0-5c38c1498002","resolution":{"observed_at":"2026-08-06T11:01:15.584421Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T11:01:14.943561Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.943561Z"},"links":{"citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:9cd0e5714cd42a63e75af59936161eb4813ba4202c2c1f4a482acaf1a60281ff","observation_id":"7c804f6f-7361-4ebe-bee7-6c1e988777e3","resolution":{"observed_at":"2026-08-06T11:01:14.943561Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.07071","last_updated":"2024-10-03T17:26:48Z","snapshot_observed_at":"2026-08-04T03:32:56.098220Z","submitted_at":"2024-07-09T17:44:34Z","title":"Lookback Lens: Detecting and Mitigating Contextual Hallucinations in Large Language Models Using Only Attention Maps","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.07071","snapshot_observed_at":"2026-08-06T11:01:14.946851Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.946851Z"},"links":{"cited_paper":"/paper/2407.07071","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:82cda1b75c1ad9591a2e3786258cb65cec0724bc3774d1490df04dff07425f0c","observation_id":"aca8bfaa-8d87-483a-b4c1-a03cdfa3dcf1","resolution":{"observed_at":"2026-08-06T11:01:14.946851Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.08600","last_updated":"2023-10-04T13:17:38Z","snapshot_observed_at":"2026-07-06T16:19:05.495349Z","submitted_at":"2023-09-15T17:56:55Z","title":"Sparse Autoencoders Find Highly Interpretable Features in Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.08600","snapshot_observed_at":"2026-08-06T11:01:14.950678Z","title":"Sparse autoencoders find highly interpretable features in language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.950678Z"},"links":{"cited_paper":"/paper/2309.08600","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:2c5408b67cb0011d583bf893b47d1cb55c2255384bb7112aec638e2283881606","observation_id":"7d8f589a-a897-48c2-8201-98f30091193a","resolution":{"observed_at":"2026-08-06T11:01:14.950678Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2209.10652","last_updated":"2022-09-21T20:49:26Z","snapshot_observed_at":"2026-07-06T13:54:56.779166Z","submitted_at":"2022-09-21T20:49:26Z","title":"Toy Models of Superposition","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2209.10652","snapshot_observed_at":"2026-08-06T11:01:14.954567Z","title":"H., Lasenby, R., Drain, D., Chen, C., Grosse, R., McCandlish, S., Kaplan, J., Amodei, D., Wattenberg, M., and Olah, C","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.954567Z"},"links":{"cited_paper":"/paper/2209.10652","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:8ed57c8772f522b2b03c0f940291c056947140ae351905ebe8aa6154cefc98ee","observation_id":"bae3f20a-b573-41b3-837c-3d90c4fbcc55","resolution":{"observed_at":"2026-08-06T11:01:14.954567Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T11:01:15.570387Z","title":"J., Gurnee, W., and Tegmark, M","venue":null,"work_id":"3ca8bb17-1d98-4e98-ad9f-aec79b43675f","year":2024},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.957804Z"},"links":{"citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:9164746123c2d9c732f8422d51e33a1c37620468c90e3067c3148a7a7d855d6c","observation_id":"2ee00304-7644-481d-92c9-60ec99b60ba5","resolution":{"observed_at":"2026-08-06T11:01:15.573589Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T11:01:14.961171Z","title":"Detecting hallucinations in large language models using semantic entropy","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.961171Z"},"links":{"citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:9fb22918b13ca65e57eadb8cd5473a8116422d98d25d708805de42651c1d54ca","observation_id":"cf51700b-ea7a-45f0-8302-84d492486e92","resolution":{"observed_at":"2026-08-06T11:01:14.961171Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.14257","last_updated":"2025-02-08T12:50:42Z","snapshot_observed_at":"2026-07-06T19:53:52.079055Z","submitted_at":"2024-11-21T16:05:58Z","title":"Do I Know This Entity? Knowledge Awareness and Hallucinations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.14257","snapshot_observed_at":"2026-08-06T11:01:14.964076Z","title":"Do I Know This Entity? Knowledge Awareness and Hallucinations in Language Models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.964076Z"},"links":{"cited_paper":"/paper/2411.14257","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:0efe127af8593b9d53b0a1784ff1f95a7700209778492fea1b1015a9fcc74a9f","observation_id":"99c7fa8d-dbdd-45da-8611-e9028fd05045","resolution":{"observed_at":"2026-08-06T11:01:14.964076Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2101.00027","last_updated":"2020-12-31T19:00:10Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-12-31T19:00:10Z","title":"The Pile: An 800GB Dataset of Diverse Text for Language Modeling","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2101.00027","snapshot_observed_at":"2026-08-06T11:01:14.967476Z","title":"The pile: An 800gb dataset of diverse text for language modeling","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.967476Z"},"links":{"cited_paper":"/paper/2101.00027","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:5ff67cd77bc0064b654e359029bc071c2bd8375a041dd7dacf80b4da991459da","observation_id":"e8775890-8743-4f70-bf32-4465dae99fd6","resolution":{"observed_at":"2026-08-06T11:01:14.967476Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T11:01:15.559990Z","title":"spacy: Industrial-strength natural language processing in python","venue":null,"work_id":"a1a8b67e-ca00-41eb-b8cd-22c0ce7b7c2a","year":2020},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.970964Z"},"links":{"citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:5774e0ff18c48ab321aa05810263d623a6cfd756f63371fec74cb5fef041d58f","observation_id":"010c7cde-f373-46aa-947e-2f89ad41d465","resolution":{"observed_at":"2026-08-06T11:01:15.563539Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T11:01:14.973911Z","title":"A Survey on Hallucination in Large Language Models: Principles, Taxonomy, Challenges, and Open Questions","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.973911Z"},"links":{"citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:1a4f791e0af2457051a6f5d6323465e951d380d5acdd3b0734ac4f02327cd3d3","observation_id":"e448f1d3-c66a-4869-91b7-84c252dbd04e","resolution":{"observed_at":"2026-08-06T11:01:14.973911Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.05566","last_updated":"2024-01-17T20:26:01Z","snapshot_observed_at":"2026-07-06T17:14:03.152949Z","submitted_at":"2024-01-10T22:14:35Z","title":"Sleeper Agents: Training Deceptive LLMs that Persist Through Safety Training","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.05566","snapshot_observed_at":"2026-08-06T11:01:14.977196Z","title":"M., Maxwell, T., Cheng, N., et al","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.977196Z"},"links":{"cited_paper":"/paper/2401.05566","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:8c46e1af53e4b3e4833f66f77eb316fc0914bd321ee35c344822a16b21089362","observation_id":"f853e8fc-909e-401a-818b-d4bb84f6455e","resolution":{"observed_at":"2026-08-06T11:01:14.977196Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T11:01:14.980382Z","title":"Survey of Hallucination in Natural Language Generation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.980382Z"},"links":{"citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:2c5e3027ed384e545e1cc6fd8cb9958f1b9ae3e15ea472aaac9997f985274d65","observation_id":"316b27aa-2979-4ab1-bcc6-425a3813d766","resolution":{"observed_at":"2026-08-06T11:01:14.980382Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.15927","last_updated":"2024-06-22T19:46:06Z","snapshot_observed_at":"2026-07-30T04:59:24.269176Z","submitted_at":"2024-06-22T19:46:06Z","title":"Semantic Entropy Probes: Robust and Cheap Hallucination Detection in LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.15927","snapshot_observed_at":"2026-08-06T11:01:14.983343Z","title":"Semantic Entropy Probes: Robust and Cheap Hallucination Detection in LLMs , 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.983343Z"},"links":{"cited_paper":"/paper/2406.15927","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:2eda3ba55dea9221b9d9fcc5ec34e0423dbf6172ea8663820eab55a4412ce131","observation_id":"8aedb19d-a45e-4668-b3dc-72ada4946a1f","resolution":{"observed_at":"2026-08-06T11:01:14.983343Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1312.5663","last_updated":"2014-03-22T17:12:07Z","snapshot_observed_at":"2026-08-01T16:20:01.675800Z","submitted_at":"2013-12-19T17:46:46Z","title":"k-Sparse Autoencoders","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1312.5663","snapshot_observed_at":"2026-08-06T11:01:14.986525Z","title":"and Frey, B","venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.986525Z"},"links":{"cited_paper":"/paper/1312.5663","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:d1ada23c9c41951139236f03cd7d0990bc60e4eb97e689d22241dd1a8809cb90","observation_id":"f2ffc54d-ff52-4ec8-9a4c-1eca17ef4fb1","resolution":{"observed_at":"2026-08-06T11:01:14.986525Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T11:01:14.989995Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.989995Z"},"links":{"citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:eccf054cbd0878c37b6140c125b6145dd51e24d7dd2028377c63c4e0870835d3","observation_id":"9061660b-fc77-436e-94c7-068de140ee8a","resolution":{"observed_at":"2026-08-06T11:01:14.989995Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.19647","last_updated":"2025-03-27T05:44:45Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-03-28T17:56:07Z","title":"Sparse Feature Circuits: Discovering and Editing Interpretable Causal Graphs in Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.19647","snapshot_observed_at":"2026-08-06T11:01:14.993002Z","title":"J., Belinkov, Y., Bau, D., and Mueller, A","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.993002Z"},"links":{"cited_paper":"/paper/2403.19647","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:912f429298a89d135c6c6bf03cd80e7ad1764456ef8c9ccd0b9b2f727cfd49ee","observation_id":"3ea7fb9a-8e14-4c15-8014-6f4e2fe1fda6","resolution":{"observed_at":"2026-08-06T11:01:14.993002Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2005.00661","last_updated":"2020-05-02T00:09:16Z","snapshot_observed_at":"2026-07-06T09:16:59.519524Z","submitted_at":"2020-05-02T00:09:16Z","title":"On Faithfulness and Factuality in Abstractive Summarization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2005.00661","snapshot_observed_at":"2026-08-06T11:01:14.996355Z","title":"On faithfulness and factuality in abstractive summarization","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.996355Z"},"links":{"cited_paper":"/paper/2005.00661","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:125b37e94102a27d5a042a53645342aa3bc981fc73a51f8b90ad5d7703655f95","observation_id":"e78388c6-ec3c-42d3-8241-5bd3ed2b75d5","resolution":{"observed_at":"2026-08-06T11:01:14.996355Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.07404","last_updated":"2025-05-30T15:21:51Z","snapshot_observed_at":"2026-07-06T19:48:44.991187Z","submitted_at":"2024-11-11T22:22:21Z","title":"Controllable Context Sensitivity and the Knob Behind It","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.07404","snapshot_observed_at":"2026-08-06T11:01:14.999648Z","title":"Controllable Context Sensitivity and the Knob Behind It","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:14.999648Z"},"links":{"cited_paper":"/paper/2411.07404","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:5db3259fafc9b4d65acd8be99cc6828ee707f26c107b739c4e5b7f9b7e9738f2","observation_id":"f3a57042-15f0-415e-ae1b-0d9ce27ab1a2","resolution":{"observed_at":"2026-08-06T11:01:14.999648Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T11:01:15.549539Z","title":"A., and Kriegeskorte, N","venue":null,"work_id":"93b1d326-6cb2-4327-bda3-fd14350418cb","year":2009},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.002911Z"},"links":{"citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:67ba3f0321dd9f29bb8ebee0055051d9d471d717bdc42153f2d8c6c54a420f9c","observation_id":"99f8c185-00f4-4cde-8bf2-6a90e1ca62a0","resolution":{"observed_at":"2026-08-06T11:01:15.552912Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09130","last_updated":"2021-02-18T03:07:28Z","snapshot_observed_at":"2026-08-02T06:07:07.294413Z","submitted_at":"2021-02-18T03:07:28Z","title":"Entity-level Factual Consistency of Abstractive Text Summarization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.09130","snapshot_observed_at":"2026-08-06T11:01:15.005942Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.005942Z"},"links":{"cited_paper":"/paper/2102.09130","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:ca19e2f403a6a5e0ca64b7b196138cfa753f4fb3ace8329f9aa4b5f53810d755","observation_id":"010d2a4c-07c6-495d-8649-4533985454be","resolution":{"observed_at":"2026-08-06T11:01:15.005942Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.00941","last_updated":"2023-09-07T20:36:48Z","snapshot_observed_at":"2026-07-06T16:13:34.788427Z","submitted_at":"2023-09-02T13:37:34Z","title":"Emergent Linear Representations in World Models of Self-Supervised Sequence Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.00941","snapshot_observed_at":"2026-08-06T11:01:15.009218Z","title":"Emergent linear representations in world models of self-supervised sequence models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.009218Z"},"links":{"cited_paper":"/paper/2309.00941","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:b96a84327e0446f3a957ddca8cb665be45f9c09f2990d567267372fa419f7739","observation_id":"575fcf43-1924-448c-a64e-02a093430aec","resolution":{"observed_at":"2026-08-06T11:01:15.009218Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1808.08745","last_updated":"2018-08-27T09:08:18Z","snapshot_observed_at":"2026-07-06T06:57:36.335954Z","submitted_at":"2018-08-27T09:08:18Z","title":"Don't Give Me the Details, Just the Summary! Topic-Aware Convolutional Neural Networks for Extreme Summarization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1808.08745","snapshot_observed_at":"2026-08-06T11:01:15.012396Z","title":"B., and Lapata, M","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.012396Z"},"links":{"cited_paper":"/paper/1808.08745","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:9000de27ae80080a7c415e09b34741ef579cf23cdb692a28b167bbceda229033","observation_id":"18e6b3a5-e525-4361-afe7-0f8cc9432132","resolution":{"observed_at":"2026-08-06T11:01:15.012396Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.03658","last_updated":"2024-07-17T22:24:27Z","snapshot_observed_at":"2026-07-06T16:43:58.947915Z","submitted_at":"2023-11-07T01:59:11Z","title":"The Linear Representation Hypothesis and the Geometry of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.03658","snapshot_observed_at":"2026-08-06T11:01:15.015631Z","title":"J., and Veitch, V","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.015631Z"},"links":{"cited_paper":"/paper/2311.03658","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:b083f81a1fc555bbe47b275ec5e88a299847b702201275b1435ec6aca1accd46","observation_id":"bd9dcc5f-e8d4-4cfd-a90b-f22ad9631edf","resolution":{"observed_at":"2026-08-06T11:01:15.015631Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T11:01:15.018933Z","title":"A practical review of mechanistic interpretability for transformer-based language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.018933Z"},"links":{"citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:85c148895611f6184637f79f40d9b02aff54d23821954f7632be8f1c998e90df","observation_id":"2bc2c98d-2c5c-4c7f-b88b-53efa68d1e78","resolution":{"observed_at":"2026-08-06T11:01:15.018933Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T11:01:15.539025Z","title":"Hallushield: A mechanistic approach to hallucination resistant models","venue":null,"work_id":"950fcbb1-e8d1-447a-aa0a-8c4bab7268a9","year":2025},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.022626Z"},"links":{"citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:9c989d87e3d5f8bef726502b211f5f175ba7ef07b8ca2d854544ed00d8d6395f","observation_id":"0b84aaaa-522b-4386-bbfc-6f33d8465876","resolution":{"observed_at":"2026-08-06T11:01:15.542375Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1704.04368","last_updated":"2017-04-25T05:47:50Z","snapshot_observed_at":"2026-07-06T05:37:49.239583Z","submitted_at":"2017-04-14T07:55:19Z","title":"Get To The Point: Summarization with Pointer-Generator Networks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1704.04368","snapshot_observed_at":"2026-08-06T11:01:15.025985Z","title":"J., and Manning, C","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.025985Z"},"links":{"cited_paper":"/paper/1704.04368","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:1c51c650bdd0fbac588525fea94f566ed8b863fa862ab16bdd897a8eaca5a0d8","observation_id":"44e1a28c-b63a-40a4-a831-d01c2428d807","resolution":{"observed_at":"2026-08-06T11:01:15.025985Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.09971","last_updated":"2024-07-11T06:31:31Z","snapshot_observed_at":"2026-08-05T06:06:57.851845Z","submitted_at":"2024-04-15T17:48:46Z","title":"Constructing Benchmarks and Interventions for Combating Hallucinations in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.09971","snapshot_observed_at":"2026-08-06T11:01:15.029328Z","title":"Constructing benchmarks and interventions for combating hallucinations in llms","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.029328Z"},"links":{"cited_paper":"/paper/2404.09971","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:6123f6f78b974525dc91e6bb1d138b075ff45a62d60573f5bace9346bd23b7d8","observation_id":"e60da72d-a791-43c5-92f8-e6bdf59ff759","resolution":{"observed_at":"2026-08-06T11:01:15.029328Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.12964","last_updated":"2025-08-25T16:47:29Z","snapshot_observed_at":"2026-08-04T12:08:20.637547Z","submitted_at":"2025-02-18T15:46:31Z","title":"Trust Me, I'm Wrong: LLMs Hallucinate with Certainty Despite Knowing the Answer","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.12964","snapshot_observed_at":"2026-08-06T11:01:15.033088Z","title":"Trust Me, I'm Wrong: High-Certainty Hallucinations in LLMs","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.033088Z"},"links":{"cited_paper":"/paper/2502.12964","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:dd2d2a546f71bded9d217f001c2c82d636c2a0c1c77eb9990880503e2a9a6c64","observation_id":"29f924ff-d768-4776-ab16-9827381fb2ac","resolution":{"observed_at":"2026-08-06T11:01:15.033088Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.11877","last_updated":"2023-11-12T13:14:36Z","snapshot_observed_at":"2026-07-06T16:34:54.717394Z","submitted_at":"2023-10-18T11:01:09Z","title":"The Curious Case of Hallucinatory (Un)answerability: Finding Truths in the Hidden States of Over-Confident Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.11877","snapshot_observed_at":"2026-08-06T11:01:15.036213Z","title":"The curious case of hallucinatory (un) answerability: Finding truths in the hidden states of over-confident large language models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.036213Z"},"links":{"cited_paper":"/paper/2310.11877","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:c8f66cf6f15e5732991621f25e16920dd6da1a13be6d607284815f9c273d29bc","observation_id":"447bea56-d038-44d7-9624-3c81ac77f357","resolution":{"observed_at":"2026-08-06T11:01:15.036213Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T11:01:15.039601Z","title":"Redeep: Detecting hallucination in retrieval augmented generation via mechanistic interpretability","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.039601Z"},"links":{"citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:bc2b21ea9b7dd80fa4dcb8c817148efe878d252bcbf4723c585e53cfac444b88","observation_id":"0ba1cff9-5cf9-4d5a-9248-59565109c9fc","resolution":{"observed_at":"2026-08-06T11:01:15.039601Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21424","last_updated":"2024-08-09T11:58:55Z","snapshot_observed_at":"2026-07-06T18:54:54.118506Z","submitted_at":"2024-07-31T08:19:06Z","title":"Cost-Effective Hallucination Detection for LLMs","version":2},"cited_work":{"arxiv_id":"2407.21424","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.21424","snapshot_observed_at":"2026-08-06T11:01:15.141632Z","title":"Cost-Effective Hallucination Detection for LLMs","venue":"cs.CL","work_id":"c967a09e-f3b9-4c71-a30c-977121571869","year":2024},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.042383Z"},"links":{"cited_paper":"/paper/2407.21424","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:822024549d56fc20f3fc4e84b3f520f51164e41bfc36a240e9f20af4d4c17a24","observation_id":"2d009fdf-20bb-4135-8758-51881219073e","resolution":{"observed_at":"2026-08-06T11:01:15.147227Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T11:01:15.045482Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.045482Z"},"links":{"citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:fdc3bdeb059b264069f730dbc9656d0a032d3cc35fdaaa88ca3e569087fcbefa","observation_id":"c43006d9-345a-43c0-b362-fbd2b37b9da3","resolution":{"observed_at":"2026-08-06T11:01:15.045482Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.15098","last_updated":"2024-04-17T04:25:21Z","snapshot_observed_at":"2026-07-06T16:23:57.143116Z","submitted_at":"2023-09-26T17:48:55Z","title":"Attention Satisfies: A Constraint-Satisfaction Lens on Factual Errors of Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.15098","snapshot_observed_at":"2026-08-06T11:01:15.050200Z","title":"Attention satisfies: A constraint-satisfaction lens on factual errors of language models, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.050200Z"},"links":{"cited_paper":"/paper/2309.15098","citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:0b6b212920bbfcb1875b2a16006c45356fcfa1e869c4fd392d136c9722a47f5b","observation_id":"59842f00-0487-43b2-9205-a8bc4677d5e6","resolution":{"observed_at":"2026-08-06T11:01:15.050200Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T11:01:15.053772Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-06T11:01:15.053772Z"},"links":{"citing_paper":"/paper/2507.23221"},"observation_digest":"sha256:143e80a851a7e0aaa70714ce232654c55bc70d9fd8f0d8bc6b3c339b65f246fe","observation_id":"dbcffd29-bb0b-4789-b68e-c050bb0b1f89","resolution":{"observed_at":"2026-08-06T11:01:15.053772Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2507.23221","last_updated":"2025-07-31T03:26:57Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-07T01:51:07.704268Z","submitted_at":"2025-07-31T03:26:57Z","title":"A Single Direction of Truth: An Observer Model's Linear Residual Probe Exposes and Steers Contextual Hallucinations"},"reference_resolution":{"displayed":37,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":31,"verified_exact":1,"verified_fuzzy":5},"total_outbound_references":37},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 37 of 37 outbound references and 2 inbound Pith citation observations for arXiv:2507.23221."}