{"as_of":"2026-08-07T11:38:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:528ae3c2bfe76ddbd13ba67d56e5a14f33abc12570862420d7cb9a1a49530bb0","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":18,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":18,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":18,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":18,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T05:41:25.526186Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":44,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":"2304.11082","doi":"10.48550/arxiv.2304.11082","metadata_source":"arxiv_reference","pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Fundamental limitations of alignment in large language models","venue":"arXiv (Cornell University)","work_id":"818b2f89-488c-4c42-beee-dc4ed4223978","year":2024},"citing_paper":{"arxiv_id":"2307.02483","last_updated":"2023-07-05T17:58:10Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-07-05T17:58:10Z","title":"Jailbroken: How Does LLM Safety Training Fail?","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-14T18:17:42.752997Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2307.02483"},"observation_digest":"sha256:eff3823354de6f9d126110b0f516c1059b9117226353becffa001727f0e4953b","observation_id":"bc298fbf-da94-4833-8e0d-3f5871613c44","resolution":{"observed_at":"2026-05-14T18:17:42.828873Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":"2304.11082","doi":"10.48550/arxiv.2304.11082","metadata_source":"arxiv_reference","pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Fundamental limitations of alignment in large language models","venue":"arXiv (Cornell University)","work_id":"818b2f89-488c-4c42-beee-dc4ed4223978","year":2024},"citing_paper":{"arxiv_id":"2307.15043","last_updated":"2023-12-20T20:48:57Z","snapshot_observed_at":"2026-07-06T15:59:23.019044Z","submitted_at":"2023-07-27T17:49:12Z","title":"Universal and Transferable Adversarial Attacks on Aligned Language Models","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-24T07:42:09.112946Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2307.15043"},"observation_digest":"sha256:85ec88e525036e0ad90d94b0ff2185455b061767ac0552df96517f4e3442d405","observation_id":"3efae1f4-54d7-4687-aafa-7b082e8c0467","resolution":{"observed_at":"2026-05-24T07:44:08.470910Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-07T05:41:25.526186Z","title":"Fundamental limitations of alignment in large language models, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07330","last_updated":"2025-06-09T00:11:06Z","snapshot_observed_at":"2026-08-07T08:21:07.384830Z","submitted_at":"2025-06-09T00:11:06Z","title":"JavelinGuard: Low-Cost Transformer Architectures for LLM Security","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-07T05:41:25.526186Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2506.07330"},"observation_digest":"sha256:8a9142cc97f5e365eefdefd17a9913911051be6d2bfc96763473fbe48ff49faf","observation_id":"1885a77b-acf8-4ed0-bd2d-6a19a4d708c2","resolution":{"observed_at":"2026-08-07T05:41:25.526186Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-06T14:54:36.557710Z","title":"Ho, Percy Liang, and Arvind Narayanan","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.17787","last_updated":"2025-07-23T09:50:17Z","snapshot_observed_at":"2026-08-06T14:47:37.458206Z","submitted_at":"2025-07-23T09:50:17Z","title":"Hyperbolic Deep Learning for Foundation Models: A Survey","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-06T14:54:36.557710Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2507.17787"},"observation_digest":"sha256:e7fd0b4ef58d5cba4ff9d6a343a83c8262c637802890289d964ac52eda0b068a","observation_id":"6decc58e-6b91-4b12-af93-4bdc9bc1e436","resolution":{"observed_at":"2026-08-06T14:54:36.557710Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-05T14:51:03.796947Z","title":"Fundamental Limitations of Alignment in Large Language Models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.20848","last_updated":"2025-08-28T14:40:27Z","snapshot_observed_at":"2026-08-05T14:51:02.831302Z","submitted_at":"2025-08-28T14:40:27Z","title":"JADES: A Universal Framework for Jailbreak Assessment via Decompositional Scoring","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-05T14:51:03.796947Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2508.20848"},"observation_digest":"sha256:9fe39eff370b28f488657bbc9b09ace9eaae3546fccd85c55cddb60f40e87eba","observation_id":"560af3c0-046c-40eb-90de-d44ddc6a6834","resolution":{"observed_at":"2026-08-05T14:51:03.796947Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":"2304.11082","doi":"10.48550/arxiv.2304.11082","metadata_source":"arxiv_reference","pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Fundamental limitations of alignment in large language models","venue":"arXiv (Cornell University)","work_id":"818b2f89-488c-4c42-beee-dc4ed4223978","year":2024},"citing_paper":{"arxiv_id":"2512.10100","last_updated":"2026-04-07T12:47:26Z","snapshot_observed_at":"2026-07-06T22:38:40.044448Z","submitted_at":"2025-12-10T21:44:10Z","title":"Robust AI Security and Alignment: A Sisyphean Endeavor?","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-16T22:58:29.306435Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2512.10100"},"observation_digest":"sha256:fa7eaa07b90b0c0ff989c2d5693cfdede6a39c175ce0fb4bd95f8727e4144f84","observation_id":"5e0dd17a-3fe9-4cf7-97a3-67fc29cff162","resolution":{"observed_at":"2026-05-16T22:58:38.228561Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":"2304.11082","doi":"10.48550/arxiv.2304.11082","metadata_source":"arxiv_reference","pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Fundamental limitations of alignment in large language models","venue":"arXiv (Cornell University)","work_id":"818b2f89-488c-4c42-beee-dc4ed4223978","year":2024},"citing_paper":{"arxiv_id":"2604.17359","last_updated":"2026-08-05T23:18:46Z","snapshot_observed_at":"2026-08-07T10:37:49.535655Z","submitted_at":"2026-04-19T10:05:25Z","title":"Plausible Patients, Impossible Populations: Auditing Epidemiological Fidelity in Large Language Model Mental Health Simulations","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-10T06:04:24.077276Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2604.17359"},"observation_digest":"sha256:cdb8adb497ef79f3f3005d138255797598d006882a41fb15ae45da0612aa36f6","observation_id":"6a98c7df-c782-4da9-8ac3-581311a95014","resolution":{"observed_at":"2026-05-10T06:06:18.723757Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":"2304.11082","doi":"10.48550/arxiv.2304.11082","metadata_source":"arxiv_reference","pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Fundamental limitations of alignment in large language models","venue":"arXiv (Cornell University)","work_id":"818b2f89-488c-4c42-beee-dc4ed4223978","year":2024},"citing_paper":{"arxiv_id":"2605.08496","last_updated":"2026-05-08T21:21:59Z","snapshot_observed_at":"2026-07-06T23:20:43.010758Z","submitted_at":"2026-05-08T21:21:59Z","title":"Latent Personality Alignment: Improving Harmlessness Without Mentioning Harms","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-12T01:46:49.586630Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2605.08496"},"observation_digest":"sha256:81bd41fe9931c3a37a43527c3b906b86b28730413d03e7ced0b0a06a40c73db6","observation_id":"5d7d489b-4ec9-4eb0-8d06-8dab12957bd0","resolution":{"observed_at":"2026-05-12T07:51:43.581200Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":"2304.11082","doi":"10.48550/arxiv.2304.11082","metadata_source":"arxiv_reference","pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Fundamental limitations of alignment in large language models","venue":"arXiv (Cornell University)","work_id":"818b2f89-488c-4c42-beee-dc4ed4223978","year":2024},"citing_paper":{"arxiv_id":"2605.12809","last_updated":"2026-05-12T23:01:29Z","snapshot_observed_at":"2026-07-06T23:24:27.821980Z","submitted_at":"2026-05-12T23:01:29Z","title":"Correcting Influence: Unboxing LLM Outputs with Orthogonal Latent Spaces","version":1},"reference_index":206,"source":"arxiv_source","source_observed_at":"2026-05-14T20:17:01.224864Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2605.12809"},"observation_digest":"sha256:379c72ac3634b5652517145bc8c89ef57bdc3aa52f2c8aad714e8134b7e1517d","observation_id":"a9fe88b3-3c82-4db0-9281-54d53f7f3910","resolution":{"observed_at":"2026-05-14T20:17:55.556966Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":"2304.11082","doi":"10.48550/arxiv.2304.11082","metadata_source":"arxiv_reference","pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Fundamental limitations of alignment in large language models","venue":"arXiv (Cornell University)","work_id":"818b2f89-488c-4c42-beee-dc4ed4223978","year":2024},"citing_paper":{"arxiv_id":"2605.16339","last_updated":"2026-05-07T16:48:48Z","snapshot_observed_at":"2026-07-06T23:27:29.931962Z","submitted_at":"2026-05-07T16:48:48Z","title":"Preference Instability in Reward Models: Detection and Mitigation via Sparse Autoencoders","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-20T22:48:54.238767Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2605.16339"},"observation_digest":"sha256:7e3cc804075bb48678c0fbfeecb2052c2401a98203806c52320ee26bdf5705ab","observation_id":"9b99e65a-6ad6-4fdd-bcf4-65cbeffdebe5","resolution":{"observed_at":"2026-05-20T22:49:10.130331Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":"2304.11082","doi":"10.48550/arxiv.2304.11082","metadata_source":"arxiv_reference","pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Fundamental limitations of alignment in large language models","venue":"arXiv (Cornell University)","work_id":"818b2f89-488c-4c42-beee-dc4ed4223978","year":2024},"citing_paper":{"arxiv_id":"2605.20641","last_updated":"2026-05-20T02:55:56Z","snapshot_observed_at":"2026-07-06T23:31:13.042581Z","submitted_at":"2026-05-20T02:55:56Z","title":"Trusted Weights, Treacherous Optimizations? Optimization-Triggered Backdoor Attacks on LLMs","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-21T04:45:35.079192Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2605.20641"},"observation_digest":"sha256:8cedf9aa9cb6da19aabcb1e4c24031b722ddfeb5f886621af6e2adf3e0074da5","observation_id":"3de5c74b-9fdd-42c4-8631-8a55023605de","resolution":{"observed_at":"2026-05-21T04:49:35.636729Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":"2304.11082","doi":"10.48550/arxiv.2304.11082","metadata_source":"arxiv_reference","pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Fundamental limitations of alignment in large language models","venue":"arXiv (Cornell University)","work_id":"818b2f89-488c-4c42-beee-dc4ed4223978","year":2024},"citing_paper":{"arxiv_id":"2605.25739","last_updated":"2026-07-19T12:16:05Z","snapshot_observed_at":"2026-08-02T13:17:54.641684Z","submitted_at":"2026-05-25T11:51:08Z","title":"The Behavioral Credibility Trilemma: When Calibrated Autonomy Becomes Impossible","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-29T22:28:02.493124Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2605.25739"},"observation_digest":"sha256:98eddb972df7fd947c43cc9681a70390755694a8e4ba3b97341681268241279d","observation_id":"e70bd3cf-98ea-48d2-84aa-478bb6a88973","resolution":{"observed_at":"2026-06-29T22:34:01.928075Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-02T13:17:57.695515Z","title":"Fundamental limitations of alignment in large language models.arXiv preprint arXiv:2304.11082,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2605.25739","last_updated":"2026-07-19T12:16:05Z","snapshot_observed_at":"2026-08-02T13:17:54.641684Z","submitted_at":"2026-05-25T11:51:08Z","title":"The Behavioral Credibility Trilemma: When Calibrated Autonomy Becomes Impossible","version":2},"reference_index":2012,"source":"pdf_text","source_observed_at":"2026-08-02T13:17:57.695515Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2605.25739"},"observation_digest":"sha256:e4f322c852110af7ddd7bbb1c87fca0f971fdb5efe97d5b944dfb80e56169ccb","observation_id":"669a96b2-01be-4161-880d-11824b753b97","resolution":{"observed_at":"2026-08-02T13:17:57.695515Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":"2304.11082","doi":"10.48550/arxiv.2304.11082","metadata_source":"arxiv_reference","pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Fundamental limitations of alignment in large language models","venue":"arXiv (Cornell University)","work_id":"818b2f89-488c-4c42-beee-dc4ed4223978","year":2024},"citing_paper":{"arxiv_id":"2605.30169","last_updated":"2026-07-02T23:58:01Z","snapshot_observed_at":"2026-08-01T18:48:30.510137Z","submitted_at":"2026-05-28T16:20:19Z","title":"Dissociative Identity: Language Model Agents Lack Grounding for Reputation Mechanisms","version":2},"reference_index":136,"source":"pdf_text","source_observed_at":"2026-06-29T00:26:54.019256Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2605.30169"},"observation_digest":"sha256:1197bff53eeb26e82f1ee9d5002dea8edf0531e7f9a64b619956e9d13a941eb9","observation_id":"7fab3428-33eb-4396-bbc3-6788746fba32","resolution":{"observed_at":"2026-06-29T00:32:53.150041Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":"2304.11082","doi":"10.48550/arxiv.2304.11082","metadata_source":"arxiv_reference","pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Fundamental limitations of alignment in large language models","venue":"arXiv (Cornell University)","work_id":"818b2f89-488c-4c42-beee-dc4ed4223978","year":2024},"citing_paper":{"arxiv_id":"2606.00485","last_updated":"2026-05-30T02:37:18Z","snapshot_observed_at":"2026-08-06T04:47:49.161962Z","submitted_at":"2026-05-30T02:37:18Z","title":"Confused ChatGPT: Cross-App Context Poisoning via First-Party APIs","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-06-28T18:55:45.465260Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2606.00485"},"observation_digest":"sha256:1351c450c54aa5506aec0d8d1bb60b9ec9277137fe72af4617dfbc836c492d77","observation_id":"8fc14387-4f84-447c-8574-00d7ebe85e8d","resolution":{"observed_at":"2026-06-28T19:32:35.659788Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":"2304.11082","doi":"10.48550/arxiv.2304.11082","metadata_source":"arxiv_reference","pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Fundamental limitations of alignment in large language models","venue":"arXiv (Cornell University)","work_id":"818b2f89-488c-4c42-beee-dc4ed4223978","year":2024},"citing_paper":{"arxiv_id":"2606.08367","last_updated":"2026-06-06T22:59:27Z","snapshot_observed_at":"2026-08-02T14:05:19.547758Z","submitted_at":"2026-06-06T22:59:27Z","title":"Emergence World: A Platform for Evaluating Long-Horizon Multi-Agent Autonomy","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-06-27T18:36:44.265273Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2606.08367"},"observation_digest":"sha256:3f9ea24855746a0ea2e5508bb9952b6c4f1911f3646548fdb34ab2420a0453a1","observation_id":"def98b02-acfb-4a38-bbab-17645d140d20","resolution":{"observed_at":"2026-07-02T22:57:25.953535Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":"2304.11082","doi":"10.48550/arxiv.2304.11082","metadata_source":"arxiv_reference","pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Fundamental limitations of alignment in large language models","venue":"arXiv (Cornell University)","work_id":"818b2f89-488c-4c42-beee-dc4ed4223978","year":2024},"citing_paper":{"arxiv_id":"2606.12234","last_updated":"2026-06-10T15:42:15Z","snapshot_observed_at":"2026-08-05T00:08:40.346158Z","submitted_at":"2026-06-10T15:42:15Z","title":"On The Effectiveness-Fluency Trade-Off In LLM Conditioning: A Systematic Study","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-06-27T09:40:48.736006Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2606.12234"},"observation_digest":"sha256:c72d1d66b3d6005e0da081e4989fb83dae91901bedefa500486cd818a9e911a4","observation_id":"92bd060d-0b28-4cbd-9348-c6f9896922f7","resolution":{"observed_at":"2026-07-03T11:18:03.179730Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.11082","snapshot_observed_at":"2026-08-01T07:28:31.376035Z","title":"arXiv preprint arXiv:2304.11082 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.21453","last_updated":"2026-07-24T02:14:44Z","snapshot_observed_at":"2026-08-03T14:01:50.133228Z","submitted_at":"2026-07-23T15:55:29Z","title":"Test-Time Scaling via Error Localization","version":2},"reference_index":132,"source":"arxiv_source","source_observed_at":"2026-08-01T07:28:31.376035Z"},"links":{"cited_paper":"/paper/2304.11082","citing_paper":"/paper/2607.21453"},"observation_digest":"sha256:1637b8a79507cd443a1e376459ef3f31c7fdf7e572220a230d9c0a19948fa2d4","observation_id":"d459a3d0-42eb-458d-81b9-10a3b9baff94","resolution":{"observed_at":"2026-08-01T07:28:31.376035Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2304.11082/citation-record","integrity":"/paper/2304.11082/integrity","json":"/paper/2304.11082/citation-record.json","paper":"/paper/2304.11082"},"outbound":[],"paper":{"arxiv_id":"2304.11082","last_updated":"2024-06-03T12:19:16Z","latest_version":6,"primary_category":"cs.CL","snapshot_observed_at":"2026-07-06T15:18:26.300068Z","submitted_at":"2023-04-19T17:50:09Z","title":"Fundamental Limitations of Alignment in Large Language Models"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 18 inbound Pith citation observations for arXiv:2304.11082."}