{"as_of":"2026-08-07T20:57:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:3a7ce62d9fe32e888f036d3f4b108e99522213433895ceecabf1885615502e5b","coverage":[{"denominator":16,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":16,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T11:17:36.764467Z","state":"measured"},{"denominator":16,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":16,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2510.06096/citation-record","integrity":"/paper/2510.06096/integrity","json":"/paper/2510.06096/citation-record.json","paper":"/paper/2510.06096"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2402.12264","last_updated":"2025-05-20T19:59:52Z","snapshot_observed_at":"2026-08-04T17:39:07.204509Z","submitted_at":"2024-02-19T16:26:00Z","title":"Uncertainty quantification in fine-tuned LLMs using LoRA ensembles","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12264","snapshot_observed_at":"2026-08-04T11:17:36.682867Z","title":"Uncertainty quantification in fine-tuned llms using lora ensembles.arXiv preprint arXiv:2402.12264,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.682867Z"},"links":{"cited_paper":"/paper/2402.12264","citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:52d231eb673bf946a5dfe049ca024970190da24c7701ceffb6e9e31362194b88","observation_id":"1a4c187e-5c6c-4a98-9c82-14f8146df2af","resolution":{"observed_at":"2026-08-04T11:17:36.682867Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15217","last_updated":"2023-09-11T17:25:24Z","snapshot_observed_at":"2026-08-03T19:11:09.671782Z","submitted_at":"2023-07-27T22:29:25Z","title":"Open Problems and Fundamental Limitations of Reinforcement Learning from Human Feedback","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.15217","snapshot_observed_at":"2026-08-04T11:17:36.700516Z","title":"Open problems and fundamental limitations of reinforcement learning from human feedback.arXiv preprint arXiv:2307.15217,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.700516Z"},"links":{"cited_paper":"/paper/2307.15217","citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:615a81148496cd98ef92b5e08f5d01963e9d6f693ddb49b4b33de6624eda1f53","observation_id":"b2174ce8-2c66-4898-a84a-1f6dcd1579e8","resolution":{"observed_at":"2026-08-04T11:17:36.700516Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2209.10652","last_updated":"2022-09-21T20:49:26Z","snapshot_observed_at":"2026-07-06T13:54:56.779166Z","submitted_at":"2022-09-21T20:49:26Z","title":"Toy Models of Superposition","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2209.10652","snapshot_observed_at":"2026-08-04T11:17:36.706487Z","title":"Nelson Elhage, Tristan Hume, Catherine Olsson, Nicholas Schiefer, Tom Henighan, Shauna Kravec, Zac Hatfield-Dodds, Robert Lasenby, Dawn Drain, Carol Chen, et al","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.706487Z"},"links":{"cited_paper":"/paper/2209.10652","citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:8d53ebf5b524b638e6432d565194457145a9477687a3e3450f98813a3a5dac88","observation_id":"8e0b2168-f3dc-4494-9489-80b906e6c9ba","resolution":{"observed_at":"2026-08-04T11:17:36.706487Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2103.14659","last_updated":"2021-03-26T18:01:48Z","snapshot_observed_at":"2026-07-06T10:53:51.296957Z","submitted_at":"2021-03-26T18:01:48Z","title":"Alignment of Language Agents","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2103.14659","snapshot_observed_at":"2026-08-04T11:17:36.718355Z","title":"Zachary Kenton, Tom Everitt, Laura Weidinger, Iason Gabriel, Vladimir Mikulik, and Geoffrey Irving","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.718355Z"},"links":{"cited_paper":"/paper/2103.14659","citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:3519ce064a387bfe0f852b0e2d1dc5bef515746d9c780994673a30f33c333546","observation_id":"b9c2b000-f501-40e1-85e4-21a600ba3d39","resolution":{"observed_at":"2026-08-04T11:17:36.718355Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.07236","last_updated":"2025-09-05T17:54:18Z","snapshot_observed_at":"2026-08-06T18:43:28.146997Z","submitted_at":"2025-07-09T19:13:25Z","title":"Simple Yet Effective: An Information-Theoretic Approach to Multi-LLM Uncertainty Quantification","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.07236","snapshot_observed_at":"2026-08-04T11:17:36.724299Z","title":"An information-theoretic perspective on multi-llm uncertainty estimation.arXiv preprint arXiv:2507.07236,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.724299Z"},"links":{"cited_paper":"/paper/2507.07236","citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:a9145dd7087f83a77ad7a37fae35c29c750d0b3c046c6fb6aa88912dcb572004","observation_id":"b8c570dd-d43a-4150-b0d1-544d1c28b749","resolution":{"observed_at":"2026-08-04T11:17:36.724299Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15124","last_updated":"2025-04-14T22:39:09Z","snapshot_observed_at":"2026-07-06T19:55:37.400185Z","submitted_at":"2024-11-22T18:44:04Z","title":"Tulu 3: Pushing Frontiers in Open Language Model Post-Training","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.15124","snapshot_observed_at":"2026-08-04T11:17:36.730234Z","title":"Tulu 3: Pushing frontiers in open language model post-training.arXiv preprint arXiv:2411.15124,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.730234Z"},"links":{"cited_paper":"/paper/2411.15124","citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:5e8ab86ad7bf5b5dd1a4a0828092045e533791a7c9a22492dd1211f654d9c943","observation_id":"7206c980-9ad0-4ebd-b63c-4191a50e06af","resolution":{"observed_at":"2026-08-04T11:17:36.730234Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T11:17:36.748639Z","title":"Bayesian prompt ensembles: Model uncertainty estimation for black-box large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.748639Z"},"links":{"citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:797417a3664d59b0eb5289dcf7e783d823dd1f2a863fece17cf461b8995013b7","observation_id":"a8960377-10e3-4407-a63f-80af2ab5a6f5","resolution":{"observed_at":"2026-08-04T11:17:36.748639Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.13210","last_updated":"2024-07-03T00:23:41Z","snapshot_observed_at":"2026-07-06T17:33:00.716054Z","submitted_at":"2024-02-20T18:20:59Z","title":"Bayesian Reward Models for LLM Alignment","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.13210","snapshot_observed_at":"2026-08-04T11:17:36.758976Z","title":"Adam X Yang, Maxime Robeyns, Thomas Coste, Zhengyan Shi, Jun Wang, Haitham Bou- Ammar, and Laurence Aitchison","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.758976Z"},"links":{"cited_paper":"/paper/2402.13210","citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:7927a25544f8a14809edc13eab0a19217cb6ba0bfb47f3d70811cba40e42b998","observation_id":"4440dea5-165d-4523-8693-b37bfd7be1bc","resolution":{"observed_at":"2026-08-04T11:17:36.758976Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T11:17:36.764467Z","title":"toget back at fuckboys","venue":null,"work_id":null,"year":1994},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.764467Z"},"links":{"citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:55a7d2f5891bde5112d7a22f36d5400548cf10977e675dff8d7c7d2535a12c68","observation_id":"6feb97e1-b97a-4e3f-a7d9-b00823355b0e","resolution":{"observed_at":"2026-08-04T11:17:36.764467Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.10060","last_updated":"2026-04-17T19:21:30Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-06-11T18:00:00Z","title":"Textual Bayes: Quantifying Prompt Uncertainty in LLM-Based Systems","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.10060","snapshot_observed_at":"2026-08-04T11:17:36.735976Z","title":"Textual bayes: Quantifying uncertainty in llm-based systems.arXiv preprint arXiv:2506.10060,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":2007,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.735976Z"},"links":{"cited_paper":"/paper/2506.10060","citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:f7ccf5849e42e311183936aaccfb9d7b2676d9bc9977581cb5efea9271a25e91","observation_id":"86eb37f9-d4fa-4adb-8efd-97e1592fb523","resolution":{"observed_at":"2026-08-04T11:17:36.735976Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2209.07858","last_updated":"2022-11-22T19:12:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-08-23T23:37:14Z","title":"Red Teaming Language Models to Reduce Harms: Methods, Scaling Behaviors, and Lessons Learned","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2209.07858","snapshot_observed_at":"2026-08-04T11:17:36.712466Z","title":"Red teaming language models to reduce harms: Methods, scaling behaviors, and lessons learned.arXiv preprint arXiv:2209.07858,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":2017,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.712466Z"},"links":{"cited_paper":"/paper/2209.07858","citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:2d566b7b1d4f4ad17fc43fc43e8d6a4224f7e2da7ee3250da621140f1e4901a0","observation_id":"69cb4738-ee36-4eff-815f-889f1f09edca","resolution":{"observed_at":"2026-08-04T11:17:36.712466Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.13158","last_updated":"2025-07-17T14:22:24Z","snapshot_observed_at":"2026-08-07T07:16:46.127762Z","submitted_at":"2025-07-17T14:22:24Z","title":"Inverse Reinforcement Learning Meets Large Language Model Post-Training: Basics, Advances, and Opportunities","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.13158","snapshot_observed_at":"2026-08-04T11:17:36.742733Z","title":"Inverse reinforcement learning meets large language model post-training: Basics, advances, and opportunities.arXiv preprint arXiv:2507.13158,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.742733Z"},"links":{"cited_paper":"/paper/2507.13158","citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:7a83625011656cb86f8b80db0136aedf297f0e18fa4780b36bbf536b5c3b5932","observation_id":"58ce233a-7689-4a88-977a-8ab27b2358d0","resolution":{"observed_at":"2026-08-04T11:17:36.742733Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T11:17:36.688704Z","title":"Emergent misalignment: Narrow finetuning can produce broadly misaligned llms.arXiv preprint arXiv:2502.17424,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.688704Z"},"links":{"citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:8b4a970a05db19b1fbc9df50855938bae9feb220b73b1d2567149aa144efb98d","observation_id":"4e9a5758-72b3-4456-9766-bff79fea411a","resolution":{"observed_at":"2026-08-04T11:17:36.688704Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2204.05862","last_updated":"2022-04-12T15:02:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-04-12T15:02:38Z","title":"Training a Helpful and Harmless Assistant with Reinforcement Learning from Human Feedback","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2204.05862","snapshot_observed_at":"2026-08-04T11:17:36.676213Z","title":"URL https://huggingface.co/datasets/allenai/real-toxicity-prompts","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.676213Z"},"links":{"cited_paper":"/paper/2204.05862","citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:0202eca110236b957c937ab2b4bffc4a48c821e222ed604e53f74cb4ef875eeb","observation_id":"a7416fa0-5e01-4096-8f7b-9953471000e4","resolution":{"observed_at":"2026-08-04T11:17:36.676213Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2112.04359","last_updated":"2021-12-08T16:09:48Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-12-08T16:09:48Z","title":"Ethical and social risks of harm from Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2112.04359","snapshot_observed_at":"2026-08-04T11:17:36.753954Z","title":"Taxonomy of risks posed by language models.arXiv preprint arXiv:2112.04359,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.753954Z"},"links":{"cited_paper":"/paper/2112.04359","citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:6774d848f1248da365189ec916bec80c8b57bc080177b578d3342c6891f4c4c8","observation_id":"cd63256c-b9d3-431d-8e54-c3cdc8957097","resolution":{"observed_at":"2026-08-04T11:17:36.753954Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2108.07258","last_updated":"2022-07-12T23:45:14Z","snapshot_observed_at":"2026-08-02T09:20:40.804790Z","submitted_at":"2021-08-16T17:50:08Z","title":"On the Opportunities and Risks of Foundation Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2108.07258","snapshot_observed_at":"2026-08-04T11:17:36.694949Z","title":"Hudson, Ehsan Adeli, Russ Altman, and et al","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives","version":3},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-04T11:17:36.694949Z"},"links":{"cited_paper":"/paper/2108.07258","citing_paper":"/paper/2510.06096"},"observation_digest":"sha256:a4cf511f3928d1bd32fb616eea6b5d03699387c69ee91ad3896f1c1065235c3b","observation_id":"ce3d0d09-b88d-4b46-a180-347335720062","resolution":{"observed_at":"2026-08-04T11:17:36.694949Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2510.06096","last_updated":"2026-06-26T18:48:41Z","latest_version":3,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-07T07:17:08.647250Z","submitted_at":"2025-10-07T16:25:14Z","title":"The Alignment Auditor: A Bayesian Framework for Verifying and Refining LLM Objectives"},"reference_resolution":{"displayed":16,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":16,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":16},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 16 of 16 outbound references and 0 inbound Pith citation observations for arXiv:2510.06096."}