{"as_of":"2026-08-08T13:36:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:ee92986d02fde60ca950dfd245cdf2d3da65b1302335778bd6f95963afc83197","coverage":[{"denominator":50,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":50,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T05:43:07.337382Z","state":"measured"},{"denominator":50,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":50,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.13774/citation-record","integrity":"/paper/2506.13774/integrity","json":"/paper/2506.13774/citation-record.json","paper":"/paper/2506.13774"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.257158Z","title":"Artificial Intelligence, Values, and Alignment","venue":null,"work_id":"7712a691-d171-4a90-82a5-77521c0d31f6","year":2020},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.102830Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:e6f7401c0b917f13c92944ddbaf7f62783887ccf8ec17884780da091d32d48b6","observation_id":"03ffc49a-1a7b-429c-8020-54a96502feb0","resolution":{"observed_at":"2026-08-07T05:43:08.261527Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.243311Z","title":"Artificial Morality: Top -down, Bottom-up, and Hybrid Approaches","venue":null,"work_id":"3383f635-6bd5-4feb-91f4-9c7e3516488e","year":2005},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.113520Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:9bed11efb4ec4c91cd7d4f17e98e8bea0d6ce6d9c61544a656f41c3b9823fbdb","observation_id":"93026cbf-d874-413a-9bd0-b2651e5c8a3b","resolution":{"observed_at":"2026-08-07T05:43:08.247705Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.228150Z","title":"Translating Principles into Practices of Digital Ethics: Five Risks of Being Unethical","venue":null,"work_id":"1b4a8211-5c62-45bf-b96c-c438d5de7c77","year":2019},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.119100Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:d46741bbcc05a1ab157ff95665fa077c1856683a64174e3141f793c8dd02ea50","observation_id":"e4ec19f7-68a5-4123-90ff-cb5f9267b597","resolution":{"observed_at":"2026-08-07T05:43:08.232427Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15217","last_updated":"2023-09-11T17:25:24Z","snapshot_observed_at":"2026-08-03T19:11:09.671782Z","submitted_at":"2023-07-27T22:29:25Z","title":"Open Problems and Fundamental Limitations of Reinforcement Learning from Human Feedback","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.15217","snapshot_observed_at":"2026-08-07T05:43:07.123823Z","title":"Open Problems and Fundamental Limitations of Reinforcement Learning from Human Feedback","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.123823Z"},"links":{"cited_paper":"/paper/2307.15217","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:950d97bac0dadd549e557d551d14502657cd3189fbe16a5265314d0e163e85f0","observation_id":"cab35d18-a447-4352-ac03-fd37e5153ac7","resolution":{"observed_at":"2026-08-07T05:43:07.123823Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.09269","last_updated":"2024-11-07T16:43:01Z","snapshot_observed_at":"2026-08-04T19:48:49.986548Z","submitted_at":"2024-02-14T15:55:30Z","title":"Personalized Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.09269","snapshot_observed_at":"2026-08-07T05:43:07.128902Z","title":"Personalized Large Language Models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.128902Z"},"links":{"cited_paper":"/paper/2402.09269","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:ee97c7148beb417d6182afa7c8df21129ed252216db6ebfa6cb351a3b9a78236","observation_id":"c5354f90-fad6-4bb2-96cf-a66ca53ebf2d","resolution":{"observed_at":"2026-08-07T05:43:07.128902Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.214393Z","title":"Towards an End -to-End Personal Fine -Tuning Framework for AI Value Alignment","venue":null,"work_id":"298d7b8a-51ae-4263-8c8e-ffe4a0adaecd","year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.133884Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:ec79bcf3d95e86e54976a67ee17354702ec9cb276cbaa0304ef2379032e49821","observation_id":"49e8007a-9ada-423f-8dbc-f1bc1335f848","resolution":{"observed_at":"2026-08-07T05:43:08.218849Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.200406Z","title":"Safer Agentic AI","venue":null,"work_id":"a47d4a23-5e7c-480c-a21c-e0e93c50984e","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.139059Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:4169ceac4cdcb1c480617adce39de2e47205d4fefcbc8d1e6dbfa70d3c0afac6","observation_id":"a1f78d6b-2a43-4697-b583-f52683d34c35","resolution":{"observed_at":"2026-08-07T05:43:08.204958Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.186337Z","title":"Introducing the Model Context Protocol","venue":null,"work_id":"9c952372-d8d1-49a3-8321-2ea5769348b5","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.143540Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:771ccf3ec5211124d0c7997f0ea37fb41d951941f3750d861e0fd20b73a25675","observation_id":"ab4fc287-4bcb-4053-b2df-466b9746a33d","resolution":{"observed_at":"2026-08-07T05:43:08.190940Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.172155Z","title":"Deep Reinforcement Learning from Human Preferences","venue":null,"work_id":"3319e636-628e-4f09-9e7c-ff0a59bf88c6","year":2017},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.147943Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:1963fecc4be4df99ec0495ce4b821dee7568552e89e13672384cfc3cf9b8c9a4","observation_id":"0b0b6157-c207-4c1b-b451-66d573f3884b","resolution":{"observed_at":"2026-08-07T05:43:08.176703Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.04175","last_updated":"2024-06-25T18:37:19Z","snapshot_observed_at":"2026-08-06T07:49:06.782482Z","submitted_at":"2024-06-06T15:32:29Z","title":"Confabulation: The Surprising Value of Large Language Model Hallucinations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.04175","snapshot_observed_at":"2026-08-07T05:43:07.152276Z","title":"Confabulation: The Surprising Value of Large Language Model Hallucinations","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.152276Z"},"links":{"cited_paper":"/paper/2406.04175","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:44586319aed738e3c2d67d7352417949b065428ebe1a51aa5e4d337549a02011","observation_id":"c1a6a662-4abe-47f4-b593-49c0d6ede551","resolution":{"observed_at":"2026-08-07T05:43:07.152276Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.157190Z","title":"Choice Vectors: Streamlining Personal AI Alignment Through Binary Selection","venue":null,"work_id":"1be5211c-f217-4e1a-a39c-d81b9e5497ca","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.157090Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:98fef6c41b0a0258efdfc65405991546592f74bba577bc227f8660769f59a65a","observation_id":"75f881bc-d2a2-4cc7-a983-2f2387095612","resolution":{"observed_at":"2026-08-07T05:43:08.161821Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.143284Z","title":null,"venue":null,"work_id":"e7267fa3-3ec8-4ed7-be5c-c3da057702c7","year":1923},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.161465Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:3c6713d3324e8ac101b0c45584d3de4a191c8989db4d44632f82b7ab807ffeca","observation_id":"f8843945-5b73-4fb6-8fdd-6eacf9e521a7","resolution":{"observed_at":"2026-08-07T05:43:08.147667Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.165649Z","title":"Universality of Representation in Biologic al and Artificial Neural Networks","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.165649Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:0531853e7e02bc5a824e76a9564aa36f569900503b2a440ef0ab395031bdf310","observation_id":"232768fa-1ba1-48b4-b37f-24a77dad24c0","resolution":{"observed_at":"2026-08-07T05:43:07.165649Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.129041Z","title":"The neural bases of cognitive conflict and control in mo ral judgment","venue":null,"work_id":"59b595d4-f6a1-4922-a7e4-bb7d2d224a63","year":2004},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.169921Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:ab755b1038726ed779c726f533aa90d891da2e26201964582c558945ac9f37be","observation_id":"b5015ad3-26fb-48f4-8a6f-ff1489e4c7a9","resolution":{"observed_at":"2026-08-07T05:43:08.133544Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.113746Z","title":"The neural basis of human social values: Evidence from functional MRI","venue":null,"work_id":"0115e97c-bb9f-4421-aab6-1dd1cbd9de5e","year":2009},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.174334Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:1a4be3402b2e6c13b399cc36cb3b53471ee68bee2dc942641efb188029e73b0f","observation_id":"5f808b93-0ce8-402b-a21e-fc5f0480cb73","resolution":{"observed_at":"2026-08-07T05:43:08.118460Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.098656Z","title":"A Cognitive Theory of Consciousness: The Workspace of the Mind; Cambridge University Press: Cambridge, UK, 1988","venue":null,"work_id":"eb5a53f3-d80a-4612-b3f0-3535426229c5","year":1988},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.179006Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:2ba499fb0d949e0a3f628ad85ab85510be95ab2bcc9e4820a1fc66cc8d3f72e9","observation_id":"2d831fb6-b9b9-40aa-af80-9648e02cea1c","resolution":{"observed_at":"2026-08-07T05:43:08.103481Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.082983Z","title":"Unified Theories of Cognition; Harvard University Press: Cambridge, MA, USA, 1990","venue":null,"work_id":"b2683094-f6bd-491d-b95c-631fdd1058c0","year":1990},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.183568Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:ee3bdc0a1ede23c46702093d6db6c0d9048abc928251d97f5eb8581c53d505ab","observation_id":"6886e3ee-5210-46b2-b28e-efef9e830027","resolution":{"observed_at":"2026-08-07T05:43:08.088209Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2505.08662","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.772777Z","title":"Revealing economic facts: LLMs know more than they say","venue":null,"work_id":"b29812ab-6be9-4477-95f1-2d525898babc","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.188289Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:ab2f04188ea59cc76c6e6e71bfc80f646beb02e233cee5e1a04dd7af8d7ec059","observation_id":"bec20e91-f154-4787-a276-9d29e7bc19d6","resolution":{"observed_at":"2026-08-07T05:43:07.780589Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.01081","last_updated":"2025-04-08T18:38:04Z","snapshot_observed_at":"2026-08-07T16:16:42.934012Z","submitted_at":"2025-04-01T18:00:20Z","title":"ShieldGemma 2: Robust and Tractable Image Content Moderation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.01081","snapshot_observed_at":"2026-08-07T05:43:07.193249Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.193249Z"},"links":{"cited_paper":"/paper/2504.01081","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:d7600a6311d9d121542732c3330c9ca35484f30e0de974763bb563461922d727","observation_id":"565c71ff-9edc-481f-a63f-4b03b833470d","resolution":{"observed_at":"2026-08-07T05:43:07.193249Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.068714Z","title":"Superego -Agent LGDemo (Branch: Fastapi_Mcp)","venue":null,"work_id":"cf36d267-18f0-4732-afff-e270a768737a","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.198921Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:6c4b8eb44ddc6d76cec4b7959fd7226d5a44fc60ccfb39b1ebffa50186726d5c","observation_id":"40c700b9-bd04-4265-a0d1-417a7d2802c3","resolution":{"observed_at":"2026-08-07T05:43:08.073012Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.04249","last_updated":"2024-02-27T04:43:08Z","snapshot_observed_at":"2026-07-06T17:26:23.067923Z","submitted_at":"2024-02-06T18:59:08Z","title":"HarmBench: A Standardized Evaluation Framework for Automated Red Teaming and Robust Refusal","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.04249","snapshot_observed_at":"2026-08-07T05:43:07.203855Z","title":"HarmBench: A standardized evaluation framework for automated red teaming and robust refusal","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.203855Z"},"links":{"cited_paper":"/paper/2402.04249","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:bd29590b4448291174740f99a91a64a02e552af510de36c4ec9adeaa11a4f9fa","observation_id":"c55222a1-896d-459a-9f3e-dfd48bfa9033","resolution":{"observed_at":"2026-08-07T05:43:07.203855Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.053714Z","title":"AgentHarm: A benchmark for measuring harmfulness of LLM agents","venue":null,"work_id":"0dc17e2f-f693-4623-968e-0e5a27f52f23","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.209304Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:27c85c39ef24195cbf32c759f39db7f2652c5e6104675b128eba247c97953be4","observation_id":"218dadac-60a5-4f85-909e-cf5e6a25d19e","resolution":{"observed_at":"2026-08-07T05:43:08.058305Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.038900Z","title":"Do the rewards justify the means? Measuring trade -offs between rewards and ethical behavior in the Machiavelli benchmark","venue":null,"work_id":"ca0c0869-c5b6-4046-95ff-417ba22d957e","year":2023},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.213631Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:fefd4df8634c623804b8770e600383744c8a5983d89a695d94111ac709af01dd","observation_id":"0d21dcae-dcab-4951-b7a8-6e0cf673c985","resolution":{"observed_at":"2026-08-07T05:43:08.043769Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.12272","last_updated":"2024-04-18T15:45:27Z","snapshot_observed_at":"2026-08-05T08:58:26.872416Z","submitted_at":"2024-04-18T15:45:27Z","title":"Who Validates the Validators? Aligning LLM-Assisted Evaluation of LLM Outputs with Human Preferences","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.12272","snapshot_observed_at":"2026-08-07T05:43:07.218428Z","title":"Who validates the validators? Aligning LLM-assisted evaluation of LLM outputs with human preferences","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.218428Z"},"links":{"cited_paper":"/paper/2404.12272","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:b88fa9e33c5407ffd5d1b0b4f67b50b4f3e0d63a7b45770490c6619b5fe9a975","observation_id":"7d230248-3583-4771-9434-10d2f6bcc2e5","resolution":{"observed_at":"2026-08-07T05:43:07.218428Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.024084Z","title":"Vijil Test Library: Evaluating LLM Trustworthiness Across Eight Dimensions","venue":null,"work_id":"6da1e6c9-9e9e-418f-8b96-d6c0462510c1","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.223161Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:370f3f7a295779e98e102d79483d2bc15354815ea51e8513877ed9573bfa53b9","observation_id":"b9a1bd80-b170-4952-ac58-bb35175f738c","resolution":{"observed_at":"2026-08-07T05:43:08.029110Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.009296Z","title":"INSPECT: An Extensible Toolkit for AI Behavior Evaluation","venue":null,"work_id":"a21e96d2-fc8f-490a-848c-69a4f04affb3","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.227370Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:bafebee85b6a967ce2c720e781b900b4ce75354963d646a30c98b5d41c1bbfdd","observation_id":"dd65809f-b519-4dd9-aca7-bc452cb6adae","resolution":{"observed_at":"2026-08-07T05:43:08.013655Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.994918Z","title":"Governance in Agentic Workflows: Leveraging LLMs as Oversight Agents","venue":null,"work_id":"b8b4767a-3323-41e3-99a6-276eec52e91e","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.231834Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:0cad9894e395512f72aef27efaddd36b20f30eb13683b41992c7f906dd854bd1","observation_id":"5c5cdae8-08c2-4c38-a138-b11566074059","resolution":{"observed_at":"2026-08-07T05:43:07.999467Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.08968","last_updated":"2025-03-03T22:10:04Z","snapshot_observed_at":"2026-08-04T22:28:55.011921Z","submitted_at":"2024-10-11T16:38:01Z","title":"Controllable Safety Alignment: Inference-Time Adaptation to Diverse Safety Requirements","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.08968","snapshot_observed_at":"2026-08-07T05:43:07.235972Z","title":"Controllable Safety Alignment: Inference-Time Adaptation to Diverse Safety Requirements","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.235972Z"},"links":{"cited_paper":"/paper/2410.08968","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:8d3822c9efdea63377dacae0bc7f3a74a96951813dcfc095d7a34176c1a59d53","observation_id":"f34152c1-3bdc-4796-b437-1ecc6eaf0dc8","resolution":{"observed_at":"2026-08-07T05:43:07.235972Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.00166","last_updated":"2026-04-30T20:54:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-30T19:11:52Z","title":"Disentangled Safety Adapters Enable Efficient Guardrails and Flexible Inference-Time Alignment","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.00166","snapshot_observed_at":"2026-08-07T05:43:07.240615Z","title":"Disentangled Safety Adapters Enable Efficient Guardrails and Flexible Inference-Time Alignment","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.240615Z"},"links":{"cited_paper":"/paper/2506.00166","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:9565cc2e445757f8b6cf3e4922d3709adc99397c5c9d04219cfbb922614f3237","observation_id":"c56c63fe-b2f0-40db-8b6f-319be09db8c4","resolution":{"observed_at":"2026-08-07T05:43:07.240615Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.11206","last_updated":"2024-01-20T10:41:03Z","snapshot_observed_at":"2026-08-06T01:29:18.626417Z","submitted_at":"2024-01-20T10:41:03Z","title":"InferAligner: Inference-Time Alignment for Harmlessness through Cross-Model Guidance","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.11206","snapshot_observed_at":"2026-08-07T05:43:07.245116Z","title":"InferAligner: Inference -Time Alignment for Harmlessness through Cross-Model Guidance","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.245116Z"},"links":{"cited_paper":"/paper/2401.11206","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:83eccd0dd4effe37bf7ad73f3cb2e9a95149ddc71991a118098308be005e8312","observation_id":"01f226e5-346c-4313-bf58-63314a39d76c","resolution":{"observed_at":"2026-08-07T05:43:07.245116Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01208","last_updated":"2025-06-20T10:54:05Z","snapshot_observed_at":"2026-08-08T13:31:24.117834Z","submitted_at":"2025-02-03T09:59:32Z","title":"On Almost Surely Safe Alignment of Large Language Models at Inference-Time","version":3},"cited_work":{"arxiv_id":"2502.01208","doi":null,"metadata_source":"pith","pith_arxiv_id":"2502.01208","snapshot_observed_at":"2026-08-07T05:43:07.622958Z","title":"On Almost Surely Safe Alignment of Large Language Models at Inference-Time","venue":"cs.LG","work_id":"0ad26be9-3465-4b1c-8504-3dccd9876074","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.249475Z"},"links":{"cited_paper":"/paper/2502.01208","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:d20e85bef4fc726e0144ac86ea2d4155fb52a9ed795b91914720978b0a460de1","observation_id":"f47fd653-d640-4acc-b755-e93a1f270232","resolution":{"observed_at":"2026-08-07T05:43:07.630186Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.02039","last_updated":"2025-06-02T22:37:06Z","snapshot_observed_at":"2026-08-07T23:06:33.439528Z","submitted_at":"2025-03-03T20:32:05Z","title":"Dynamic Search for Inference-Time Alignment in Diffusion Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.02039","snapshot_observed_at":"2026-08-07T05:43:07.253673Z","title":"Dynamic Search for Inference-Time Alignment in Diffusion Models (DSearch)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.253673Z"},"links":{"cited_paper":"/paper/2503.02039","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:f3c79977d0e3c6fa2f6c044bc8646d18fd401a5a608ff0a76934291bf2cc379d","observation_id":"535df10b-43a8-4d3e-afd6-35257bf5828b","resolution":{"observed_at":"2026-08-07T05:43:07.253673Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.18837","last_updated":"2025-01-31T01:09:32Z","snapshot_observed_at":"2026-07-06T20:28:47.519113Z","submitted_at":"2025-01-31T01:09:32Z","title":"Constitutional Classifiers: Defending against Universal Jailbreaks across Thousands of Hours of Red Teaming","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.18837","snapshot_observed_at":"2026-08-07T05:43:07.258162Z","title":"Constit utional Classifiers: Defending against Universal Jailbreaks across Thousands of Hours of Red Teaming","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.258162Z"},"links":{"cited_paper":"/paper/2501.18837","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:8d41976cb96366aaa68515352150b0fdcf529b7359dbca69480315517d14682f","observation_id":"ad083c49-e7a3-48d4-8c78-de4e9196656b","resolution":{"observed_at":"2026-08-07T05:43:07.258162Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2212.08073","last_updated":"2022-12-15T06:19:23Z","snapshot_observed_at":"2026-08-02T04:53:58.766070Z","submitted_at":"2022-12-15T06:19:23Z","title":"Constitutional AI: Harmlessness from AI Feedback","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2212.08073","snapshot_observed_at":"2026-08-07T05:43:07.262772Z","title":"Constitutional AI: Harmlessness from AI Feedback","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.262772Z"},"links":{"cited_paper":"/paper/2212.08073","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:d53878783b789272ea848c6b9cb2d6a0d71d4541fbe37fffc1022904fb277279","observation_id":"b7f7c262-ade5-4959-a0a1-84e014de2c97","resolution":{"observed_at":"2026-08-07T05:43:07.262772Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2204.05862","last_updated":"2022-04-12T15:02:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-04-12T15:02:38Z","title":"Training a Helpful and Harmless Assistant with Reinforcement Learning from Human Feedback","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2204.05862","snapshot_observed_at":"2026-08-07T05:43:07.267381Z","title":"T raining a Helpful and Harmless Assistant with Reinforcement Learning from Human Feedback","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.267381Z"},"links":{"cited_paper":"/paper/2204.05862","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:7c33cb19d910f2d71536427b0fdd171d8372a13a5851e36128e5801ffbf97e11","observation_id":"7ae553b9-9542-4f09-986a-e2bb36766dc6","resolution":{"observed_at":"2026-08-07T05:43:07.267381Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.980210Z","title":"Monitoring Reasoning Models for Misbehavior and the Risks of Promoting Obfuscation","venue":null,"work_id":"81d3f66c-2d50-4fc0-85ca-76f2cdf707bc","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.272296Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:4de4533ea5cf00027ef6650d9eb93460895b76542933300b2392d52e7e122bdc","observation_id":"0307194a-196c-4c9a-adca-c6c142cca4c6","resolution":{"observed_at":"2026-08-07T05:43:07.984868Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.06942","last_updated":"2024-07-23T06:47:13Z","snapshot_observed_at":"2026-08-06T12:56:48.193649Z","submitted_at":"2023-12-12T02:34:06Z","title":"AI Control: Improving Safety Despite Intentional Subversion","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.06942","snapshot_observed_at":"2026-08-07T05:43:07.278246Z","title":"AI Control: Improving Safety Despite Intentional Subversion","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.278246Z"},"links":{"cited_paper":"/paper/2312.06942","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:22f4a0e296364097196bf338fc069b38be3b02887efb2042cc2585020d467655","observation_id":"171bc1d9-efa4-4d45-85cc-b610962d7f10","resolution":{"observed_at":"2026-08-07T05:43:07.278246Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.965984Z","title":"OpenAI x DFT: The First Moral Graph","venue":null,"work_id":"a49e7d4e-9ee5-412b-b110-1a33995fc6d8","year":2023},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.283191Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:7a8f4e686fe85356972f9b596145a44cfe6131af3e4d298a9198d3a994deff17","observation_id":"f222863a-22fb-447c-badb-4127e9498e77","resolution":{"observed_at":"2026-08-07T05:43:07.970374Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.951344Z","title":"Model Integrity","venue":null,"work_id":"e06c8563-8feb-4a7e-a507-c8e5dc561609","year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.287785Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:bfa919bd42a21d6dde17817ceaf5ef5416c5c48ba2d0b9f56f45e37da11c6787","observation_id":"8ff4d278-26db-4e2f-a2be-693040d26acf","resolution":{"observed_at":"2026-08-07T05:43:07.956049Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.936583Z","title":"The Global Landscape of AI Ethics Guidelines","venue":null,"work_id":"c36e5993-a8bf-47d8-90ea-2e903e4f3069","year":2019},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.292294Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:191f1b1beb34f77d843cc1a338de1ee23190756475f5ad0bb401dfa119b93f16","observation_id":"27bab82c-a0dc-4c0c-ad58-3d049b54f333","resolution":{"observed_at":"2026-08-07T05:43:07.941746Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.921150Z","title":"WhatsApp MCP Exploited: Exfiltrating Your Message History via MCP","venue":null,"work_id":"12f67775-f8ae-4b0e-a319-9ca9c3cae0f0","year":null},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.296584Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:711a2844b438d59c5f19286063cd383931c97c10ac380437a8e1c4248e87fc20","observation_id":"5d78eeaf-2263-4a86-9a54-c9b9b492120d","resolution":{"observed_at":"2026-08-07T05:43:07.926146Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.890964Z","title":"MCP Security Notification: Tool Poisoning Attacks","venue":null,"work_id":"44f1ce5f-dd10-41da-b7f7-e1ed40a33143","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.306568Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:6f6407760dc0861343466a63a0fe9248e42866dacc0f6d5a431eace0d3293655","observation_id":"48e2fdca-9791-46d1-a005-6bdfaca00cf9","resolution":{"observed_at":"2026-08-07T05:43:07.895246Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.311471Z","title":"Emergent misalignment: Narrow finetuning can produce broadly misaligned LLMs","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.311471Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:89e21ca2d7f94cddefe6ed34d6be80f213a6bb9753d2d643d3d94cdb4a69ebb7","observation_id":"20244967-148d-46a2-bb57-cdfbdc0be7ec","resolution":{"observed_at":"2026-08-07T05:43:07.311471Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.875882Z","title":"On Emergent Misalignment","venue":null,"work_id":"fd1df4c5-7e8b-4d6e-aa4a-9710dce04976","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.315617Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:0af311189b57859f53206a0b3182616830930d530e32c9d758f5ad9340f98124","observation_id":"31272682-141c-4702-a791-551e22302342","resolution":{"observed_at":"2026-08-07T05:43:07.880226Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.860100Z","title":"Model Plurality","venue":null,"work_id":"18089c75-6a9a-4013-842e-0cc5d2f5d45d","year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.319836Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:d3755a5fa13c313c95624ab994bd467e7b35917f23cb12c75978494ea06e48b4","observation_id":"5f40ad55-bb11-4b30-bc91-b86935cea64a","resolution":{"observed_at":"2026-08-07T05:43:07.864809Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.844679Z","title":"Model Plurality: A Taxonomy for Pluralistic AI","venue":null,"work_id":"3806dc4b-430c-4fb0-9a78-b552894fc032","year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.324191Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:c48019f8225a80d38fdfc15a15dde578c54f69671574a1e1efdc2a465af7304e","observation_id":"37621a40-2612-4ca3-812b-e77b1cdf5562","resolution":{"observed_at":"2026-08-07T05:43:07.849749Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.17950","last_updated":"2025-04-24T21:28:16Z","snapshot_observed_at":"2026-08-07T15:59:22.987847Z","submitted_at":"2025-04-24T21:28:16Z","title":"Collaborating Action by Action: A Multi-agent LLM Framework for Embodied Reasoning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.17950","snapshot_observed_at":"2026-08-07T05:43:07.328340Z","title":"Collaborating Action by Action: A Multi-agent LLM Framework for Embodied Reasoning","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.328340Z"},"links":{"cited_paper":"/paper/2504.17950","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:10793edbc3226c0fff9d4331cc1d79ed9bf7ec9212a2f0e74db9f76b3815176f","observation_id":"6f6e36b4-7098-4f6c-9730-c2637ec23bd0","resolution":{"observed_at":"2026-08-07T05:43:07.328340Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.11581","last_updated":"2025-03-23T13:00:03Z","snapshot_observed_at":"2026-08-02T06:52:32.786538Z","submitted_at":"2024-11-18T13:57:35Z","title":"OASIS: Open Agent Social Interaction Simulations with One Million Agents","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.11581","snapshot_observed_at":"2026-08-07T05:43:07.332394Z","title":"OASIS: Open Agent Social Interaction Simulations with One Million Agents","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.332394Z"},"links":{"cited_paper":"/paper/2411.11581","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:69ca7686fae471dab3cffc8fd291b10a1060ae3535f509f6911ba13aee391b6a","observation_id":"dc7ee6ef-2dca-4126-922c-1c188a5c2c2e","resolution":{"observed_at":"2026-08-07T05:43:07.332394Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.00114","last_updated":"2024-10-31T18:11:22Z","snapshot_observed_at":"2026-08-08T05:37:29.265864Z","submitted_at":"2024-10-31T18:11:22Z","title":"Project Sid: Many-agent simulations toward AI civilization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.00114","snapshot_observed_at":"2026-08-07T05:43:07.337382Z","title":"Project Sid: Many- agent simulations toward AI civilization","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.337382Z"},"links":{"cited_paper":"/paper/2411.00114","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:a49e6eeea1370a2374fa361aca712735fdc26b7d93f178de9cdba358c4466699","observation_id":"1b915a17-c0ea-45d4-93aa-e2fcfdfb4806","resolution":{"observed_at":"2026-08-07T05:43:07.337382Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.905785Z","title":null,"venue":null,"work_id":"43b30316-97dc-4362-a2f2-5b1c307d160d","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.301344Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:3766f5b843fa10fc57a8a32ed9a52ac157960e7d942bf33d45d7bd0f1c2a8d76","observation_id":"106163d8-0eb8-450e-9b6f-081131884779","resolution":{"observed_at":"2026-08-07T05:43:07.910198Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","latest_version":2,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-08T03:32:46.630570Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values"},"reference_resolution":{"displayed":50,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":21,"verified_exact":2,"verified_fuzzy":27},"total_outbound_references":50},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 50 of 50 outbound references and 0 inbound Pith citation observations for arXiv:2506.13774."}