{"as_of":"2026-08-20T05:57:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2f781f3c92e3231649e53c88b996ec20f411b67ed939187572d850d19873c307","coverage":[{"denominator":37,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":37,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T10:44:44.310314Z","state":"measured"},{"denominator":38,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":38,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-19T06:32:44.657259+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T11:35:11.412956Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-16T11:35:12.372679Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"cited_work":{"arxiv_id":"2412.16325","doi":null,"metadata_source":"pith","pith_arxiv_id":"2412.16325","snapshot_observed_at":"2026-08-16T11:35:12.372679Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","venue":"cs.AI","work_id":"81c66582-fdc4-4c20-afa2-8e076b46c87e","year":2024},"citing_paper":{"arxiv_id":"2504.15125","last_updated":"2025-08-18T10:09:08Z","snapshot_observed_at":"2026-08-18T17:48:11.814383Z","submitted_at":"2025-04-21T14:20:49Z","title":"Contemplative Artificial Intelligence","version":3},"reference_index":506,"source":"pdf_text","source_observed_at":"2026-08-16T11:35:11.412956Z"},"links":{"cited_paper":"/paper/2412.16325","citing_paper":"/paper/2504.15125"},"observation_digest":"sha256:5f854c5a1015c34163345d32260f322a1986eb3de9682542ab30341581d67578","observation_id":"64a96ae2-078d-45b4-b7f8-2ba2f1f6dd26","resolution":{"observed_at":"2026-08-16T11:35:12.376510Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2412.16325/citation-record","integrity":"/paper/2412.16325/integrity","json":"/paper/2412.16325/citation-record.json","paper":"/paper/2412.16325"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.446491Z","title":"Unsolved problems in ml safety","venue":null,"work_id":"b2877b67-6688-4dc1-a008-07a6e765e9cf","year":2021},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:43.944126Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:7e76d08111a2ddbd2c87202e3a388a5e5c437efe9fe897e43d998eebae40faa0","observation_id":"b137f163-56a9-4707-82a7-a87ba70b3d47","resolution":{"observed_at":"2026-08-11T10:44:45.452387Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.410262Z","title":"Toward trustworthy ai development: Mechanisms for supporting verifiable claims","venue":null,"work_id":"76bb2a6e-060b-41a7-b0fb-2c356d77fc2d","year":2020},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:43.955285Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:efb70629f6af45fc358253f1e971fc3aa2814375508e28fc6cab4c08e172c7cd","observation_id":"48444c6d-dd7f-4327-a6b8-ece2bbb5690c","resolution":{"observed_at":"2026-08-11T10:44:45.419072Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.382070Z","title":"Deception analysis with artificial intelligence: An interdisciplinary perspective","venue":null,"work_id":"3f1f2a2c-4351-4fea-8899-fa8432927040","year":2024},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:43.968120Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:1d7806a4cad22f8ef2a6720541d6bd22220e6b416ae25ee8e307e9d8c3248486","observation_id":"5238e710-8330-473f-94cf-75ecb29db56f","resolution":{"observed_at":"2026-08-11T10:44:45.389584Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.359890Z","title":"Unmasking the shadows of ai: Investigating deceptive capabilities in large language models","venue":null,"work_id":"4870986d-211d-46c5-ae89-50c678229ac9","year":2024},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:43.977594Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:58764f889383da50038ae323858b78786dd9168e243db07519d9d8ce06d2505b","observation_id":"17402991-8706-4046-a5e1-330864d132cb","resolution":{"observed_at":"2026-08-11T10:44:45.365826Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.340901Z","title":"Human-level play in the game of diplomacy by combining language models with strategic reasoning","venue":null,"work_id":"9d6432f5-26cd-42d5-b7b5-a1767509ce87","year":2022},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:43.984689Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:2d4bca95e8ea7920e14c9c1cdfc2fab1646153e2a3be02a164c5a783028f1146","observation_id":"9de0a34c-9312-4e5a-926b-2b7024e769dd","resolution":{"observed_at":"2026-08-11T10:44:45.346769Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.316372Z","title":"Hendricks, M.R","venue":null,"work_id":"5131fe6d-ea0e-4de0-b778-d23d689df9b9","year":2023},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:43.991644Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:6aa93e784ea4740fea5ff5702f0fa74bdc9aa68d29a823407febe29b24217bb7","observation_id":"40e46928-84a1-43a1-b842-8bc1206a3565","resolution":{"observed_at":"2026-08-11T10:44:45.323249Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.296494Z","title":"Collective constitutional ai: Aligning a language model with public input","venue":null,"work_id":"52eca17c-fe48-4c6b-bb77-e173f9242f04","year":2021},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:43.998339Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:f6161e442fb07430ed341b7e25a505db6ee800fc1f86a0f9cd7c36288aca7da2","observation_id":"e7614278-62d3-4b53-9339-8f2fdf2c71f5","resolution":{"observed_at":"2026-08-11T10:44:45.302447Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.270577Z","title":"Constitutional ai: Harmlessness from ai feedback","venue":null,"work_id":"dfb9cd2a-da46-4b9e-a6db-87b29543e35c","year":2022},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.006696Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:2b3907953865e54eb0b27ff913b362df154e38c508b2ca525a494fbf2b6e4758","observation_id":"14f108e6-a78e-4dc9-b621-26c06a763c70","resolution":{"observed_at":"2026-08-11T10:44:45.279057Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.240943Z","title":"Truthful ai: Developing and governing ai that does not lie","venue":null,"work_id":"87b3b4e2-e254-4434-a81d-9a0ab377f4c3","year":2021},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.015463Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:e5209c1e18535669524521afbf7b4ab285d707f43ae3ba856ee46514d4646d3a","observation_id":"9e72724f-ee72-4532-a3bf-77e75543b47e","resolution":{"observed_at":"2026-08-11T10:44:45.250357Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.201905Z","title":"Language models represent beliefs of self and others","venue":null,"work_id":"8cda6187-5fbe-486e-bd40-cfa1a4fc3002","year":2024},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.023399Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:bb8e053cd7788135119ccebe36bab2833e49ce71cece63b382397b1bddaad766","observation_id":"36de0d67-2e0f-4dcf-a358-0a91332d7a02","resolution":{"observed_at":"2026-08-11T10:44:45.211940Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.174650Z","title":"Premakumar, Michael Vaiana, Florin Pop, Judd Rosenblatt, Diogo Schwerz de Lucena, Kirsten Ziman, and Michael S","venue":null,"work_id":"8a4e17f2-a758-4370-abd9-fff81daaf154","year":2024},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.031285Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:c60e5469865ebbc4b8123bd44304c1641633d3a1d464db6b6b828b8c6d94c69b","observation_id":"d3dc3d51-5a6d-4ffd-9a91-64bd9cba8d47","resolution":{"observed_at":"2026-08-11T10:44:45.184641Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.139397Z","title":"Predicting vs","venue":null,"work_id":"d4912276-9633-463a-8593-41ba85afd2f7","year":2024},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.040063Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:f0109e88b4a1d42a35473693aa19a1239d8b7554d9b180edb66b2da1974c9943","observation_id":"ff581c8b-40d4-421b-a9f1-d956110cc66a","resolution":{"observed_at":"2026-08-11T10:44:45.150329Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.103385Z","title":null,"venue":null,"work_id":"bba3dbce-39b0-42db-acb9-6c542cf6f48f","year":2017},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.047169Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:a64a2609593b64b88ffb415b0a2ad70ee44f2ff88caefbc79010724bde3813fe","observation_id":"7e01b7c9-552a-4324-8a96-967aaa4617c5","resolution":{"observed_at":"2026-08-11T10:44:45.113038Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.068578Z","title":"Brethel-Haurwitz, Elise M","venue":null,"work_id":"d6eb98b4-4e3d-4db1-8365-8160071d1961","year":2018},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.056293Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:188d09fd4683ed31bff071b3cce538ed75f131584673deb844ad1d2aae5089ac","observation_id":"7294465e-5088-4713-9a1a-950442ea7f3f","resolution":{"observed_at":"2026-08-11T10:44:45.077331Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.036994Z","title":"Brethel-Haurwitz, Elise M","venue":null,"work_id":"821ca3c9-ac80-4bbe-9149-432fbe0654c8","year":2019},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.067209Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:5919c2abf42a6903536aa31bac1beddc2cd43232f29e29aa3933275c11aae650","observation_id":"9697ab3b-33d0-4106-b3e7-9d92fd04f44e","resolution":{"observed_at":"2026-08-11T10:44:45.045415Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:45.000677Z","title":"Do altruists lie less? Journal of Economic Behavior & Organization, 157:560–579, 2019","venue":null,"work_id":"57945a5d-c442-422c-9f49-016d91e43084","year":2019},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.079318Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:878094d2779ecb7ab5b4c485e301aa9c3cf00c043b8766e749b0331e862265b3","observation_id":"2cd2f173-7e91-4d50-9852-68eaf49cec22","resolution":{"observed_at":"2026-08-11T10:44:45.010656Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.963693Z","title":"O’Connell, Shawn A","venue":null,"work_id":"43a65597-c184-4b6c-8abd-c0d08aace1f9","year":2020},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.093737Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:55eedc0b92d5de386db9e2fda62d85ab3bb75927e3b8ed135e6373573ce4f9af","observation_id":"9270e5fe-131a-4111-99f3-7de25208f163","resolution":{"observed_at":"2026-08-11T10:44:44.975134Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.929670Z","title":null,"venue":null,"work_id":"518dc09d-88f6-478d-a2f6-647744aac751","year":2013},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.106615Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:8275e1d274f2f45d584dcc96381dc30683129af970d8b3be5529ec0a5f6948e7","observation_id":"5d9e9d4c-2ef2-4b3b-a71f-3c7fcb8d5f3e","resolution":{"observed_at":"2026-08-11T10:44:44.936691Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.905675Z","title":"Jonason, Minna Lyons, Holly M","venue":null,"work_id":"9af7fdb2-d325-4fbd-9037-7805c96ae724","year":2014},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.120476Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:d1c4f21153efb76a3dbafbff5ad331c7c9ed417207e268a61bcf32a69b9bea64","observation_id":"6bb15a6f-7c1f-4d01-ac4f-121e1c4b12c9","resolution":{"observed_at":"2026-08-11T10:44:44.912098Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.875874Z","title":"Towards empathic deep q-learning","venue":null,"work_id":"46be4dec-0a8d-46bd-a97e-c15db255702e","year":2019},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.131620Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:a2206976717a6a2491fc8dd42f0402f737a44c20eaa26f7a1b6da44e9dff8fab","observation_id":"386ceeb9-3847-4a9e-8ad8-43030e1d798e","resolution":{"observed_at":"2026-08-11T10:44:44.883165Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.835073Z","title":"Modeling others using oneself in multi-agent reinforcement learning","venue":null,"work_id":"b0c3d6c5-82ff-4cda-aa06-13483e698f31","year":2018},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.145499Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:c116bebb78778c76fdc02db535aeffacd657b9dc79f568df44b39b35ad54ac6b","observation_id":"2bcabce1-db1d-422a-b3d7-290a2bd7e52b","resolution":{"observed_at":"2026-08-11T10:44:44.848221Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.800775Z","title":"Zou, Colin Raffel, Chris Callison-Burch, Yan Cao, Dzmitry Bahdanau, Gregory Diamos, and Jacob Steinhardt","venue":null,"work_id":"f23d393a-f074-4ff0-bfec-bce758876b4e","year":2023},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.156539Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:068fb7270a9f4660d5658fcca2d53e041b3f43c6c3dab3dbbfe91928021f8aaf","observation_id":"ab1e7bda-58a3-48f1-9c74-df8bccc7fb8e","resolution":{"observed_at":"2026-08-11T10:44:44.813247Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.758843Z","title":"Path-specific objectives for safer agent incentives","venue":null,"work_id":"d03f5f2d-a4af-4257-ab15-461fa204b793","year":2022},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.169401Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:1a1593c2e23f509b8a6410a1485b67fe7cf82b152519c7f1f99385b24974ea69","observation_id":"312a6920-8674-41d1-a3f9-9d86e988c482","resolution":{"observed_at":"2026-08-11T10:44:44.772250Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.717302Z","title":"Ortega, Elizabeth Barnes, and Shane Legg","venue":null,"work_id":"1c3b01a5-00c5-4396-b0b9-487f591db04d","year":2019},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.184667Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:372458f18eb6cacec2840c1ef900ae213f51955dc5488cd72f5253fcc708b399","observation_id":"51aa5b73-1ead-4862-8ab3-b26d7c3b3d51","resolution":{"observed_at":"2026-08-11T10:44:44.726973Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.681320Z","title":"Honesty is the best policy: Defining and mitigating ai deception","venue":null,"work_id":"541932a0-c1d4-467b-b85e-a419c8da97fd","year":2023},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.194963Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:770e51583ab03fcce493d452f4e7e137a8be8e079c5c8ba91c66bec0f35ee6a2","observation_id":"bc973fb2-c6b5-45cc-8d2c-81b2914c8c86","resolution":{"observed_at":"2026-08-11T10:44:44.689673Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.651551Z","title":"The history and risks of reinforcement learning and human feedback","venue":null,"work_id":"e7ff61fb-78c4-4e3f-988a-d1bb831d211d","year":2022},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.204251Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:f8eca9c43998f03c8f29896fc0b34e7d96b0c856ead9e6a440051504a58d7d1a","observation_id":"d1e973be-c953-4d8b-961b-8dfa99b0c82c","resolution":{"observed_at":"2026-08-11T10:44:44.662316Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.628250Z","title":"Deception abilities emerged in large language models","venue":null,"work_id":"9cba5afe-ae7c-4957-bbc3-7211b97e6545","year":2024},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.215819Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:0a3852cdd60c02845500b78ef9837cb803b3b0244b7b9ed584bd2cb548adbfdd","observation_id":"5954792c-479b-4402-b0a2-a0a36e484fb0","resolution":{"observed_at":"2026-08-11T10:44:44.635817Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.605192Z","title":"Xing, Hao Zhang, Joseph E","venue":null,"work_id":"c65afa38-8cce-4e67-be67-308f85b6af55","year":2023},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.228992Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:58e0dc5cc9e1c21aaee0269ad2f9f48108524c7ad5b7ac74712a8b52a7863968","observation_id":"4f6b15b4-5811-40ec-b669-2a3cb30c1726","resolution":{"observed_at":"2026-08-11T10:44:44.611128Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.576634Z","title":"Physical-deception: An implementation of multi-agent deep deterministic policy gradient in pytorch to solve the physical deception environment from openai, 2023","venue":null,"work_id":"8aa028df-660a-458d-814a-d36fea31d09e","year":2023},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.240552Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:b8d03b9fb98a41899e16864a3879df5e129bd7f54268674aeb1d619c225978d4","observation_id":"8ec5ae5c-d92c-4b1d-97ae-c186b2a92284","resolution":{"observed_at":"2026-08-11T10:44:44.585522Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.555076Z","title":"Multi-agent actor-critic for mixed cooperative-competitive environments","venue":null,"work_id":"4b1dd96e-e392-44bc-b72d-157190d31bfc","year":2017},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.251103Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:66293f0ea1502eb633cc8da8bf3e7a4609babdef14ea571db906878f586e3b26","observation_id":"bfcaeb9b-4a37-4fa4-bc8a-bdc87e678ef1","resolution":{"observed_at":"2026-08-11T10:44:44.561180Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.257213Z","title":"Human Compatible: Artificial Intelligence and the Problem of Control","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.257213Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:d74e745aed63287e8a1d02d47dbf13378d1a79411d2e4e3fb42e1d71f9727f53","observation_id":"5dfa4582-d139-4fe9-90ca-2106764c518b","resolution":{"observed_at":"2026-08-11T10:44:44.257213Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.503152Z","title":"Risks from learned optimization in advanced machine learning systems","venue":null,"work_id":"0a96f75b-542b-4bf2-8541-5cf39264b4bc","year":2019},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.264947Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:9f701acb20dc114e7f878ed8d61805f179947dd69d7c6cea3a216b61d3492164","observation_id":"a06c1aae-01bf-4c97-9480-4a8a436d9615","resolution":{"observed_at":"2026-08-11T10:44:44.510954Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.479285Z","title":"Training a helpful and harmless assistant with reinforcement learning from human feedback","venue":null,"work_id":"d813873a-3e2b-4308-bd7c-2600f7e2cc21","year":2022},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.274554Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:c97ac4f14f7bc714e956ffffba64e847aa32daf9a34abaf21d8aa1dc4aeaba13","observation_id":"de3ef56c-131a-44ec-9dc8-2079ec1f7e83","resolution":{"observed_at":"2026-08-11T10:44:44.486557Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.447718Z","title":"Ziegler, Tim Maxwell, Newton Cheng, et al","venue":null,"work_id":"a13199b6-3b03-426b-99d8-46a2a793bcab","year":2024},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.287167Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:e71ef5531ddd48e10e00166573ad55476a79afa72b3611220a811bd6296e60e3","observation_id":"57346059-457a-4b6c-b1ba-8b23c2432b3a","resolution":{"observed_at":"2026-08-11T10:44:44.456945Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.424337Z","title":null,"venue":null,"work_id":"72bf511b-add9-43e2-a2fb-57880f9bdb98","year":2016},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.292880Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:954d7a956a2d0742b9de606b81417ed66ad4525cb1c77e1c5ffd4d2530f422c2","observation_id":"f182b524-7d34-4a95-8516-bc3dc2c32418","resolution":{"observed_at":"2026-08-11T10:44:44.430758Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.397057Z","title":"Chain of thought prompting elicits reasoning in large language models","venue":null,"work_id":"3132276d-2bb1-4799-9140-43507111f9d1","year":2022},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.300573Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:0150dd2f26d478537bcd9116fe0bb9ed1a9cfd41db3c35f0435f14e2d3a6cf70","observation_id":"c995afbc-b5ea-47a3-b2a6-53afb6342ed8","resolution":{"observed_at":"2026-08-11T10:44:44.403307Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:44:44.369541Z","title":"Only respond with the room name, no other text","venue":null,"work_id":"e4e324e2-2ba4-429d-89b8-ff698c0fdf18","year":2023},"citing_paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-11T10:44:44.310314Z"},"links":{"citing_paper":"/paper/2412.16325"},"observation_digest":"sha256:8e9929597d781452d509c66a35051c9d45a16b10a3b9a294be7297b828c42d32","observation_id":"7cd99554-e09e-404d-8e41-8a34ea973aae","resolution":{"observed_at":"2026-08-11T10:44:44.379471Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2412.16325","last_updated":"2024-12-20T20:23:52Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-17T23:35:13.598113Z","submitted_at":"2024-12-20T20:23:52Z","title":"Towards Safe and Honest AI Agents with Neural Self-Other Overlap"},"reference_resolution":{"displayed":37,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":4,"verified_exact":0,"verified_fuzzy":33},"total_outbound_references":37},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"thesis":"As of 20 August 2026, this Paper Citation Record lists 37 of 37 outbound references and 1 inbound Pith citation observation for arXiv:2412.16325."}