{"as_of":"2026-07-28T20:18:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:6ebe5d5e5a1982e1eb0fb1c5f8c68b5d5a9038b34a0614f335be8be70d80fa1b","coverage":[{"denominator":42,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":42,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-17T15:01:41.161487Z","state":"measured"},{"denominator":89,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":89,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-07-28T06:31:03.373048+00:00","state":"measured"},{"denominator":47,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":47,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-15T12:59:04.484292Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-10T18:37:31.166799Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2307.13702","last_updated":"2023-07-17T01:08:39Z","snapshot_observed_at":"2026-07-06T15:58:27.523039Z","submitted_at":"2023-07-17T01:08:39Z","title":"Measuring Faithfulness in Chain-of-Thought Reasoning","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-11T20:51:38.390234Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2307.13702"},"observation_digest":"sha256:23e0fe22358875d337466e2a91b7764f888d05341520c4ecb9c1885cc829edb1","observation_id":"310a7225-23d5-403c-b16f-07a60dbc2a9a","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2308.03958","last_updated":"2024-02-15T01:03:13Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-08-07T23:48:36Z","title":"Simple synthetic data reduces sycophancy in large language models","version":2},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-05-16T14:48:08.508109Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2308.03958"},"observation_digest":"sha256:689edeec61e861a1c997132164c4cf4372826baee6a2c80f3926cb21f7eee515","observation_id":"727345d7-2b3b-41a1-9510-ab5cc62c3528","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2310.10631","last_updated":"2024-03-15T19:14:39Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-16T17:54:07Z","title":"Llemma: An Open Language Model For Mathematics","version":3},"reference_index":129,"source":"arxiv_source","source_observed_at":"2026-05-19T08:17:46.055279Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2310.10631"},"observation_digest":"sha256:7bbb9e448b9506701df0bd4d4b53d0f09cb89c6cb0953de0c7ef13e197566935","observation_id":"677afb01-6fb3-4c8a-9183-1496bbf70eb3","resolution":{"observed_at":"2026-05-19T08:17:46.335169Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:3b28a11248d9646af58da430065c39dc6cce686c07b93eec723e0b1a13f71570","observation_id":"b8b401bf-361d-43d4-993c-33504914f439","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2311.05232","last_updated":"2024-11-19T12:42:45Z","snapshot_observed_at":"2026-07-06T16:45:07.733095Z","submitted_at":"2023-11-09T09:25:37Z","title":"A Survey on Hallucination in Large Language Models: Principles, Taxonomy, Challenges, and Open Questions","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-13T02:46:26.957539Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2311.05232"},"observation_digest":"sha256:e5041b4cd330aa8d331e00ec0985759f7b7270fa7aba91482f9380cc2c6eac64","observation_id":"e7055ce4-1ef6-49b0-a8c9-a3b6dba9cd98","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2401.05561","last_updated":"2024-09-30T10:17:12Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-01-10T22:07:21Z","title":"TrustLLM: Trustworthiness in Large Language Models","version":6},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-05-18T11:17:08.108565Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2401.05561"},"observation_digest":"sha256:728f145fd8f8bcb3553fdbb8adabe5bca02cefad7bebf2f73fa5ed9038b7b8ae","observation_id":"8601443d-1632-4d69-907c-c0eeff835971","resolution":{"observed_at":"2026-05-18T11:17:08.347452Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2402.05070","last_updated":"2024-08-20T19:14:31Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-02-07T18:21:17Z","title":"A Roadmap to Pluralistic Alignment","version":3},"reference_index":231,"source":"arxiv_source","source_observed_at":"2026-05-16T14:37:53.279275Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2402.05070"},"observation_digest":"sha256:a87000a50585461f803ccde014285f9b4cb67b47b8be55b873b7acc22ea608d9","observation_id":"901d418f-491f-47db-8464-82980c25dbad","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2406.10162","last_updated":"2024-06-29T00:28:47Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-14T16:26:20Z","title":"Sycophancy to Subterfuge: Investigating Reward-Tampering in Large Language Models","version":3},"reference_index":108,"source":"arxiv_source","source_observed_at":"2026-05-17T14:43:29.496457Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2406.10162"},"observation_digest":"sha256:ec39d82738eab8f35b31b8a66c34b98675b93dabdd714fc46512f70fbf0580a2","observation_id":"2223dc72-1d69-4bda-a082-004977665a4a","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2408.00724","last_updated":"2025-03-03T07:53:32Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-08-01T17:16:04Z","title":"Inference Scaling Laws: An Empirical Analysis of Compute-Optimal Inference for Problem-Solving with Language Models","version":3},"reference_index":267,"source":"arxiv_source","source_observed_at":"2026-05-18T06:38:36.517935Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2408.00724"},"observation_digest":"sha256:57ee108dbda80fedb67fb94e2695bd6b7fcdaa3980095d1233a78d928c13c822","observation_id":"3e608c7d-1aef-462d-9636-e631bea76776","resolution":{"observed_at":"2026-05-18T06:38:37.090656Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2408.12935","last_updated":"2026-05-13T07:56:42Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-08-23T09:33:48Z","title":"AI Safety Landscape for Large Language Models: Taxonomy, State-of-the-art, and Future Directions","version":4},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-05-23T21:54:26.670284Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2408.12935"},"observation_digest":"sha256:30584b8d7becdb30a003a6dbf5797c7323bfc072c2966a5dd2d6f723a4927104","observation_id":"bb63dc13-39a3-4e88-93b4-b7a2e01e40bf","resolution":{"observed_at":"2026-05-23T21:55:50.388446Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2503.11926","last_updated":"2025-03-14T23:50:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-03-14T23:50:34Z","title":"Monitoring Reasoning Models for Misbehavior and the Risks of Promoting Obfuscation","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-21T07:24:12.845841Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2503.11926"},"observation_digest":"sha256:c083314cec1783bf356f6b8a069ca226714f8ec2c4542237801a867eccb8f104","observation_id":"9e5401f7-c3a6-44e1-b232-ae7e53c35dca","resolution":{"observed_at":"2026-05-21T07:24:12.970748Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2506.06414","last_updated":"2026-04-20T20:58:30Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-06-06T17:33:33Z","title":"Benchmarking Misuse Mitigation Against Covert Adversaries","version":2},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-05-19T10:29:05.104520Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2506.06414"},"observation_digest":"sha256:801c0fd86ede444d08aa5d0aaf25e46a26114dfccedf7be63e6c7d396d417e86","observation_id":"48e7a18c-07af-44b3-b642-2af018737cc0","resolution":{"observed_at":"2026-05-19T10:32:14.636622Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2510.07239","last_updated":"2026-05-17T13:01:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-10-08T17:06:20Z","title":"Red-Bandit: Test-Time Adaptation for LLM Red-Teaming via Bandit-Guided LoRA Experts","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-21T20:53:58.198974Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2510.07239"},"observation_digest":"sha256:de28655f8c3c2e94566533184311498b926811f9ae6c0d6bc6743f5dbe2ab548","observation_id":"3be9b6ac-ece3-479a-b547-8d6ba6fe8cfc","resolution":{"observed_at":"2026-05-21T20:54:21.555096Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-15T12:59:04.484292Z","title":"Measuring progress on scalable oversight for large language models.arXiv preprint arXiv:2211.03540,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2603.13356","last_updated":"2026-07-14T13:24:15Z","snapshot_observed_at":"2026-07-17T23:18:57.221804Z","submitted_at":"2026-03-09T01:35:37Z","title":"Learning When to Trust in Contextual Social Bandits","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-15T12:59:04.484292Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2603.13356"},"observation_digest":"sha256:d463abfb8c4db2704aa727fea770f3827bd0977ea3c3bd6efe933f1a760b6e7d","observation_id":"ca93db08-baeb-405b-b3ec-e9cc1a04fb04","resolution":{"observed_at":"2026-07-15T12:59:04.484292Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2604.08606","last_updated":"2026-04-08T13:57:38Z","snapshot_observed_at":"2026-07-06T22:57:40.773778Z","submitted_at":"2026-04-08T13:57:38Z","title":"Extrapolating Volition with Recursive Information Markets","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-10T17:18:34.660366Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2604.08606"},"observation_digest":"sha256:75364bd8ea436d21407ce27e7a1ee19b0283a6d4d91b3f8df51fc4e922b9c5ae","observation_id":"800a08c5-2321-45c7-838a-1a34b6af4792","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2604.20070","last_updated":"2026-04-22T00:32:35Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-22T00:32:35Z","title":"Auditing and Controlling AI Agent Actions in Spreadsheets","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-10T00:18:56.460027Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2604.20070"},"observation_digest":"sha256:eacf562dca086136700da82f49987c6a9b01a9948562d7b7faf7dd52635823d8","observation_id":"3665eacb-86e0-463d-b7d0-64c798b4d4bc","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2604.21718","last_updated":"2026-04-26T23:28:29Z","snapshot_observed_at":"2026-07-06T23:08:14.631453Z","submitted_at":"2026-04-22T09:01:04Z","title":"Building a Precise Video Language with Human-AI Oversight","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-10T00:37:31.858728Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2604.21718"},"observation_digest":"sha256:99499c90013bc4e9b47ecc5c8f72fed9e25dce5cefc32123b5889306af7b9ecc","observation_id":"50419c15-4d9e-43c2-949d-5d1335d125af","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2604.23646","last_updated":"2026-04-26T10:31:13Z","snapshot_observed_at":"2026-07-06T23:09:48.137856Z","submitted_at":"2026-04-26T10:31:13Z","title":"Structural Enforcement of Goal Integrity in AI Agents via Separation-of-Powers Architecture","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-08T06:21:46.364082Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2604.23646"},"observation_digest":"sha256:f32d1b8392f1801f4aff05b831d031d69c629a92bcdce0408d4e21dd503148f3","observation_id":"a05e8937-e2ad-401a-9a3e-907f6cbceaf4","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2605.03808","last_updated":"2026-05-05T14:35:47Z","snapshot_observed_at":"2026-07-06T23:16:40.201701Z","submitted_at":"2026-05-05T14:35:47Z","title":"Agentic-imodels: Evolving agentic interpretability tools via autoresearch","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-07T16:37:43.371592Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2605.03808"},"observation_digest":"sha256:dcbdc83fecf871d891d1ec5027172e3b5cf1426bfdec7eb6987bffc42a7ebc07","observation_id":"b25ff9da-9ea7-4955-b92d-3df7b37fdc41","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2605.06840","last_updated":"2026-05-22T00:29:09Z","snapshot_observed_at":"2026-07-06T23:19:15.382283Z","submitted_at":"2026-05-07T18:45:46Z","title":"Extracting Search Trees from LLM Reasoning Traces Reveals Myopic Planning","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-11T00:52:59.406190Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2605.06840"},"observation_digest":"sha256:6350f1f508b84f3d0d49f1a0fa7e342844c7572c68016074650e2c3536b77a98","observation_id":"6b7f1291-242a-4bdd-9a84-c26283012f02","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2605.06840","last_updated":"2026-05-22T00:29:09Z","snapshot_observed_at":"2026-07-06T23:19:15.382283Z","submitted_at":"2026-05-07T18:45:46Z","title":"Extracting Search Trees from LLM Reasoning Traces Reveals Myopic Planning","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-12T02:23:55.323250Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2605.06840"},"observation_digest":"sha256:39222e567f5d87d96488945b047c4ceff18d96de64d688b2b7d085bf4b2e026b","observation_id":"e88a34e5-5099-4b7d-b469-1edb27a7e224","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2605.06840","last_updated":"2026-05-22T00:29:09Z","snapshot_observed_at":"2026-07-06T23:19:15.382283Z","submitted_at":"2026-05-07T18:45:46Z","title":"Extracting Search Trees from LLM Reasoning Traces Reveals Myopic Planning","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-13T06:21:00.353334Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2605.06840"},"observation_digest":"sha256:8a0d859cda7bbb8aa378742186465b3082985023c68713c1c2be15cb7c81f64d","observation_id":"52ba6201-1c33-4492-8663-24ecf82b4e34","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2605.06840","last_updated":"2026-05-22T00:29:09Z","snapshot_observed_at":"2026-07-06T23:19:15.382283Z","submitted_at":"2026-05-07T18:45:46Z","title":"Extracting Search Trees from LLM Reasoning Traces Reveals Myopic Planning","version":4},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-14T20:55:31.770238Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2605.06840"},"observation_digest":"sha256:5a1fe65d5567d5b7944b0419427f3cd5511527df405d602979746a14072a7e13","observation_id":"75d69348-83be-4026-bc55-6aff42c4b851","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2605.06840","last_updated":"2026-05-22T00:29:09Z","snapshot_observed_at":"2026-07-06T23:19:15.382283Z","submitted_at":"2026-05-07T18:45:46Z","title":"Extracting Search Trees from LLM Reasoning Traces Reveals Myopic Planning","version":5},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-25T05:57:58.487109Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2605.06840"},"observation_digest":"sha256:ae17e910c441eeb3a6e734cb5b9019acf9458d67275652099fb8f3c3b8bd9581","observation_id":"96615e44-6514-4e9a-aaa8-e9727849299b","resolution":{"observed_at":"2026-05-25T06:00:23.382677Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2605.07021","last_updated":"2026-05-20T02:19:39Z","snapshot_observed_at":"2026-07-06T23:19:25.538514Z","submitted_at":"2026-05-07T23:05:50Z","title":"Behavior Cue Reasoning: Monitorable Reasoning Improves Efficiency and Safety through Oversight","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-11T00:54:25.549158Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2605.07021"},"observation_digest":"sha256:14a7e170870766464fab5174b4eca8c1f5654b72eb2e020c27514b36b27cac91","observation_id":"6265a63d-b96a-473b-91b3-fdccb56ba11f","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2605.07021","last_updated":"2026-05-20T02:19:39Z","snapshot_observed_at":"2026-07-06T23:19:25.538514Z","submitted_at":"2026-05-07T23:05:50Z","title":"Behavior Cue Reasoning: Monitorable Reasoning Improves Efficiency and Safety through Oversight","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-21T08:29:09.122055Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2605.07021"},"observation_digest":"sha256:48a834ea296996698037f68d3df13929ce7e6285d06b76491d8c00c72f0d34af","observation_id":"a7e5db1f-8aee-4d85-8204-4e1a350632ed","resolution":{"observed_at":"2026-05-21T08:29:52.668438Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2605.08321","last_updated":"2026-05-08T16:23:47Z","snapshot_observed_at":"2026-07-06T23:20:33.816373Z","submitted_at":"2026-05-08T16:23:47Z","title":"LLM Wardens: Mitigating Adversarial Persuasion with Third-Party Conversational Oversight","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-12T00:51:17.434889Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2605.08321"},"observation_digest":"sha256:8049ae90c0553f2cfddcf86044bd85eba744715a6513dcca0f0bf900eb57c039","observation_id":"a000633a-0eea-4ccf-98db-9324d8d58369","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2605.12809","last_updated":"2026-05-12T23:01:29Z","snapshot_observed_at":"2026-07-06T23:24:27.821980Z","submitted_at":"2026-05-12T23:01:29Z","title":"Correcting Influence: Unboxing LLM Outputs with Orthogonal Latent Spaces","version":1},"reference_index":211,"source":"arxiv_source","source_observed_at":"2026-05-14T20:17:01.224864Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2605.12809"},"observation_digest":"sha256:4185f6bc475400b99cfbd33305da36ee54aea4bc86de7e0000436ef93bd9bca8","observation_id":"7e6c204e-c786-4fb9-ac09-3ce41289aba7","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2605.12991","last_updated":"2026-05-16T19:42:13Z","snapshot_observed_at":"2026-07-06T23:24:37.915639Z","submitted_at":"2026-05-13T04:45:08Z","title":"Not Just RLHF: Why Alignment Alone Won't Fix Multi-Agent Sycophancy","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-14T20:04:57.638215Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2605.12991"},"observation_digest":"sha256:6dd3bcc9540c8717ca60aeb363f0eac8c410d484e79a43f13182b6c8f23f0f2c","observation_id":"cbbff684-9dc2-4824-9d4a-d88844acec80","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2605.12991","last_updated":"2026-05-16T19:42:13Z","snapshot_observed_at":"2026-07-06T23:24:37.915639Z","submitted_at":"2026-05-13T04:45:08Z","title":"Not Just RLHF: Why Alignment Alone Won't Fix Multi-Agent Sycophancy","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-20T21:30:30.384184Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2605.12991"},"observation_digest":"sha256:e5c9821e1d5409d9a398bdfd79dc322690beffc4f04a2a19fb23e24dbc5f3c20","observation_id":"f17d7638-63a1-41b3-a6bd-51446d02fd38","resolution":{"observed_at":"2026-05-20T21:33:46.582396Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2605.13625","last_updated":"2026-05-13T14:52:40Z","snapshot_observed_at":"2026-07-06T23:25:11.623026Z","submitted_at":"2026-05-13T14:52:40Z","title":"How to Interpret Agent Behavior","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-14T18:23:25.269217Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2605.13625"},"observation_digest":"sha256:12131a4e006c1d0013ee8ae0610c008c64a76f6d948c730ec7e4cc3799f9219b","observation_id":"8053ac46-4d95-4aba-8a75-6439a1ad7f3b","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2605.25739","last_updated":"2026-07-19T12:16:05Z","snapshot_observed_at":"2026-07-23T23:20:12.103762Z","submitted_at":"2026-05-25T11:51:08Z","title":"The Behavioral Credibility Trilemma: When Calibrated Autonomy Becomes Impossible","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-29T22:28:02.493124Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2605.25739"},"observation_digest":"sha256:83bb4ebc171b84e85c7466404c842391946252d8bd534d596ff408bfe83fc2e3","observation_id":"760313a2-9ef8-4878-92e5-3c63ea3fc415","resolution":{"observed_at":"2026-06-29T22:34:01.931399Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2605.29192","last_updated":"2026-05-28T00:08:35Z","snapshot_observed_at":"2026-07-06T23:38:37.178378Z","submitted_at":"2026-05-28T00:08:35Z","title":"ReasonOps: Operator Segmentation for LLM Reasoning Traces","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-29T08:04:44.270536Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2605.29192"},"observation_digest":"sha256:d945a4669fa1dc467d99126ac4d771e991da04240893c9a95c42d12838a741b3","observation_id":"ada9729b-3f19-4131-b615-3ed97031409a","resolution":{"observed_at":"2026-06-29T08:13:15.919114Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2606.00235","last_updated":"2026-05-29T18:10:00Z","snapshot_observed_at":"2026-07-06T23:40:56.510371Z","submitted_at":"2026-05-29T18:10:00Z","title":"Civilizational Metamaterials: Engineering Coordination Under Capability Gradients and Structural Turbulence","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-28T19:29:13.317260Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2606.00235"},"observation_digest":"sha256:b1fa4c7b61c0f81e039d83a1262babc7530502c5b9a3d89639c806641969c97f","observation_id":"b082e248-ed58-4ef8-9852-9153da2db123","resolution":{"observed_at":"2026-06-28T19:32:34.344128Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2606.02866","last_updated":"2026-06-01T20:29:47Z","snapshot_observed_at":"2026-07-06T23:43:12.737017Z","submitted_at":"2026-06-01T20:29:47Z","title":"When Helping Hurts and How to Fix It: Multi-Agent Debate for Data Cleaning","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-06-28T14:05:49.696342Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2606.02866"},"observation_digest":"sha256:e6011a9cd2fed65dbd5eb241c8388298c02bd83d4c84a527105d3354dcf5b9db","observation_id":"888bdece-b8ad-438b-ba15-21aa95d4705e","resolution":{"observed_at":"2026-07-01T23:36:23.859141Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2606.07602","last_updated":"2026-05-29T09:31:25Z","snapshot_observed_at":"2026-07-06T23:47:15.041928Z","submitted_at":"2026-05-29T09:31:25Z","title":"Sample-Efficient Post-Training for LEGO Spatial-Physics Reasoning","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-06-28T23:38:38.127345Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2606.07602"},"observation_digest":"sha256:52fd07b4030fbb878238fa33270adc9ea4ded29c0a1c0950f64f02331e5bed39","observation_id":"fee3e933-890f-416d-82ff-4a0a763fc890","resolution":{"observed_at":"2026-06-28T23:42:49.528693Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2606.08044","last_updated":"2026-06-06T08:10:56Z","snapshot_observed_at":"2026-07-06T23:47:33.680573Z","submitted_at":"2026-06-06T08:10:56Z","title":"When Behavioral Safety Evaluation Fails: A Representation-Level Perspective","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-06-27T20:04:17.744876Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2606.08044"},"observation_digest":"sha256:fd71365110466708284f0e1d7e7bd5c925c0570fdb94fb7bd9a1b3b0ce05778f","observation_id":"b38cc38a-d24e-42aa-8656-65f8a2ee65ed","resolution":{"observed_at":"2026-07-02T20:57:23.088762Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2606.09711","last_updated":"2026-06-08T16:32:54Z","snapshot_observed_at":"2026-07-06T23:49:03.237958Z","submitted_at":"2026-06-08T16:32:54Z","title":"Proxy Reward Internalization and Mechanistic Exploitation: A Learned Precursor to Reward Hacking and Its Generalization","version":1},"reference_index":123,"source":"arxiv_source","source_observed_at":"2026-06-27T16:26:34.918099Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2606.09711"},"observation_digest":"sha256:d774224163c64415aec58a5992e438d8ded9fba78a0c62796fbdcd1271316159","observation_id":"340197aa-715d-48ad-8231-e0ea82a8d34f","resolution":{"observed_at":"2026-06-27T16:31:02.718795Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2606.10747","last_updated":"2026-06-09T11:57:02Z","snapshot_observed_at":"2026-07-06T23:49:56.404171Z","submitted_at":"2026-06-09T11:57:02Z","title":"The Arbiter Agent: Continually Monitoring Multi-Agent Conversations to Detect Emergent Misalignment","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-27T13:26:14.457195Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2606.10747"},"observation_digest":"sha256:9a2b5408905dc8dc10718f550391736cc8889ce41f6e4e0a281100fccdb81f47","observation_id":"d21ccf50-e626-4102-b7c9-9ccce5a3d31b","resolution":{"observed_at":"2026-06-27T13:30:56.585818Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2606.18671","last_updated":"2026-06-17T04:13:20Z","snapshot_observed_at":"2026-07-06T23:54:04.739534Z","submitted_at":"2026-06-17T04:13:20Z","title":"HANSEL: Extracting Breadcrumbs from Web Agent Trajectories for Interactive Verification","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-26T19:58:23.151980Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2606.18671"},"observation_digest":"sha256:61cc10ce8e430bd34215003d27bd5de2f2e4bff6b6082a57191a1ba3637f35d2","observation_id":"3bd557f4-6cc1-46f9-858a-d0aee14bd328","resolution":{"observed_at":"2026-07-04T02:09:22.850175Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-12T03:28:32.737784Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.03292","last_updated":"2026-07-03T13:01:56Z","snapshot_observed_at":"2026-07-12T03:28:32.224203Z","submitted_at":"2026-07-03T13:01:56Z","title":"Regulating AI: Where U.S. State Policy and HCI (Mis)align","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-12T03:28:32.737784Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2607.03292"},"observation_digest":"sha256:3436a4b65f73bd9ed2b0ffa31f9ab8c25f042f220bf6b6e2c78b54dc5099a7dc","observation_id":"03ece39d-ccb5-4cef-bbaa-cb848f64945b","resolution":{"observed_at":"2026-07-12T03:28:32.737784Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-11T16:47:52.768236Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.04590","last_updated":"2026-07-06T01:33:12Z","snapshot_observed_at":"2026-07-11T16:47:52.180168Z","submitted_at":"2026-07-06T01:33:12Z","title":"Attention Limited Reward Learning","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-11T16:47:52.768236Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2607.04590"},"observation_digest":"sha256:f5413da3c1d34966a5171b9468f9188c0abc818ea18234fcb90878b8f641a1e0","observation_id":"58f09301-6d3f-449b-852e-96cf1488e19c","resolution":{"observed_at":"2026-07-11T16:47:52.768236Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2607.05394","last_updated":"2026-07-08T17:57:18Z","snapshot_observed_at":"2026-07-11T23:19:57.063758Z","submitted_at":"2026-07-06T17:59:58Z","title":"Weak-to-Strong Generalization via Direct On-Policy Distillation","version":1},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-07-07T12:31:42.224094Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2607.05394"},"observation_digest":"sha256:860cada30c0db88930ba79a54c78e7562823c84abdae37215a49f0fb279b362c","observation_id":"6af3d5c1-69af-47a7-b4d8-1a4627d85359","resolution":{"observed_at":"2026-07-07T12:33:44.994534Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-11T07:01:56.628017Z","title":"Measuring progress on scalable oversight for large language models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.05394","last_updated":"2026-07-08T17:57:18Z","snapshot_observed_at":"2026-07-11T23:19:57.063758Z","submitted_at":"2026-07-06T17:59:58Z","title":"Weak-to-Strong Generalization via Direct On-Policy Distillation","version":2},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-07-11T07:01:56.628017Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2607.05394"},"observation_digest":"sha256:51323521d198f3141f2200bc49efd988746884b1c59a64b16598279ff71ddbd9","observation_id":"efc4af27-2ee3-4a29-9127-9092046319d2","resolution":{"observed_at":"2026-07-11T07:01:56.628017Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2607.07040","last_updated":"2026-07-08T06:19:10Z","snapshot_observed_at":"2026-07-11T23:18:50.679052Z","submitted_at":"2026-07-08T06:19:10Z","title":"Measuring Intelligence Beyond Human Scale","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-09T21:21:39.305904Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2607.07040"},"observation_digest":"sha256:8b66a1cbff0c0f41a76ce121661ffa65e88ae4f14c2317c2d6f918ab38b1b5c7","observation_id":"d08d1af3-c0a9-48b4-9d10-950748120de9","resolution":{"observed_at":"2026-07-09T21:26:34.379823Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2607.07695","last_updated":"2026-07-08T17:53:56Z","snapshot_observed_at":"2026-07-11T23:20:20.966286Z","submitted_at":"2026-07-08T17:53:56Z","title":"Institutional Red-Teaming: Deployment Rules, Not Just Models, Causally Shape Multi-Agent AI Safety","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-07-09T02:17:34.473738Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2607.07695"},"observation_digest":"sha256:3544bdd47ae93a5a37b6a26c104c546604996a27dcae260fc025b6ac1b407350","observation_id":"63bcabe1-a474-4187-a067-0c06b3a8122b","resolution":{"observed_at":"2026-07-09T02:25:55.820061Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-07-10T18:37:31.166799Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2607.07774","last_updated":"2026-07-08T16:46:06Z","snapshot_observed_at":"2026-07-12T23:17:59.702153Z","submitted_at":"2026-07-08T16:46:06Z","title":"ScopeJudge: Cost-Aware Pre-Execution Gating for Offensive Security Agents","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-10T18:29:50.731038Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2607.07774"},"observation_digest":"sha256:173b00fa81efef295109c3952c1be431c82360ad065e3f29078ae5fd01a34c47","observation_id":"7d790367-a3a2-4cfc-9264-d510447bdbd1","resolution":{"observed_at":"2026-07-10T18:37:31.168237Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2211.03540/citation-record","integrity":"/paper/2211.03540/integrity","json":"/paper/2211.03540/citation-record.json","paper":"/paper/2211.03540"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The case for aligning narrowly superhuman models , url=","venue":null,"work_id":"65c0dfdd-8094-41fa-a3e3-8e1ff9983674","year":null},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:9fae079c95866e9e1ab5f46ef898336a7e6b5f102adaf52ab3f2539637fc7ce0","observation_id":"e4174cb8-408e-4da0-b551-8d1f93f2ef5b","resolution":{"observed_at":"2026-05-17T15:01:41.420621Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"c9a23a44-916a-409a-873f-4454e7e1f53e","year":null},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:f5df6b60326061b658c7bf0dad5527c91a48d0713e863ca9460591a39eff542e","observation_id":"43f90ad9-77ef-4b3d-8ef8-da54d4337960","resolution":{"observed_at":"2026-05-17T15:01:41.387249Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"f7c32c41-1100-4d81-a520-fa449c400ad0","year":null},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:416fc3ff6ceaed0de92b9eabc9d1f92003e8685f95b58dfc66e603a611641344","observation_id":"1fc12978-fdcc-4fdf-8f40-3a7f0935865b","resolution":{"observed_at":"2026-05-17T15:01:41.379940Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Weld , journal=","venue":null,"work_id":"d43c6246-e8c2-4e4b-b05d-871df008ca0e","year":null},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:3dbf9736aeede63f325e2a18a5dace994f5ff692e4a68b9312f601a829c2f863","observation_id":"e61c81fe-cc1d-41c0-ba01-0f0747937c12","resolution":{"observed_at":"2026-05-17T15:01:41.357536Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T15:16:18.971970Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":"4a77e424-cf16-4c16-8d94-2ee44db893d1","year":null},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:69154cdc33bfc7d219dfa87cf5efd9065cb76cf6deb2d96ee230e0d3c3d2cd16","observation_id":"6c904c71-1b06-4520-a528-11806f77ca97","resolution":{"observed_at":"2026-05-17T15:01:41.398455Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":"b202dcb7-0590-4a1b-b41f-72cd5085cc57","year":null},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:0b0cc55133c1dae344c511684a37290151e2732e73935626d1fbabef59a563e1","observation_id":"e38530dd-4572-4342-b936-787ee2a1ad4f","resolution":{"observed_at":"2026-05-17T15:01:41.412471Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2014 , isbn =","venue":null,"work_id":"d659a874-9049-4030-9db9-0051ae4aac0a","year":2014},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:cdaafcc0ebbc638b30824f9b390e4e38ec44e8bbd3abb85e98c1dd483f3bbc63","observation_id":"68d621d7-74b2-410e-a62b-39fdb16acce7","resolution":{"observed_at":"2026-05-17T15:01:41.361299Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"ad1e04e1-694b-47e7-8c77-9d79995e2dd3","year":null},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:ce2aa3d1bee0f6fe2202abed6b29e0a3d2437725ebaca75d579efe5f5ffec8ad","observation_id":"7f12dfe8-a82b-4e83-ae87-a8f1fa0e797c","resolution":{"observed_at":"2026-05-17T15:01:41.364577Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Organizational behavior and human performance , volume=","venue":null,"work_id":"c89953a5-fdfa-4d60-8d3f-d63fa6e36708","year":1980},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:4fd26f869cbe6d9d49e78f48af06ff8bf7cd118012faa0f3eafe61910bf42f84","observation_id":"0e030b36-4d66-468e-904f-da803314b8ce","resolution":{"observed_at":"2026-05-17T15:01:41.368635Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"110d7e0a-a7a6-489a-95e8-f31b6b600581","year":null},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:d1253cea81410c92a319d270b01d36a857ed8d2c22d459428a78247db4458529","observation_id":"56ee0fd4-2e0e-40eb-ab9f-be6117d81079","resolution":{"observed_at":"2026-05-17T15:01:41.372508Z","resolver_source":"raw_fallback","status":"parse_uncertain"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Submitted to The Eleventh International Conference on Learning Representations , year=","venue":null,"work_id":"551329bb-c30d-4161-b253-3c7ae54e2c31","year":null},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:a9f072678aa8b9bbaae2a0a3fb4c5f36014d9050e55548a73176691b95ffa7da","observation_id":"ebf72b88-e288-4613-9b41-1c5355bd2a76","resolution":{"observed_at":"2026-05-17T15:01:41.376177Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1606.06565","last_updated":"2016-07-25T17:23:29Z","snapshot_observed_at":"2026-07-06T05:00:46.434335Z","submitted_at":"2016-06-21T13:37:05Z","title":"Concrete Problems in AI Safety","version":2},"cited_work":{"arxiv_id":"1606.06565","doi":"10.48550/arxiv.1606.06565","metadata_source":"pith","pith_arxiv_id":"1606.06565","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Concrete Problems in AI Safety","venue":"cs.AI","work_id":"c8d14fbe-6eab-464a-95b3-778aabd82fa3","year":2016},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/1606.06565","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:8033e94aa07340a610a92866b76fd227b0648c94c1ede0fcf5263e092bfbe848","observation_id":"48075c22-a281-43ab-b94b-91eba9ddc028","resolution":{"observed_at":"2026-05-17T15:01:41.251250Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:52.215761+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:52.215761+00:00","source":"openalex_status_cache"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2112.00861","last_updated":"2021-12-09T21:40:22Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-12-01T22:24:34Z","title":"A General Language Assistant as a Laboratory for Alignment","version":3},"cited_work":{"arxiv_id":"2112.00861","doi":"10.48550/arxiv.2112.00861","metadata_source":"pith","pith_arxiv_id":"2112.00861","snapshot_observed_at":"2026-07-10T15:37:20.432112Z","title":"A General Language Assistant as a Laboratory for Alignment","venue":"cs.CL","work_id":"a43f9ea0-01be-47d5-b8ee-a1a9f73381c5","year":2021},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/2112.00861","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:d0ce7dc8d5322f7e408bcb335bf4813bdcc5e769762770d962bc8c36f1d9ee2c","observation_id":"af034535-36a8-4349-aee9-185ad063330c","resolution":{"observed_at":"2026-05-17T15:01:41.258042Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-05-20T18:52:13.802987+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T18:52:13.802987+00:00","source":"openalex_status_cache"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2204.05862","last_updated":"2022-04-12T15:02:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-04-12T15:02:38Z","title":"Training a Helpful and Harmless Assistant with Reinforcement Learning from Human Feedback","version":1},"cited_work":{"arxiv_id":"2204.05862","doi":"10.1016/j.respol.2005.01.014","metadata_source":"pith","pith_arxiv_id":"2204.05862","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Training a Helpful and Harmless Assistant with Reinforcement Learning from Human Feedback","venue":"cs.CL","work_id":"a1f2574b-a899-4713-be60-c87ba332656c","year":2022},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/2204.05862","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:87c098735b01077f80772728af6fe6f77fae8ec1e82299c0e568880f9d9b3c64","observation_id":"068e9cb5-7109-4594-a000-6428372bbef8","resolution":{"observed_at":"2026-05-17T15:01:41.264900Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"e2735ff0-fbfc-4989-a9a3-5f9b2f57c0b0","year":2021},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:bfda038a2e12c2a74203723598b7c05ebbd27470d5a52e7b52929e89a2ca4e84","observation_id":"63ac058c-2294-42fd-a11c-5fe3c6dd8d34","resolution":{"observed_at":"2026-05-17T15:01:41.391002Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"2f2ee7d5-4928-4af1-8882-8f8de283ce00","year":2014},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:e63dbaddb73372c25e12a9902d304d108d695b91b2e210e6090a163f42d21903","observation_id":"0f41ae8f-6573-400f-b128-5808eaa93162","resolution":{"observed_at":"2026-05-17T15:01:41.394544Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1810.08575","last_updated":"2018-10-19T16:30:48Z","snapshot_observed_at":"2026-07-06T07:09:24.322232Z","submitted_at":"2018-10-19T16:30:48Z","title":"Supervising strong learners by amplifying weak experts","version":1},"cited_work":{"arxiv_id":"1810.08575","doi":"10.1177/0162243915579283","metadata_source":"pith","pith_arxiv_id":"1810.08575","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Supervising strong learners by amplifying weak experts","venue":"cs.LG","work_id":"5d6c7456-8e86-44d2-b157-ac0daeb9587c","year":2018},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/1810.08575","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:b7d0dfefe75dfbec985e485db2e3c03c641b25d3c9ccc7d169edac84712059b8","observation_id":"770ab2e0-d5bd-4de9-b653-926b1649a85f","resolution":{"observed_at":"2026-05-17T15:01:41.271749Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"a88be2b4-23af-4455-95f0-f75817c930d0","year":2021},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:c1af636d72f58a2d18e1d0562d12f36f42a1c20d0ef42351b8f81f8d515d8d99","observation_id":"d627035e-63ff-4ce6-a56d-6bd6c823a23a","resolution":{"observed_at":"2026-05-17T15:01:41.403924Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"621b9daf-7ada-4e20-9b5c-aad562c58763","year":2017},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:f6d7594ead57679f0ea6730aeb0fb6d001370d19ac314a51dc8638b1924620d2","observation_id":"d83291cd-833c-4417-aacd-50548c71d4c3","resolution":{"observed_at":"2026-05-17T15:01:41.408201Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"7481.345064","doi":"10.1145/3397481.3450644","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"In: 26th Inter- national Conference on Intelligent User Interfaces, pp","venue":null,"work_id":"f31d7193-7f47-46e0-b7ad-548a7e369132","year":2021},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:6846344a33b755b9f3dc6463c9b3716c7038224c6f2947af99881e457669e1d3","observation_id":"2630f559-5066-4a0a-b0ac-fdb86d473c89","resolution":{"observed_at":"2026-05-17T15:01:41.208363Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"70a66136-a8fc-440f-acfb-c427e2df9d4f","year":2021},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:7c4f76f87468789ea98f4f3e61b1a092e6fc6facf04eb3818148a8557aae1607","observation_id":"2082097a-2392-4c89-b1bc-af4e73fabce2","resolution":{"observed_at":"2026-05-17T15:01:41.416665Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2009.03300","last_updated":"2021-01-12T18:57:11Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-09-07T17:59:25Z","title":"Measuring Massive Multitask Language Understanding","version":3},"cited_work":{"arxiv_id":"2009.03300","doi":"10.48550/arxiv.2009.03300","metadata_source":"pith","pith_arxiv_id":"2009.03300","snapshot_observed_at":"2026-07-11T01:37:42.666395Z","title":"Measuring Massive Multitask Language Understanding","venue":"cs.CY","work_id":"e87ec49a-544b-4ec8-8991-75298c64ff5e","year":2020},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/2009.03300","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:88a8f6c3799b65930b04f77522410c57230457e73dd7e755c57ccf9c0e6304a1","observation_id":"ab1ac030-0194-48be-8b88-645a90aa4394","resolution":{"observed_at":"2026-05-17T15:01:41.277357Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-11T02:19:14.061878+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T02:19:14.061878+00:00","source":"openalex_status_cache"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"83eeb547-de7a-4d4c-a9ec-07e5146322c1","year":2020},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:e1c762316e89377a9e7f1d92954ed3b1582bd59bd3698dac3713a0658c2294b9","observation_id":"77c61ac6-aff5-4770-b065-45719dbc9a48","resolution":{"observed_at":"2026-05-17T15:01:41.345384Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1906.01820","last_updated":"2021-12-01T11:22:52Z","snapshot_observed_at":"2026-07-06T07:58:02.684338Z","submitted_at":"2019-06-05T04:43:25Z","title":"Risks from Learned Optimization in Advanced Machine Learning Systems","version":3},"cited_work":{"arxiv_id":"1906.01820","doi":"10.48550/arxiv.1906.01820","metadata_source":"pith","pith_arxiv_id":"1906.01820","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Risks from Learned Optimization in Advanced Machine Learning Systems","venue":"cs.AI","work_id":"871c0bb7-e08b-4d8b-be76-610707c748dd","year":2019},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/1906.01820","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:64e7bc7c7c01367cee818dc6a5ce2056c5da417c22b459b9ca627ca77d81cad2","observation_id":"219991db-4899-45af-b29f-a371f39ce014","resolution":{"observed_at":"2026-05-17T15:01:41.284183Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"03915460-90f2-4752-860e-52375213cae3","year":2019},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:4723fc2c1287ca441803da164b911b25a3b6a85adbd5da087cb67364fb4d0863","observation_id":"a16fa239-61ad-4915-8a27-069ed6c94bb6","resolution":{"observed_at":"2026-05-17T15:01:41.353088Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1805.00899","last_updated":"2018-10-22T17:36:07Z","snapshot_observed_at":"2026-07-06T06:37:02.085757Z","submitted_at":"2018-05-02T16:27:32Z","title":"AI safety via debate","version":2},"cited_work":{"arxiv_id":"1805.00899","doi":"10.48550/arxiv.1805.00899","metadata_source":"pith","pith_arxiv_id":"1805.00899","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"AI safety via debate","venue":"stat.ML","work_id":"13c1ec37-af93-438a-bdf0-f2eafaee5635","year":2018},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/1805.00899","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:a8b066e26199e64723aa0a8567298d47977f5fb38901d62c6a4eb0a2626cd6e7","observation_id":"616e3a07-4f6d-4b1a-b950-b71a2e5d4011","resolution":{"observed_at":"2026-05-17T15:01:41.289796Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2207.05221","last_updated":"2022-11-21T16:38:35Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-07-11T22:59:39Z","title":"Language Models (Mostly) Know What They Know","version":4},"cited_work":{"arxiv_id":"2207.05221","doi":"10.1145/3618260.3649777","metadata_source":"pith","pith_arxiv_id":"2207.05221","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Language Models (Mostly) Know What They Know","venue":"cs.CL","work_id":"8ca58a10-da41-4f70-baae-7e449512e345","year":2022},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/2207.05221","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:5080ea7508ef41b2d651f066cce104dee31b8641e797f5fbdb2f4d8ddf0f5651","observation_id":"0c268239-7b4f-4d82-b546-4bdca85a5e61","resolution":{"observed_at":"2026-05-17T15:01:41.295310Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-05-21T07:53:13.382372+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-21T07:53:13.382372+00:00","source":"openalex_status_cache"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.11916","last_updated":"2023-01-29T05:14:17Z","snapshot_observed_at":"2026-07-06T13:13:18.623287Z","submitted_at":"2022-05-24T09:22:26Z","title":"Large Language Models are Zero-Shot Reasoners","version":4},"cited_work":{"arxiv_id":"2205.11916","doi":"10.48550/arxiv.2205.11916","metadata_source":"pith","pith_arxiv_id":"2205.11916","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Large Language Models are Zero-Shot Reasoners","venue":"cs.CL","work_id":"d9b7eb1a-7165-46ff-9f06-d2f0b9d6f95d","year":2022},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/2205.11916","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:8631b87cda4ddfd834eff86c8146670a0a2b90f32942ed416046af6f5d4a6c74","observation_id":"5d9aba1e-ff68-465f-8c5b-fd7298f3886a","resolution":{"observed_at":"2026-05-17T15:01:41.301227Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2112.11471","last_updated":"2021-12-21T19:00:02Z","snapshot_observed_at":"2026-07-06T12:21:13.440357Z","submitted_at":"2021-12-21T19:00:02Z","title":"Towards a Science of Human-AI Decision Making: A Survey of Empirical Studies","version":1},"cited_work":{"arxiv_id":"2112.11471","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2112.11471","snapshot_observed_at":"2026-07-04T21:00:09.994882Z","title":"Vera Liao, Alison Smith-Renner, and Chenhao Tan","venue":null,"work_id":"4de75202-9f61-494a-a586-25504d0e772c","year":2021},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/2112.11471","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:fa8d7ae8172b417d5927edfbccf51cfce49bf797f398c0b1da508ce772e73748","observation_id":"bb04b6c5-6c29-4d52-ad20-af9f0fb2442f","resolution":{"observed_at":"2026-05-17T15:01:41.307326Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"9672.293987","doi":"10.1145/2939672.2939874","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Bach, and Jure Leskovec","venue":null,"work_id":"e6f0aa41-986c-46d4-bbd5-a8134e1339d1","year":2016},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:479ba36d90cbe6b0a3f4affd172820e45e8494f012e781831725470cec2aba75","observation_id":"a056445e-b407-4ff3-a9f0-721f4559b0bd","resolution":{"observed_at":"2026-05-17T15:01:41.233339Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1811.07871","last_updated":"2018-11-19T18:48:04Z","snapshot_observed_at":"2026-07-06T07:15:51.575438Z","submitted_at":"2018-11-19T18:48:04Z","title":"Scalable agent alignment via reward modeling: a research direction","version":1},"cited_work":{"arxiv_id":"1811.07871","doi":"10.48550/arxiv.1811.07871","metadata_source":"pith","pith_arxiv_id":"1811.07871","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scalable agent alignment via reward modeling: a research direction","venue":"cs.LG","work_id":"1e9c4f6d-b369-4bd2-8e6e-1ec00318c924","year":2018},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/1811.07871","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:dd5a18eb0f31f7d2d703ae00a2e2d3ebc0cc9aaae79aba9cbce6facc6808a9fc","observation_id":"bc09f87c-8d8a-4b89-85ed-780447feea9e","resolution":{"observed_at":"2026-05-17T15:01:41.314000Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"5a01971b-0c0b-44fc-aa32-67c56ab7ed2c","year":1980},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:35e2cc276f439d41229828e78f2b8a96e858afc4f73fc65fa5218c6142dab440","observation_id":"bad5145c-a75a-4df8-b393-d975c9763547","resolution":{"observed_at":"2026-05-17T15:01:41.349709Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2022.acl-long.229","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T15:37:20.355771Z","title":"URLhttps://doi.org/10.18653/v1/2022.acl-long.229","venue":"Proceedings of the 60th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","work_id":"414143a1-235c-49ec-a733-5cff2faefa32","year":2022},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:16a2eb67be0068dfd9b8e657ad223a285549a9b30e3c2375c2ce32e1da975ee4","observation_id":"253a3c6d-02c9-40ab-b123-de61f5ba7ff9","resolution":{"observed_at":"2026-05-17T15:01:41.226048Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-17T23:20:56.309351+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-17T23:20:56.309351+00:00","source":"openalex_status_cache"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1145/3479552","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"b50e70f2-733e-403c-a094-242836cbfd81","year":2021},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:690e1de1403bcf5ab15fd7cad554d9b140f582a991a29c4f03df6f905cba2ca2","observation_id":"51b8fa22-8dd0-4bcd-b1cc-1ac700418a71","resolution":{"observed_at":"2026-05-17T15:01:41.221106Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-13T22:51:22.212958+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T22:51:22.212958+00:00","source":"openalex_status_cache"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2112.00114","last_updated":"2021-11-30T21:32:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-11-30T21:32:46Z","title":"Show Your Work: Scratchpads for Intermediate Computation with Language Models","version":1},"cited_work":{"arxiv_id":"2112.00114","doi":"10.48550/arxiv.2112.00114","metadata_source":"pith","pith_arxiv_id":"2112.00114","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Show Your Work: Scratchpads for Intermediate Computation with Language Models","venue":"cs.LG","work_id":"a05b1e60-8e76-4f26-9bea-28927a5f8620","year":2021},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/2112.00114","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:04193e01c9f2636f290ec0ac6367eaff770f8e534f7f204e0fecd27b4c76979c","observation_id":"2c94a520-7597-45c9-96ba-0b3323890368","resolution":{"observed_at":"2026-05-17T15:01:41.319669Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-05-23T16:25:23.975867+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T16:25:23.975867+00:00","source":"openalex_status_cache"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2022.naacl-main.391","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Q u ALITY : Question Answering with Long Input Texts, Yes!","venue":null,"work_id":"9cdca079-cddd-4bca-a7cd-777d06cd27b8","year":2022},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:5cf0dd555aac4c1c20fa7a5c57cb856faddb0ee40361a2c087a5bbc8c7590b42","observation_id":"90bfc13f-24b3-4db8-a064-b31d33d74640","resolution":{"observed_at":"2026-05-17T15:01:41.215737Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-13T13:50:31.305152+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T13:50:31.305152+00:00","source":"openalex_status_cache"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.10860","last_updated":"2022-10-19T19:48:50Z","snapshot_observed_at":"2026-07-06T14:07:55.304725Z","submitted_at":"2022-10-19T19:48:50Z","title":"Two-Turn Debate Doesn't Help Humans Answer Hard Reading Comprehension Questions","version":1},"cited_work":{"arxiv_id":"2210.10860","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2210.10860","snapshot_observed_at":"2026-06-30T11:14:37.708645Z","title":"Parrish, H","venue":null,"work_id":"766d43f1-2185-4dcc-958a-ce4f13091673","year":2022},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/2210.10860","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:4ba3c32a9952eba45ed807a26acd780b69e12fead5552b621d76954c55b81c87","observation_id":"c03a36e3-1208-41f4-9048-799a93ffb05e","resolution":{"observed_at":"2026-05-17T15:01:41.325393Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2204.05212","last_updated":"2022-04-13T13:46:13Z","snapshot_observed_at":"2026-07-06T12:59:12.318367Z","submitted_at":"2022-04-11T15:56:34Z","title":"Single-Turn Debate Does Not Help Humans Answer Hard Reading-Comprehension Questions","version":2},"cited_work":{"arxiv_id":"2204.05212","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2204.05212","snapshot_observed_at":"2026-07-01T23:36:23.849142Z","title":null,"venue":null,"work_id":"e3c4be4c-472a-428e-a279-ea51b987310d","year":2022},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/2204.05212","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:386b9b22a84b23c243559514eab65b2e6b6689ca56a73f3189e207f0633606ba","observation_id":"5c48183f-f710-4029-920c-05f742cf5ff8","resolution":{"observed_at":"2026-05-17T15:01:41.330513Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.05802","last_updated":"2022-06-14T01:16:24Z","snapshot_observed_at":"2026-07-06T13:19:58.934755Z","submitted_at":"2022-06-12T17:40:53Z","title":"Self-critiquing models for assisting human evaluators","version":2},"cited_work":{"arxiv_id":"2206.05802","doi":"10.48550/arxiv.2206.05802","metadata_source":"pith","pith_arxiv_id":"2206.05802","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Self-critiquing models for assisting human evaluators","venue":"cs.CL","work_id":"3fcefdd1-22ab-4648-a683-cb1555e7a50e","year":2022},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/2206.05802","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:61a4006786303fd3b17a67e481ed0876d244b8536c4f3fec11d45b4b1249e34a","observation_id":"e3ebd67e-6b71-48ec-9f1b-4f3fc098f26e","resolution":{"observed_at":"2026-05-17T15:01:41.336065Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"902de4bc-c44a-4a92-b499-3d24a1f8fce0","year":2020},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:a8c001e5280bdfd6722e2ae19760442089bcedab910ead9da94d8192ba7eadab","observation_id":"5d3449ba-ecfb-45eb-b5ca-99e8eb70b6d8","resolution":{"observed_at":"2026-05-17T15:01:41.384078Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2201.11903","last_updated":"2023-01-10T23:07:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-01-28T02:33:07Z","title":"Chain-of-Thought Prompting Elicits Reasoning in Large Language Models","version":6},"cited_work":{"arxiv_id":"2201.11903","doi":"10.48550/arxiv.2201.11903","metadata_source":"pith","pith_arxiv_id":"2201.11903","snapshot_observed_at":"2026-07-10T21:57:37.821767Z","title":"Chain-of-Thought Prompting Elicits Reasoning in Large Language Models","venue":"cs.CL","work_id":"d1cf6693-a082-403c-ada9-dac7b96341f9","year":2022},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/2201.11903","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:189d877e7f295bc8eed029b4e8bbd77b8332085cdb35b8053303bfdda284bc8f","observation_id":"fe0086a4-287f-4f52-b76b-75c9e0f80dde","resolution":{"observed_at":"2026-05-17T15:01:41.341241Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-18T08:21:06.391775+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-18T08:21:06.391775+00:00","source":"openalex_status_cache"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2109.10862","last_updated":"2021-09-27T21:12:49Z","snapshot_observed_at":"2026-07-06T11:50:20.455700Z","submitted_at":"2021-09-22T17:34:18Z","title":"Recursively Summarizing Books with Human Feedback","version":2},"cited_work":{"arxiv_id":"2109.10862","doi":null,"metadata_source":"pith","pith_arxiv_id":"2109.10862","snapshot_observed_at":"2026-07-03T09:07:47.851004Z","title":"Recursively Summarizing Books with Human Feedback","venue":"cs.CL","work_id":"aeb94d9f-1a59-4bbb-b591-0ccd732aa0f8","year":2021},"citing_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-05-17T15:01:41.161487Z"},"links":{"cited_paper":"/paper/2109.10862","citing_paper":"/paper/2211.03540"},"observation_digest":"sha256:906da01947053310234f36cb00de2e0cab2e1771c56962421d19a52b9cec39b1","observation_id":"ccb6a507-96d6-47ff-926c-880bb6cb9849","resolution":{"observed_at":"2026-05-17T15:28:19.819648Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","latest_version":2,"primary_category":"cs.HC","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models"},"reference_resolution":{"displayed":42,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":1,"unresolved":0,"verified_exact":22,"verified_fuzzy":19},"total_outbound_references":42},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"thesis":"As of 28 July 2026, this Paper Citation Record lists 42 of 42 outbound references and 47 inbound Pith citation observations for arXiv:2211.03540."}