{"as_of":"2026-08-11T02:40:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:02e069ecf580b184001475b013ee2e0201eaedd16879bb4c26cdf92623626bf1","coverage":[{"denominator":29,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":29,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-11T06:26:29.196349Z","state":"measured"},{"denominator":129,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":129,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":255,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":100,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T01:04:06.622696Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"pith","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":91,"observed_at":"2026-08-05T02:28:24.338817Z","source":"pith"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2311.05232","last_updated":"2024-11-19T12:42:45Z","snapshot_observed_at":"2026-07-06T16:45:07.733095Z","submitted_at":"2023-11-09T09:25:37Z","title":"A Survey on Hallucination in Large Language Models: Principles, Taxonomy, Challenges, and Open Questions","version":2},"reference_index":287,"source":"pdf_text","source_observed_at":"2026-05-13T02:46:26.957539Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2311.05232"},"observation_digest":"sha256:89517d3dc842c2d2ec05e9638f6e04e714140b355a61cc1c54da5837c1508955","observation_id":"fe7dadc4-83c5-4090-a177-5e3f7b6e53a6","resolution":{"observed_at":"2026-05-13T02:47:07.663376Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-11T01:04:06.622696Z","title":"Sharma, M","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.19010","last_updated":"2024-12-26T00:54:03Z","snapshot_observed_at":"2026-08-11T00:55:33.535269Z","submitted_at":"2024-12-26T00:54:03Z","title":"A theory of appropriateness with applications to generative artificial intelligence","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:06.622696Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2412.19010"},"observation_digest":"sha256:f4421770682222958dc39fb2ba4688184f59c4a8cb8286a79f2ecffafc3ec39b","observation_id":"3a7f5de1-0feb-4bf7-b453-e335b530db3a","resolution":{"observed_at":"2026-08-11T01:04:06.622696Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-10T21:24:10.556478Z","title":"Bowman, Newton Cheng, Esin Durmus, Zac Hatfield-Dodds, Scott R","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.04952","last_updated":"2025-01-09T03:59:10Z","snapshot_observed_at":"2026-08-10T21:19:04.243505Z","submitted_at":"2025-01-09T03:59:10Z","title":"Open Problems in Machine Unlearning for AI Safety","version":1},"reference_index":113,"source":"arxiv_source","source_observed_at":"2026-08-10T21:24:10.556478Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2501.04952"},"observation_digest":"sha256:e3574b6d8275cff450ac9e4002714bece7faeeaf9857fb5be4b5197e0d32a293","observation_id":"00ec60a4-1eda-4a43-9875-3b600fd84f18","resolution":{"observed_at":"2026-08-10T21:24:10.556478Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-10T21:20:27.477105Z","title":"arXiv preprint arXiv:2310.13548 (2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.05171","last_updated":"2025-05-21T03:51:09Z","snapshot_observed_at":"2026-08-10T21:11:17.010609Z","submitted_at":"2025-01-09T11:45:05Z","title":"Emergence of human-like polarization among large language model agents","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-10T21:20:27.477105Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2501.05171"},"observation_digest":"sha256:17dee670db65c4818780c420bb84f02a99a3998658e0e29ccdbf91e5b5febf30","observation_id":"69380e8a-76c8-426c-be72-1cde0bbabcd1","resolution":{"observed_at":"2026-08-10T21:20:27.477105Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-10T20:05:12.562654Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09431","last_updated":"2025-01-16T09:59:45Z","snapshot_observed_at":"2026-08-10T19:59:54.145648Z","submitted_at":"2025-01-16T09:59:45Z","title":"A Survey on Responsible LLMs: Inherent Risk, Malicious Use, and Mitigation Strategy","version":1},"reference_index":176,"source":"pdf_text","source_observed_at":"2026-08-10T20:05:12.562654Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2501.09431"},"observation_digest":"sha256:14c727554a10aae661fd381fcfc1cd1b7cbb5c1dc1bd81952fd79e2cd81828a8","observation_id":"e186f360-090f-4349-b80a-c3a84bc0adab","resolution":{"observed_at":"2026-08-10T20:05:12.562654Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-10T19:55:50.737082Z","title":"Towards understanding sycophancy in language models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.09620","last_updated":"2025-05-29T02:21:03Z","snapshot_observed_at":"2026-08-10T19:47:49.643885Z","submitted_at":"2025-01-16T16:00:37Z","title":"Beyond Reward Hacking: Causal Rewards for Large Language Model Alignment","version":2},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-10T19:55:50.737082Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2501.09620"},"observation_digest":"sha256:18b10482071857c2931f3e82197e4d87e6cf988bb0a37cbe24bf2011c50a1e02","observation_id":"b82dd49e-e50e-46c1-9825-f29f98496273","resolution":{"observed_at":"2026-08-10T19:55:50.737082Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-10T15:18:19.150186Z","title":"Towards understanding sycophancy in language models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.14294","last_updated":"2025-03-02T06:49:21Z","snapshot_observed_at":"2026-08-10T15:12:35.412010Z","submitted_at":"2025-01-24T07:24:23Z","title":"Examining Alignment of Large Language Models through Representative Heuristics: The Case of Political Stereotypes","version":3},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-10T15:18:19.150186Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2501.14294"},"observation_digest":"sha256:aeeba69f70ba8ba69b9c84741a03dbcff1fa585e77656b0cfb71ac71bc4cd95b","observation_id":"9e5e7b8b-25c5-47bc-8ab3-87be46909143","resolution":{"observed_at":"2026-08-10T15:18:19.150186Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-10T04:47:29.302683Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.17348","last_updated":"2025-01-31T17:51:30Z","snapshot_observed_at":"2026-08-10T04:41:47.507327Z","submitted_at":"2025-01-28T23:50:02Z","title":"Better Slow than Sorry: Introducing Positive Friction for Reliable Dialogue Systems","version":2},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-10T04:47:29.302683Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2501.17348"},"observation_digest":"sha256:2487d001a1b972ed1f62067f3f4747167c822b2f54a3f5be4f14fe87915405f9","observation_id":"c2b50e1b-3ec9-4c3d-93e6-c568538e35dd","resolution":{"observed_at":"2026-08-10T04:47:29.302683Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-09T11:53:43.820296Z","title":"R., Cheng, N., Durmus, E., Hatfield-Dodds, Z., Johnston, S","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.02528","last_updated":"2025-02-04T17:50:08Z","snapshot_observed_at":"2026-08-09T17:51:48.077638Z","submitted_at":"2025-02-04T17:50:08Z","title":"Why human-AI relationships need socioaffective alignment","version":1},"reference_index":116,"source":"arxiv_source","source_observed_at":"2026-08-09T11:53:43.820296Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2502.02528"},"observation_digest":"sha256:aaf7016bf6c33b428d8e8ad7fa4335bf923afd8b15b38532a523292659d1028b","observation_id":"299e982f-0e7f-499e-ba11-40b46849b9fd","resolution":{"observed_at":"2026-08-09T11:53:43.820296Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-07T22:21:52.698209Z","title":"R., Cheng, N., Durmus, E., Hatfield-Dodds, Z., Johnston, S","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.09192","last_updated":"2025-05-27T17:24:38Z","snapshot_observed_at":"2026-08-08T23:01:32.478617Z","submitted_at":"2025-02-13T11:32:09Z","title":"Thinking beyond the anthropomorphic paradigm benefits LLM research","version":2},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-07T22:21:52.698209Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2502.09192"},"observation_digest":"sha256:34ce308b3f213f299e98bf162d4dd7a3948299838bda088f1e7824d92512dc4b","observation_id":"df83911b-8f55-40ed-a75a-cbe9207ffd91","resolution":{"observed_at":"2026-08-07T22:21:52.698209Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2502.16810","last_updated":"2026-05-01T23:52:08Z","snapshot_observed_at":"2026-07-31T15:13:25.308192Z","submitted_at":"2025-02-24T03:36:57Z","title":"AI Realtor: Towards Grounded Persuasive Language Generation for Automated Copywriting","version":6},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-05-23T02:55:50.650423Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2502.16810"},"observation_digest":"sha256:8fc3c4a6b9de789acf1ae28506e26c87a6e3fde438edad60ea9dd11b49f70b26","observation_id":"337a7563-7d91-41b5-b62c-e65f2f6bb52e","resolution":{"observed_at":"2026-05-23T02:57:26.367489Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-07T14:52:55.899283Z","title":"Towards Understanding Sycophancy in Language Models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.17380","last_updated":"2025-05-23T01:33:04Z","snapshot_observed_at":"2026-08-08T04:10:31.310818Z","submitted_at":"2025-05-23T01:33:04Z","title":"AI-Augmented LLMs Achieve Therapist-Level Responses in Motivational Interviewing","version":1},"reference_index":138,"source":"pdf_text","source_observed_at":"2026-08-07T14:52:55.899283Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2505.17380"},"observation_digest":"sha256:58a9261e97f3fdc02e37de2c2ac0281f706d96464afa980886c237b5bf5ec753","observation_id":"7fc3778f-ab1c-4727-9ff9-6212de168d44","resolution":{"observed_at":"2026-08-07T14:52:55.899283Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-07T14:28:03.507897Z","title":"arXiv:2310.13548","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.18779","last_updated":"2025-05-24T16:30:53Z","snapshot_observed_at":"2026-08-10T11:14:24.759244Z","submitted_at":"2025-05-24T16:30:53Z","title":"Evaluating Intra-firm LLM Alignment Strategies in Business Contexts","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T14:28:03.507897Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2505.18779"},"observation_digest":"sha256:5a5fbc6b2066fdafa51042e8edc986b23046bdadea4e0d632da1bfea197c41a3","observation_id":"239de3a9-6934-42b7-88d3-c67f103a764f","resolution":{"observed_at":"2026-08-07T14:28:03.507897Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-07T13:43:12.346980Z","title":"https://arxiv","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.21091","last_updated":"2025-06-23T06:43:45Z","snapshot_observed_at":"2026-08-10T22:53:26.697431Z","submitted_at":"2025-05-27T12:19:08Z","title":"Position is Power: System Prompts as a Mechanism of Bias in Large Language Models (LLMs)","version":3},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:12.346980Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2505.21091"},"observation_digest":"sha256:8a105b540023d9bf6af6d24dce96d4f1ff2d2d1b70e836e9c313f8a7f8a3a2f8","observation_id":"db85ade9-021f-4d08-ba0c-b583862b3326","resolution":{"observed_at":"2026-08-07T13:43:12.346980Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-07T05:26:07.247180Z","title":"Bowman, Newton Cheng, Esin Durmus, Zac Hatfield-Dodds, Scott R","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.08120","last_updated":"2025-06-09T18:20:18Z","snapshot_observed_at":"2026-08-07T05:16:51.602131Z","submitted_at":"2025-06-09T18:20:18Z","title":"Conservative Bias in Large Language Models: Measuring Relation Predictions","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T05:26:07.247180Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2506.08120"},"observation_digest":"sha256:62c48da5515af37deac1866cbe7daff4dc81f250bec171e3b9cf0f1275ffa236","observation_id":"376821a5-690f-42ca-9a69-e0cbea111ca5","resolution":{"observed_at":"2026-08-07T05:26:07.247180Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-07T04:36:07.179321Z","title":"Bowman, Newton Cheng, Esin Durmus, Zac Hatfield-Dodds, Scott R","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.10297","last_updated":"2025-06-12T02:21:43Z","snapshot_observed_at":"2026-08-10T09:03:45.646285Z","submitted_at":"2025-06-12T02:21:43Z","title":"\"Check My Work?\": Measuring Sycophancy in a Simulated Educational Context","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T04:36:07.179321Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2506.10297"},"observation_digest":"sha256:847a6e7a0a92a6ee5a9a85b2f6651f5e9e004ab1f56f2cc0e9689fd23480a761","observation_id":"eb3c9b11-77cd-4515-9821-a008d0c291a1","resolution":{"observed_at":"2026-08-07T04:36:07.179321Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-07T05:48:18.743441Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.11110","last_updated":"2025-06-08T14:08:22Z","snapshot_observed_at":"2026-08-07T05:38:21.402950Z","submitted_at":"2025-06-08T14:08:22Z","title":"AssertBench: A Benchmark for Evaluating Self-Assertion in Large Language Models","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T05:48:18.743441Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2506.11110"},"observation_digest":"sha256:53cd52d12e2f128b13433ff2296eecc4475cdbc11cf6f58865469b04c4c4142a","observation_id":"f80abd4b-c4ec-4e6a-a5db-cc57c0a0b564","resolution":{"observed_at":"2026-08-07T05:48:18.743441Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2506.12605","last_updated":"2026-05-04T21:39:43Z","snapshot_observed_at":"2026-07-06T21:42:20.767332Z","submitted_at":"2025-06-14T19:00:37Z","title":"The Rise of AI Companions: Interaction with AI Companions and Psychological Well-being","version":5},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-19T09:18:54.918640Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2506.12605"},"observation_digest":"sha256:30b16c58e7d1b731146ebce3e93348cccf29ca69e2835cb790bc7b41266fbf00","observation_id":"270d24db-fe13-475b-95c7-9967678625b7","resolution":{"observed_at":"2026-05-19T09:22:16.032256Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-06T23:49:26.199238Z","title":"To- wards understanding sycophancy in language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.16064","last_updated":"2025-06-19T06:42:35Z","snapshot_observed_at":"2026-08-08T11:18:17.446928Z","submitted_at":"2025-06-19T06:42:35Z","title":"Self-Critique-Guided Curiosity Refinement: Enhancing Honesty and Helpfulness in Large Language Models via In-Context Learning","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T23:49:26.199238Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2506.16064"},"observation_digest":"sha256:07ec7e8c917dc3d33204f6d658eeff1daa8863bea7c7a1cc470c2cc7c7b95bb1","observation_id":"8712f5f9-4c4e-492b-8b6d-989ca7ade56e","resolution":{"observed_at":"2026-08-06T23:49:26.199238Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-06T20:25:54.638875Z","title":"Sharma, M","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.03120","last_updated":"2025-07-03T18:57:43Z","snapshot_observed_at":"2026-08-09T13:48:48.934149Z","submitted_at":"2025-07-03T18:57:43Z","title":"How Overconfidence in Initial Choices and Underconfidence Under Criticism Modulate Change of Mind in Large Language Models","version":1},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-06T20:25:54.638875Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2507.03120"},"observation_digest":"sha256:23256421fd3ce3e049f97df89789adcd093a2dd163a8c016fadcd14591341b9d","observation_id":"3bc85edd-9b99-450c-b761-5bdc94cb39b2","resolution":{"observed_at":"2026-08-06T20:25:54.638875Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-06T17:17:24.416683Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.11299","last_updated":"2025-07-20T15:15:56Z","snapshot_observed_at":"2026-08-07T07:35:29.605608Z","submitted_at":"2025-07-15T13:26:49Z","title":"Dr.Copilot: A Multi-Agent Prompt Optimized Assistant for Improving Patient-Doctor Communication in Romanian","version":2},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-06T17:17:24.416683Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2507.11299"},"observation_digest":"sha256:2fa59661f5b0466a6b0e00f5f504ca7b5870940d06be4dac618efae144697415","observation_id":"5778aff4-0b18-4d77-9a21-4daeff189325","resolution":{"observed_at":"2026-08-06T17:17:24.416683Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-06T16:12:50.339267Z","title":"Towards understanding sycophancy in language models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.14293","last_updated":"2025-07-18T18:06:27Z","snapshot_observed_at":"2026-08-06T15:57:20.673152Z","submitted_at":"2025-07-18T18:06:27Z","title":"WebGuard: Building a Generalizable Guardrail for Web Agents","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T16:12:50.339267Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2507.14293"},"observation_digest":"sha256:7c912794923725d52d8e22149929cf468f84826da9e18fe100e17a370bf88079","observation_id":"8649ffeb-edf7-41dd-bf06-e35c8d85d5df","resolution":{"observed_at":"2026-08-06T16:12:50.339267Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-06T15:14:22.951866Z","title":"Mrinank Sharma, Meg Tong, Tomasz Korbak, David Duvenaud, Amanda Askell, Samuel R Bowman, Newton Cheng, Esin Durmus, Zac Hatfield-Dodds, Scott R Johnston, et al","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.16534","last_updated":"2025-07-26T12:33:42Z","snapshot_observed_at":"2026-08-10T02:11:35.629269Z","submitted_at":"2025-07-22T12:44:38Z","title":"Frontier AI Risk Management Framework in Practice: A Risk Analysis Technical Report","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T15:14:22.951866Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2507.16534"},"observation_digest":"sha256:4802453c3fd2d39093af13a84779012972117fcab97d4e6e3dc6178ea14f7e77","observation_id":"1a1a5579-dcf3-487f-b11b-651db1286378","resolution":{"observed_at":"2026-08-06T15:14:22.951866Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-07T00:23:55.651236Z","title":"Towards Understanding Sycophancy in Language Models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.21082","last_updated":"2025-06-17T07:13:35Z","snapshot_observed_at":"2026-08-10T01:01:06.386623Z","submitted_at":"2025-06-17T07:13:35Z","title":"Safety Features for a Centralised AGI Project","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T00:23:55.651236Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2507.21082"},"observation_digest":"sha256:d448b58c282cc1303ff8fc8b11071e929a2587c1afa3f1deb5d2b1792f659173","observation_id":"a0db4843-088f-4a35-bb1b-a0df4e4128de","resolution":{"observed_at":"2026-08-07T00:23:55.651236Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-06T11:04:47.211412Z","title":"Noam Shazeer","venue":null,"work_id":null,"year":1911},"citing_paper":{"arxiv_id":"2507.23170","last_updated":"2025-08-02T05:05:26Z","snapshot_observed_at":"2026-08-11T01:48:11.924653Z","submitted_at":"2025-07-31T00:51:16Z","title":"BAR Conjecture: the Feasibility of Inference Budget-Constrained LLM Services with Authenticity and Reasoning","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T11:04:47.211412Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2507.23170"},"observation_digest":"sha256:52ffc15ce398e444306a600ea90ebc19d5eea469173da0ef5aeffd382e658d5e","observation_id":"3a8320f0-f3eb-4b85-ae25-8c320865a0cf","resolution":{"observed_at":"2026-08-06T11:04:47.211412Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-06T05:05:02.851620Z","title":"Towards understanding sycophancy in language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.02312","last_updated":"2025-08-04T11:28:34Z","snapshot_observed_at":"2026-08-07T17:54:58.202973Z","submitted_at":"2025-08-04T11:28:34Z","title":"A Survey on Data Security in Large Language Models","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-06T05:05:02.851620Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2508.02312"},"observation_digest":"sha256:86b5f5eacf259ebc5d12055d9381c943cc1ef64c258f0b0116f62c4f2df85fec","observation_id":"66f6d03e-1bfd-4b2e-bd6b-bdd19e6b916d","resolution":{"observed_at":"2026-08-06T05:05:02.851620Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T21:48:01.694534Z","title":"Towards understanding sycophancy in language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.07966","last_updated":"2025-08-11T13:26:48Z","snapshot_observed_at":"2026-08-10T08:31:43.576342Z","submitted_at":"2025-08-11T13:26:48Z","title":"Exploring the Challenges and Opportunities of AI-assisted Codebase Generation","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-05T21:48:01.694534Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2508.07966"},"observation_digest":"sha256:a667280fc3d9aeb3d617c049aa24c4fb711873ac6881d78e60e294f051bff5f2","observation_id":"8dc4231d-f604-4f96-8269-22855c942db2","resolution":{"observed_at":"2026-08-05T21:48:01.694534Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T18:48:51.499083Z","title":"Towards understanding sycophancy in lan- guage models.arXiv preprint arXiv:2310.13548, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.15842","last_updated":"2025-08-19T18:20:38Z","snapshot_observed_at":"2026-08-10T19:59:32.377280Z","submitted_at":"2025-08-19T18:20:38Z","title":"Lexical Hints of Accuracy in LLM Reasoning Chains","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-05T18:48:51.499083Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2508.15842"},"observation_digest":"sha256:b17e1e5641ec556f69422542909d02453cf36c8c6dfbb7c0421c944844036493","observation_id":"42bea56d-127a-4cba-8ef5-c45f5814a891","resolution":{"observed_at":"2026-08-05T18:48:51.499083Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T17:26:48.564081Z","title":"Towards understanding sycophancy in language models.arXiv preprint arXiv:2310.13548, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.16318","last_updated":"2025-09-01T08:35:27Z","snapshot_observed_at":"2026-08-10T01:27:53.917270Z","submitted_at":"2025-08-22T11:57:55Z","title":"SATORI: Static Test Oracle Generation for REST APIs","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-05T17:26:48.564081Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2508.16318"},"observation_digest":"sha256:82a4cf282c5fb0e28ecaf2dabddddde763badc7081de467ec01dbb205b0d1a21","observation_id":"d1354682-1165-4851-a828-1c798c759a6d","resolution":{"observed_at":"2026-08-05T17:26:48.564081Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2508.16846","last_updated":"2026-05-04T16:33:58Z","snapshot_observed_at":"2026-08-06T21:23:20.360310Z","submitted_at":"2025-08-23T00:11:00Z","title":"BASIL: Bayesian Assessment of Sycophancy in LLMs","version":6},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-18T21:55:45.714195Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2508.16846"},"observation_digest":"sha256:8a32413cbc20c34240a50f58372aacf07bc9862cb3c9613679000b5df50a36ad","observation_id":"c3238fac-a24c-45a3-994d-19a9a6d254d1","resolution":{"observed_at":"2026-05-18T21:56:51.983450Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T16:57:09.417680Z","title":"Bowman, Newton Cheng, Esin Durmus, Zac Hatfield-Dodds, Scott R","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.17511","last_updated":"2025-08-24T20:23:08Z","snapshot_observed_at":"2026-08-09T21:24:47.928571Z","submitted_at":"2025-08-24T20:23:08Z","title":"School of Reward Hacks: Hacking harmless tasks generalizes to misaligned behavior in LLMs","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-05T16:57:09.417680Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2508.17511"},"observation_digest":"sha256:9729a67699ccce2b4f461dc8f5b8bd933e0406c32363e8fff7478ba363b566f3","observation_id":"2a5b2435-3baf-4cb5-982e-b69d4fdaeaf2","resolution":{"observed_at":"2026-08-05T16:57:09.417680Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2508.18473","last_updated":"2026-04-28T16:17:29Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-08-25T20:39:30Z","title":"Principled Detection of Hallucinations in Large Language Models via Multiple Testing","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-18T20:44:52.898833Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2508.18473"},"observation_digest":"sha256:4d359ac09928af1f1cec7b384a58cc86260b8bcb4d688060eaa09b0b4707737a","observation_id":"6d8cdd3d-85f0-4a46-a83f-2fc5c43bfb96","resolution":{"observed_at":"2026-05-18T20:46:51.597802Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T16:06:50.583217Z","title":"Bowman, Newton Cheng, Esin Durmus, Zac Hatfield-Dodds, Scott R","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.19316","last_updated":"2025-08-26T11:21:27Z","snapshot_observed_at":"2026-08-09T10:26:49.099378Z","submitted_at":"2025-08-26T11:21:27Z","title":"Sycophancy as compositions of Atomic Psychometric Traits","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-05T16:06:50.583217Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2508.19316"},"observation_digest":"sha256:b9562e707c7a24fbb7b4c3b288977d450462909340477e29b228652bbc0bfe39","observation_id":"32b354e8-4517-44c0-9f37-143e6d2a171a","resolution":{"observed_at":"2026-08-05T16:06:50.583217Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T15:53:53.970356Z","title":"Sharma, M","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.19461","last_updated":"2025-08-26T22:29:31Z","snapshot_observed_at":"2026-08-07T19:17:20.900756Z","submitted_at":"2025-08-26T22:29:31Z","title":"Reliable Weak-to-Strong Monitoring of LLM Agents","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-05T15:53:53.970356Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2508.19461"},"observation_digest":"sha256:a778e8ca2a1a9e1005fbfc75fd9500f25327fe8ace4dca8b7c2d549bd142b55e","observation_id":"5f9c4bf3-b53c-40b5-8b7b-a35faa0ef263","resolution":{"observed_at":"2026-08-05T15:53:53.970356Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2509.08010","last_updated":"2026-05-20T11:59:51Z","snapshot_observed_at":"2026-07-31T21:33:16.544387Z","submitted_at":"2025-09-08T16:15:07Z","title":"Measuring and mitigating overreliance to build human-compatible AI","version":2},"reference_index":107,"source":"pdf_text","source_observed_at":"2026-05-21T22:37:37.267715Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2509.08010"},"observation_digest":"sha256:4b2c462aa1cb0d598ecb2c172b6fec3eb1f5f7f6c59b053326ff15d6c9e2fe4c","observation_id":"ee0326bb-ff3d-4e86-ac02-f7e907d7b1ce","resolution":{"observed_at":"2026-05-21T22:40:43.403858Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-04T19:07:41.254077Z","title":"How AI Can Improve Access to Justice","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2509.09467","last_updated":"2025-09-11T13:50:23Z","snapshot_observed_at":"2026-08-07T04:46:04.281292Z","submitted_at":"2025-09-11T13:50:23Z","title":"Inteligencia Artificial jur\\'idica y el desaf\\'io de la veracidad: an\\'alisis de alucinaciones, optimizaci\\'on de RAG y principios para una integraci\\'on responsable","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-04T19:07:41.254077Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2509.09467"},"observation_digest":"sha256:2981de099cb06a5c623fb1a4bd9868728865e6e48504ec523596a05eeec14a7b","observation_id":"be428a8c-4059-4b1e-9c36-fc01e2112aa7","resolution":{"observed_at":"2026-08-04T19:07:41.254077Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-04T18:01:56.854499Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.10297","last_updated":"2025-09-12T14:37:57Z","snapshot_observed_at":"2026-08-09T13:15:44.594959Z","submitted_at":"2025-09-12T14:37:57Z","title":"The Morality of Probability: How Implicit Moral Biases in LLMs May Shape the Future of Human-AI Symbiosis","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-04T18:01:56.854499Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2509.10297"},"observation_digest":"sha256:42e9a3b1b6b377a1d7591173b7047e4122fab1cd017618ac57099bfd6eaa93cb","observation_id":"f4239b79-86f6-44cd-a291-d8b540c36271","resolution":{"observed_at":"2026-08-04T18:01:56.854499Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2509.21979","last_updated":"2026-05-17T21:57:50Z","snapshot_observed_at":"2026-08-10T02:52:37.322657Z","submitted_at":"2025-09-26T07:02:22Z","title":"Benchmarking and Mitigating Sycophancy in Medical Vision Language Models","version":5},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-18T14:03:41.489520Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2509.21979"},"observation_digest":"sha256:60e82152d48f93f7fb36cb8d66007c847897a03562c51ea63aad9cd8d96472ee","observation_id":"4ebdf171-434d-4d73-8962-8d0d6cf5c0bf","resolution":{"observed_at":"2026-05-18T14:06:27.398202Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2509.21979","last_updated":"2026-05-17T21:57:50Z","snapshot_observed_at":"2026-08-10T02:52:37.322657Z","submitted_at":"2025-09-26T07:02:22Z","title":"Benchmarking and Mitigating Sycophancy in Medical Vision Language Models","version":6},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-21T21:29:52.725437Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2509.21979"},"observation_digest":"sha256:c0e6aae6d9ce987418b04a1e8b51316d956a3a253857f1c01f5527651a4ca1e1","observation_id":"99d46a7e-118f-4945-95a2-a49e285fcae9","resolution":{"observed_at":"2026-05-21T21:30:39.306333Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-04T09:15:48.616354Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.16712","last_updated":"2026-06-21T20:56:35Z","snapshot_observed_at":"2026-08-08T19:02:44.134057Z","submitted_at":"2025-10-19T04:51:14Z","title":"The Chameleon Nature of LLMs: Quantifying Multi-Turn Stance Instability in Search-Enabled Language Models","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-04T09:15:48.616354Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2510.16712"},"observation_digest":"sha256:9c1d0767aaa83075b273a2c599274eb8b6728c0893d2d41a57b2e9f03c6e038d","observation_id":"2a5d1655-245d-4254-8ebe-228691202dc9","resolution":{"observed_at":"2026-08-04T09:15:48.616354Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-04T07:21:22.535843Z","title":"Sharma, M","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.26518","last_updated":"2026-06-25T13:34:02Z","snapshot_observed_at":"2026-08-10T01:21:56.171636Z","submitted_at":"2025-10-30T14:11:52Z","title":"Human-AI Complementarity: A Goal for Amplified Oversight","version":2},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-04T07:21:22.535843Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2510.26518"},"observation_digest":"sha256:d17bb7a82bb76a81efd20fe63dc5aaf6ffc645bdf132ef0ed1634d6d601078cb","observation_id":"03515745-9e3e-4822-aeee-bd4f41c47f8a","resolution":{"observed_at":"2026-08-04T07:21:22.535843Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-04T06:48:38.233633Z","title":"Bowman, Newton Cheng, Esin Durmus, Zac Hatfield-Dodds, Scott R","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.10871","last_updated":"2026-01-09T18:47:24Z","snapshot_observed_at":"2026-08-08T04:23:06.609813Z","submitted_at":"2025-11-14T00:55:28Z","title":"From Fact to Judgment: Investigating the Impact of Task Framing on LLM Conviction in Dialogue Systems","version":2},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-04T06:48:38.233633Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2511.10871"},"observation_digest":"sha256:431d34cd9ac1c6f0411a7248d959e7a9832ed7c651b5425eb79646ca6cd429aa","observation_id":"0891fa8f-241d-4331-a987-fdbc95b00893","resolution":{"observed_at":"2026-08-04T06:48:38.233633Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2511.13979","last_updated":"2026-06-17T03:30:51Z","snapshot_observed_at":"2026-08-07T16:11:28.195627Z","submitted_at":"2025-11-17T23:15:50Z","title":"Personality Pairing Improves Human-AI Collaboration","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-17T20:11:20.840369Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2511.13979"},"observation_digest":"sha256:8bb30f67fc5d8078833d85c0a4a3deedba19f3f7d4808f0c2aeaeea8c1e6a1a6","observation_id":"46f8843d-72b0-4adc-8b0a-efc6f2dab12c","resolution":{"observed_at":"2026-05-17T20:12:04.096252Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-03T21:43:43.573897Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2511.13979","last_updated":"2026-06-17T03:30:51Z","snapshot_observed_at":"2026-08-07T16:11:28.195627Z","submitted_at":"2025-11-17T23:15:50Z","title":"Personality Pairing Improves Human-AI Collaboration","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-03T21:43:43.573897Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2511.13979"},"observation_digest":"sha256:f0c5087309d073f8984f755d179bb612cf4bc80b1fe640f84869d586ab811ec1","observation_id":"803b6915-7f06-4cfc-994b-466858ea8bd7","resolution":{"observed_at":"2026-08-03T21:43:43.573897Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2511.17408","last_updated":"2026-04-19T21:58:15Z","snapshot_observed_at":"2026-08-10T23:27:19.301813Z","submitted_at":"2025-11-21T17:08:48Z","title":"The Impact of Off-Policy Training Data on Probe Generalisation","version":4},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-05-17T20:26:37.914522Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2511.17408"},"observation_digest":"sha256:3bf6789bdd00642953ed2e5d830df84e1bd111920473e6e3e803c610ba31b386","observation_id":"c1cba9f9-b963-402d-8ddd-a3f1ccecf800","resolution":{"observed_at":"2026-05-17T20:30:11.664946Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-03T20:37:41.534042Z","title":"InProceedings of the 12th International Conference on Learning Repre- sentations (ICLR 2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.19166","last_updated":"2026-07-21T15:21:25Z","snapshot_observed_at":"2026-08-03T20:37:36.623362Z","submitted_at":"2025-11-24T14:28:50Z","title":"Epistemic Familiarity is Associated With Belief Stability in Large Language Models","version":4},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-03T20:37:41.534042Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2511.19166"},"observation_digest":"sha256:280912c9abef444e699227abb8022cb76026654cc4814a091eb1f4739594a2a3","observation_id":"31237818-e317-4df1-97e4-23fbc84f44e6","resolution":{"observed_at":"2026-08-03T20:37:41.534042Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2512.10687","last_updated":"2026-04-20T11:46:57Z","snapshot_observed_at":"2026-08-02T16:01:53.254884Z","submitted_at":"2025-12-11T14:34:40Z","title":"Safe for Whom? Rethinking How We Evaluate the Safety of LLMs for Real Users","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-16T23:22:42.431997Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2512.10687"},"observation_digest":"sha256:b67a38c5f262b996f6534400896bbc8e4cc44f1616d21d887dca15c442f9a205","observation_id":"b3885686-e2a0-4dc6-b0b2-05585f67d104","resolution":{"observed_at":"2026-05-16T23:23:39.919584Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-03T16:56:19.370828Z","title":"Towards understanding sycophancy in language models.arXiv preprint arXiv:2310.13548, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2512.11935","last_updated":"2026-07-07T18:58:00Z","snapshot_observed_at":"2026-08-09T04:14:46.193894Z","submitted_at":"2025-12-12T06:28:28Z","title":"AGAPI-Agents: An Open-Access Agentic AI Platform for Accelerated Materials Design on AtomGPT.org","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-03T16:56:19.370828Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2512.11935"},"observation_digest":"sha256:a398f22d99272cce65e81b945c1ac094643a7c8dce64195e2b73c6916b6acf3c","observation_id":"d9bc3041-4c5d-4fbb-9c91-7290b4934175","resolution":{"observed_at":"2026-08-03T16:56:19.370828Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2601.10467","last_updated":"2026-05-05T15:17:14Z","snapshot_observed_at":"2026-08-02T01:41:38.660629Z","submitted_at":"2026-01-15T14:51:50Z","title":"User Detection and Response Patterns of Sycophantic Behavior in Conversational AI","version":4},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-05-16T14:03:49.595868Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2601.10467"},"observation_digest":"sha256:0c655ca9077f4b5140f083dddb4c8b459addf658dff1651c2a79808664b96399","observation_id":"bfd21e6e-8b8d-49bf-a539-0124639fe5d9","resolution":{"observed_at":"2026-05-16T14:07:58.781971Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-03T09:37:08.296310Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.13433","last_updated":"2026-05-29T00:45:48Z","snapshot_observed_at":"2026-08-03T09:37:04.326965Z","submitted_at":"2026-01-19T22:37:30Z","title":"Who Endorsed It? Measuring Authority Bias Across Expertise Levels in Language Models","version":4},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-03T09:37:08.296310Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2601.13433"},"observation_digest":"sha256:baa7a68c50db09f7a17c1f76fbe648121c80ce0e6b7f82657c5d55865945341d","observation_id":"bfb08858-2f46-4e35-a843-d7861f6d9c5e","resolution":{"observed_at":"2026-08-03T09:37:08.296310Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2601.15395","last_updated":"2026-05-06T18:07:47Z","snapshot_observed_at":"2026-07-06T22:42:36.961014Z","submitted_at":"2026-01-21T19:06:50Z","title":"Beyond Fixed Psychological Personas: State Beats Trait, but Language Models are State-Blind","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-16T12:05:28.383914Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2601.15395"},"observation_digest":"sha256:95b40540c3c1793705f554e93b733dbfcceb775ccd67f2de3b619a4d8f9bb0c2","observation_id":"3c0b1a5a-538f-4732-a27a-ee93e9bfc2c5","resolution":{"observed_at":"2026-05-16T12:07:50.633552Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2601.21350","last_updated":"2026-05-16T02:45:10Z","snapshot_observed_at":"2026-07-06T22:43:26.250071Z","submitted_at":"2026-01-29T07:18:45Z","title":"Factored Causal Representation Learning for Robust Reward Modeling in RLHF","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-21T14:18:33.768962Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2601.21350"},"observation_digest":"sha256:19f29ee0ce6fe19589936b1b784c0116a1dcd786b94e68716957e56a82cd9d19","observation_id":"2909f4c8-9d4d-4dc3-a6f2-3afebc61b437","resolution":{"observed_at":"2026-05-21T14:20:13.640388Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-03T05:52:46.614826Z","title":"Sharma, M., Tong, M., Korbak, T., Duvenaud, D., Askell, A., Bowman, S","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2602.01146","last_updated":"2026-06-02T20:45:51Z","snapshot_observed_at":"2026-08-09T04:39:20.157343Z","submitted_at":"2026-02-01T10:44:58Z","title":"PersistBench: When Should Long-Term Memories Be Forgotten by LLMs?","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-03T05:52:46.614826Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2602.01146"},"observation_digest":"sha256:da37ea34fbcb1adf41985f5e5f53894b6d25059814b16e4a23b531a3360836ae","observation_id":"45356c22-1b74-4fb5-b8e7-7bcf0775388c","resolution":{"observed_at":"2026-08-03T05:52:46.614826Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-03T04:24:14.673016Z","title":"Towards understanding sycophancy in language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.05088","last_updated":"2026-07-07T19:40:51Z","snapshot_observed_at":"2026-08-09T14:06:15.094455Z","submitted_at":"2026-02-04T22:17:04Z","title":"AI Chatbot Suicide Risk Detection and Response: Human Validation Study of the Open-Source VERA-MH Safety Evaluation","version":4},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-03T04:24:14.673016Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2602.05088"},"observation_digest":"sha256:02d605bebeb3bda1e8446079240301dc7cc3a844f61a644718f26300ea963a45","observation_id":"c8492062-8ee0-4b1d-bd58-050ae2aa5ff2","resolution":{"observed_at":"2026-08-03T04:24:14.673016Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-03T03:04:44.296777Z","title":"Towards understanding sycophancy in language models.arXiv preprint arXiv:2310.13548,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.09305","last_updated":"2026-06-28T18:13:37Z","snapshot_observed_at":"2026-08-08T13:06:41.126177Z","submitted_at":"2026-02-10T00:45:24Z","title":"Reward Modeling for Reinforcement Learning-Based LLM Reasoning: Design, Challenges, and Evaluation","version":2},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-03T03:04:44.296777Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2602.09305"},"observation_digest":"sha256:57ab65f556fa97250a3dac634bc029a8d8d2cf652365b61734b459c0745b7443","observation_id":"c76cc6e4-a9d4-42cd-841e-ba0a9a5a5ee1","resolution":{"observed_at":"2026-08-03T03:04:44.296777Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2602.15037","last_updated":"2026-01-29T06:13:44Z","snapshot_observed_at":"2026-08-07T10:11:21.061680Z","submitted_at":"2026-01-29T06:13:44Z","title":"CircuChain: Disentangling Competence and Compliance in LLM Circuit Analysis","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-16T10:05:29.093803Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2602.15037"},"observation_digest":"sha256:081c5a48ec285404638d02c60e74b843e8a8a12c7cac34276510bd5f718714cd","observation_id":"42194197-9284-4865-bb43-169eaea4c12b","resolution":{"observed_at":"2026-05-16T10:07:43.274770Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-07-15T12:59:04.484292Z","title":"Towards understanding sycophancy in language models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2603.13356","last_updated":"2026-07-14T13:24:15Z","snapshot_observed_at":"2026-08-08T07:22:43.777479Z","submitted_at":"2026-03-09T01:35:37Z","title":"Learning When to Trust in Contextual Social Bandits","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-15T12:59:04.484292Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2603.13356"},"observation_digest":"sha256:7f98afef87a3426f44ebc6436383f3cfc5d056d8bfc73957086eaa22465df735","observation_id":"ff1ecbac-5d5c-487b-875f-5e4162f64762","resolution":{"observed_at":"2026-07-15T12:59:04.484292Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2603.18373","last_updated":"2026-06-01T01:48:27Z","snapshot_observed_at":"2026-08-11T02:21:05.058789Z","submitted_at":"2026-03-19T00:15:05Z","title":"To See or To Please: Uncovering Visual Sycophancy and Split Beliefs in VLMs","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-15T09:18:34.853946Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2603.18373"},"observation_digest":"sha256:0d8458d9ccce038d9db942b982b44f10e818c78d633c7717e93d494a7142ed93","observation_id":"d2454c07-ad5e-4e22-9d1a-dd587c19482b","resolution":{"observed_at":"2026-05-15T09:19:53.751752Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2603.21735","last_updated":"2026-04-16T22:35:22Z","snapshot_observed_at":"2026-08-02T20:44:04.277476Z","submitted_at":"2026-03-23T09:24:56Z","title":"Cognitive Agency Surrender: Defending Epistemic Sovereignty via Scaffolded AI Friction","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-15T00:56:27.034133Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2603.21735"},"observation_digest":"sha256:6a3c1f278c0fb5d142cadf72001874a7025c912ae088adde9a98458946e8a99e","observation_id":"eec033d5-58d4-4972-b925-5fcc9fc9d991","resolution":{"observed_at":"2026-05-15T00:58:25.108661Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.02359","last_updated":"2026-03-20T04:31:03Z","snapshot_observed_at":"2026-07-31T07:47:28.480497Z","submitted_at":"2026-03-20T04:31:03Z","title":"Using LLM-as-a-Judge/Jury to Advance Scalable, Clinically-Validated Safety Evaluations of Model Responses to Users Demonstrating Psychosis","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-05-15T09:06:33.531027Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.02359"},"observation_digest":"sha256:0cbfefec2e6807bde92020c96c5de544ee23fcd66ef1d3f2924e8b13ec340b7a","observation_id":"86911a3e-9f5e-492f-812b-7ce0a19f4f61","resolution":{"observed_at":"2026-05-15T09:09:53.337194Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.02423","last_updated":"2026-04-02T18:00:14Z","snapshot_observed_at":"2026-08-04T00:46:56.329972Z","submitted_at":"2026-04-02T18:00:14Z","title":"SWAY: A Counterfactual Computational Linguistic Approach to Measuring and Mitigating Sycophancy","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-13T21:44:59.362174Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.02423"},"observation_digest":"sha256:f53050bd2754408a37d558656700e203861b09dd80911209e07475f69d02af54","observation_id":"fc9cfc0b-6e2c-43eb-bc07-88d5224f1da7","resolution":{"observed_at":"2026-05-13T21:48:19.237792Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.02585","last_updated":"2026-04-15T18:00:16Z","snapshot_observed_at":"2026-07-06T22:51:57.889822Z","submitted_at":"2026-04-02T23:42:20Z","title":"Mitigating LLM biases toward spurious social contexts using direct preference optimization","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-13T20:33:04.433907Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.02585"},"observation_digest":"sha256:b1aca837814eee121fde5651dc15be800089d7a8b7c182f36a4333fba56e5a65","observation_id":"bbc844bf-f483-48f4-920d-c939047f33ef","resolution":{"observed_at":"2026-05-13T20:33:14.707792Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.02686","last_updated":"2026-04-03T03:30:34Z","snapshot_observed_at":"2026-08-02T18:52:24.780545Z","submitted_at":"2026-04-03T03:30:34Z","title":"Beyond Semantic Manipulation: Token-Space Attacks on Reward Models","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-13T20:29:31.354743Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.02686"},"observation_digest":"sha256:e043b72a577064e6b0af1622af5d6a8305540618dfc55af6ca3f4b7cec48cdf2","observation_id":"7ba934ae-ebe5-49c8-ba2e-384987ed82ea","resolution":{"observed_at":"2026-05-13T20:33:16.838984Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.03888","last_updated":"2026-04-04T22:51:06Z","snapshot_observed_at":"2026-07-06T22:52:57.182669Z","submitted_at":"2026-04-04T22:51:06Z","title":"PolySwarm: A Multi-Agent Large Language Model Framework for Prediction Market Trading and Latency Arbitrage","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-13T16:55:30.415552Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.03888"},"observation_digest":"sha256:ba599b495249538761ed605b52b86bea7977d4490ae64a2bcf7bd2659eeef3f6","observation_id":"3f82efe1-a3b6-49e9-a9ab-5dc91f3a96f5","resolution":{"observed_at":"2026-05-13T17:08:01.608998Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.05274","last_updated":"2026-04-07T00:18:28Z","snapshot_observed_at":"2026-07-06T22:54:04.491386Z","submitted_at":"2026-04-07T00:18:28Z","title":"Simulating the Evolution of Alignment and Values in Machine Intelligence","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-05-10T20:15:46.311347Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.05274"},"observation_digest":"sha256:279fdffab0edaffe783243fa406d82680169f10adca0bf571a145f1ff18fb67f","observation_id":"190b03c6-7768-4af0-a9e8-c3e1cb9ca9c0","resolution":{"observed_at":"2026-05-11T06:26:29.613171Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.05279","last_updated":"2026-04-07T00:28:17Z","snapshot_observed_at":"2026-08-08T04:17:53.842643Z","submitted_at":"2026-04-07T00:28:17Z","title":"Pressure, What Pressure? Sycophancy Disentanglement in Language Models via Reward Decomposition","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-10T20:10:56.036361Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.05279"},"observation_digest":"sha256:e41979cf976ecc4a3142bdc452cca7bd7b042f627243c722149b7324d02a1bde","observation_id":"6a0e2988-53de-4915-ba8f-f4c31a1b6233","resolution":{"observed_at":"2026-05-11T06:26:29.613171Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.06188","last_updated":"2026-02-20T15:48:56Z","snapshot_observed_at":"2026-07-06T22:54:44.656835Z","submitted_at":"2026-02-20T15:48:56Z","title":"LLM Spirals of Delusion: A Benchmarking Audit Study of AI Chatbot Interfaces","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-15T20:39:14.898544Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.06188"},"observation_digest":"sha256:23457374b22ed7caec6492e12f3e2546df5bf9516069055e4d838616bb8e5628","observation_id":"1aeb082e-8208-4d37-9cf0-b414fbd7f6f6","resolution":{"observed_at":"2026-05-15T20:40:18.825441Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.07369","last_updated":"2026-04-07T05:29:13Z","snapshot_observed_at":"2026-07-06T22:55:37.787732Z","submitted_at":"2026-04-07T05:29:13Z","title":"The Role of Emotional Stimuli and Intensity in Shaping Large Language Model Behavior","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-10T20:10:12.204925Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.07369"},"observation_digest":"sha256:b7720a586cf9c387cdb5baad82717f6034ede82afe7cf6d07e1a2ff157658d41","observation_id":"a7148b16-3cd2-4c84-b14f-6cbbb7ae4bed","resolution":{"observed_at":"2026-05-11T06:26:29.613171Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.07667","last_updated":"2026-04-09T00:15:20Z","snapshot_observed_at":"2026-07-06T22:55:50.808914Z","submitted_at":"2026-04-09T00:15:20Z","title":"From Debate to Decision: Conformal Social Choice for Safe Multi-Agent Deliberation","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-10T18:37:16.493500Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.07667"},"observation_digest":"sha256:2cd4b521a565eb49bdcddb28b13a4ad1ee316c01193dd4ba4ebb5ff83cbaa382","observation_id":"aba39c2d-7340-4312-bc30-3eff64ea94f2","resolution":{"observed_at":"2026-05-11T06:26:29.613171Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.07709","last_updated":"2026-06-03T21:15:24Z","snapshot_observed_at":"2026-07-13T00:19:28.134640Z","submitted_at":"2026-04-09T01:54:33Z","title":"IatroBench: Pre-Registered Evidence of Iatrogenic Harm from AI Safety Measures","version":3},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-05-10T18:25:53.037936Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.07709"},"observation_digest":"sha256:906ab5e3173e2271fd13686203f48542d468abe7ae5bc70ec29a53dd9e8bc1d6","observation_id":"ecb60a2f-4361-4b41-a8b2-b009cec711fb","resolution":{"observed_at":"2026-05-11T06:26:29.613171Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-07-13T00:19:33.861692Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2604.07709","last_updated":"2026-06-03T21:15:24Z","snapshot_observed_at":"2026-07-13T00:19:28.134640Z","submitted_at":"2026-04-09T01:54:33Z","title":"IatroBench: Pre-Registered Evidence of Iatrogenic Harm from AI Safety Measures","version":4},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-07-13T00:19:33.861692Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.07709"},"observation_digest":"sha256:9e774ab80d3daa7af362f1bb7288a0fd89e130e0fa368ba0ff5af486e10cf2d7","observation_id":"34e7d28a-adea-47fb-a19e-23d1230e4505","resolution":{"observed_at":"2026-07-13T00:19:33.861692Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.07729","last_updated":"2026-04-09T02:25:17Z","snapshot_observed_at":"2026-08-02T17:43:45.091456Z","submitted_at":"2026-04-09T02:25:17Z","title":"Emotion Concepts and their Function in a Large Language Model","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-05-10T18:03:52.210931Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.07729"},"observation_digest":"sha256:554ee6cf3eff52bfb71e4b299379cfb1ed04098781576352c2e73f400ca06997","observation_id":"6dda1616-b12e-4ee7-910b-fb41ff247f96","resolution":{"observed_at":"2026-05-11T06:26:29.613171Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.08362","last_updated":"2026-05-21T16:20:11Z","snapshot_observed_at":"2026-08-02T15:20:13.753843Z","submitted_at":"2026-04-09T15:26:21Z","title":"Towards Real-world Human Behavior Simulation: Benchmarking Large Language Models on Long-horizon, Cross-scenario, Heterogeneous Behavior Traces","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-10T18:20:56.661012Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.08362"},"observation_digest":"sha256:69d9028bd848f14740c2b658e94c2cb7407f7245f571a3f9035b47aecf7d5612","observation_id":"2d5d9a38-8c4b-4399-82d1-ab7af92bafc3","resolution":{"observed_at":"2026-05-11T06:26:29.613171Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.08362","last_updated":"2026-05-21T16:20:11Z","snapshot_observed_at":"2026-08-02T15:20:13.753843Z","submitted_at":"2026-04-09T15:26:21Z","title":"Towards Real-world Human Behavior Simulation: Benchmarking Large Language Models on Long-horizon, Cross-scenario, Heterogeneous Behavior Traces","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-22T10:30:28.404222Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.08362"},"observation_digest":"sha256:a802b536b2253ed2ad5f6513bcce0fc4aa7cbd3ee210e102f8914be223420d58","observation_id":"4aa36881-5dd1-413d-954d-34eba8517f1e","resolution":{"observed_at":"2026-05-22T10:31:24.855785Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.10733","last_updated":"2026-04-12T17:12:55Z","snapshot_observed_at":"2026-07-06T22:59:18.756089Z","submitted_at":"2026-04-12T17:12:55Z","title":"Too Nice to Tell the Truth: Quantifying Agreeableness-Driven Sycophancy in Role-Playing Language Models","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-05-10T16:05:09.033412Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.10733"},"observation_digest":"sha256:c620d3f0da807b15cbb5f62fd698d15a15a436cb31c0de6290bf696f36a9c270","observation_id":"d4c4c4eb-d159-49af-b47c-525369c7c370","resolution":{"observed_at":"2026-05-11T09:21:00.849707Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.12076","last_updated":"2026-04-13T21:29:46Z","snapshot_observed_at":"2026-07-06T23:00:22.690154Z","submitted_at":"2026-04-13T21:29:46Z","title":"Narrative over Numbers: The Identifiable Victim Effect and its Amplification Under Alignment and Reasoning in Large Language Models","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-10T15:14:37.899736Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.12076"},"observation_digest":"sha256:b7726f238529760eeb62afc4e43579a239dc032c43a3398917a231601db271dd","observation_id":"80b312db-b4a0-448d-b104-7aba98944dcd","resolution":{"observed_at":"2026-05-11T10:56:06.024042Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.12138","last_updated":"2026-07-08T18:30:50Z","snapshot_observed_at":"2026-08-03T13:05:28.864307Z","submitted_at":"2026-04-13T23:39:39Z","title":"Retrieval-Augmented Generation Must Move Beyond Factual Grounding to Represent Diverse Opinions","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-10T15:00:05.173476Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.12138"},"observation_digest":"sha256:c0a2dcb01c54258a64b178e92fad9a5cf9dde07965ec126256a0324054230c35","observation_id":"fafac062-eb0a-42f3-8512-5e00c33cfbfa","resolution":{"observed_at":"2026-05-11T11:21:03.628157Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.13803","last_updated":"2026-04-15T12:38:51Z","snapshot_observed_at":"2026-07-06T23:01:41.643335Z","submitted_at":"2026-04-15T12:38:51Z","title":"Gaslight, Gatekeep, V1-V3: Early Visual Cortex Alignment Shields Vision-Language Models from Sycophantic Manipulation","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-05-10T13:09:35.407790Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.13803"},"observation_digest":"sha256:4cad7befb00d9b23f538df851f6ff54b6ea7eb19d06d6a5ed33f1b20244a3548","observation_id":"9e10d53b-d2fb-4149-bac9-e7aa47795f4a","resolution":{"observed_at":"2026-05-11T06:26:29.613171Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.15316","last_updated":"2026-03-01T21:55:58Z","snapshot_observed_at":"2026-07-06T23:02:53.677687Z","submitted_at":"2026-03-01T21:55:58Z","title":"Anthropomorphism and Trust in Human-Large Language Model interactions","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-15T17:41:35.364047Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.15316"},"observation_digest":"sha256:05c460582f29df5b7af4471d2a64cadfe02ceb48fbf47c67d34e092e3f6703dd","observation_id":"f1e0b457-8a5d-445e-b644-b80927ac467c","resolution":{"observed_at":"2026-05-15T17:46:23.858901Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.16320","last_updated":"2026-02-24T19:07:25Z","snapshot_observed_at":"2026-07-06T23:03:39.351412Z","submitted_at":"2026-02-24T19:07:25Z","title":"How Robustly do LLMs Understand Execution Semantics?","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-15T19:45:34.338092Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.16320"},"observation_digest":"sha256:787be54e86740de6dcb42188568d0c544806e3f304bfa6e32d3d4d6cb11255ed","observation_id":"7f269f2d-527d-4e63-97c8-cbe6410b175a","resolution":{"observed_at":"2026-05-15T19:46:33.609520Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.16399","last_updated":"2026-04-30T14:42:44Z","snapshot_observed_at":"2026-08-03T21:23:43.365960Z","submitted_at":"2026-03-31T09:48:09Z","title":"IACDM: Interactive Adversarial Convergence Development Methodology -- A Structured Framework for AI-Assisted Software Development","version":2},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-05-13T23:58:51.460784Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.16399"},"observation_digest":"sha256:af7d76a7233737301101816701977b3d37d70cfa792d0a2a58e3ca23e6e18e63","observation_id":"368e4ec1-9fd3-48ae-ba0e-edcc6879929b","resolution":{"observed_at":"2026-05-14T00:03:28.985652Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.16812","last_updated":"2026-04-28T07:38:02Z","snapshot_observed_at":"2026-08-03T08:37:41.234765Z","submitted_at":"2026-04-18T03:50:00Z","title":"Introspection Adapters: Training LLMs to Report Their Learned Behaviors","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-10T07:40:47.445360Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.16812"},"observation_digest":"sha256:f00210e853907d35a90a4e029d6a44df817f45f33a7980cc648d9b615f3b6577","observation_id":"94bd0aa9-772e-46fc-9a2c-2a5b7bae2160","resolution":{"observed_at":"2026-05-11T06:26:29.613171Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.16913","last_updated":"2026-04-18T08:46:37Z","snapshot_observed_at":"2026-07-06T23:04:10.324443Z","submitted_at":"2026-04-18T08:46:37Z","title":"The Cognitive Penalty: Ablating System 1 and System 2 Reasoning in Edge-Native SLMs for Decentralized Consensus","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T07:36:11.917011Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.16913"},"observation_digest":"sha256:9ce3cda7908f98694170d72eef338f46a46ed74317f86308f35fbc7022134e21","observation_id":"d7fd0975-c1f1-40b3-b453-8eb96578be2e","resolution":{"observed_at":"2026-05-11T06:26:29.613171Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.17596","last_updated":"2026-04-19T20:04:02Z","snapshot_observed_at":"2026-07-06T23:04:41.564912Z","submitted_at":"2026-04-19T20:04:02Z","title":"Terminal Wrench: A Dataset of 331 Reward-Hackable Environments and 3,632 Exploit Trajectories","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-10T05:48:44.687520Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.17596"},"observation_digest":"sha256:c7ae3a512c80f6f8fddcb40635b37a314fa443bce03072637d3f8b331864e23d","observation_id":"afa3076c-9c1c-48df-a0e2-bffdee48dba6","resolution":{"observed_at":"2026-05-11T06:26:29.613171Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.18803","last_updated":"2026-04-25T21:48:41Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-20T20:21:27Z","title":"LLM-as-Judge Framework for Evaluating Tone-Induced Hallucination in Vision-Language Models","version":3},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-10T05:11:01.039309Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.18803"},"observation_digest":"sha256:dc4d7a17055c845e915eee350e5d4558e4635734c893598366c42226413fa8b4","observation_id":"51f4bf77-3f7f-49b7-98ea-dd33e24ac9c9","resolution":{"observed_at":"2026-05-11T06:26:29.613171Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.18874","last_updated":"2026-04-20T21:53:39Z","snapshot_observed_at":"2026-07-06T23:05:39.724237Z","submitted_at":"2026-04-20T21:53:39Z","title":"How Adversarial Environments Mislead Agentic AI?","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-05-10T04:04:41.152756Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.18874"},"observation_digest":"sha256:3ecee86cee133b4dc064d6d337615ec67de02321717e94630d659fb25be71f39","observation_id":"db8702a5-34e1-4156-9310-9973e2a199af","resolution":{"observed_at":"2026-05-11T12:11:07.943414Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.19139","last_updated":"2026-07-05T07:17:27Z","snapshot_observed_at":"2026-08-02T13:53:17.734488Z","submitted_at":"2026-04-21T06:43:01Z","title":"The Rise of Verbal Tics in Large Language Models: A Systematic Analysis Across Frontier Models","version":2},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-05-10T02:52:49.232194Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.19139"},"observation_digest":"sha256:90f89e2c98383cc077d657ab401680182352c2baa36e7e9012bd238ff77a3199","observation_id":"0357dce7-c074-4ec2-92f8-e5d1fe5e9fd9","resolution":{"observed_at":"2026-05-11T06:26:29.613171Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.19656","last_updated":"2026-04-21T16:45:29Z","snapshot_observed_at":"2026-07-06T23:06:16.972438Z","submitted_at":"2026-04-21T16:45:29Z","title":"Pause or Fabricate? Training Language Models for Grounded Reasoning","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-05-10T03:01:58.366028Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.19656"},"observation_digest":"sha256:73aa31fa13dc29215974a97e6d5e179486b1831ccb1686cc6b25b9964adc802a","observation_id":"06652c79-d9c0-4177-bd92-2f023eb5d49e","resolution":{"observed_at":"2026-05-11T12:46:05.190917Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.20871","last_updated":"2026-03-27T12:52:20Z","snapshot_observed_at":"2026-07-06T23:07:31.559812Z","submitted_at":"2026-03-27T12:52:20Z","title":"M-CARE: Standardized Clinical Case Reporting for AI Model Behavioral Disorders, with a 20-Case Atlas and Experimental Validation","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-14T22:36:22.123368Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.20871"},"observation_digest":"sha256:6e3b23f9c51447146bda57104fffb4312e5baa7feb4400236153e47a7aaefb4d","observation_id":"616817c7-8a79-40f1-b7b1-281dc89b3b70","resolution":{"observed_at":"2026-05-14T22:38:11.332230Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.21139","last_updated":"2026-04-22T23:00:49Z","snapshot_observed_at":"2026-07-06T23:07:47.244208Z","submitted_at":"2026-04-22T23:00:49Z","title":"Slot Machines: How LLMs Keep Track of Multiple Entities","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-09T23:48:36.019590Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.21139"},"observation_digest":"sha256:f32ae0fee35d8cccc57b1867affa16862051de0b2a8082b7b9a2814892422041","observation_id":"a82001df-96a8-489f-8d12-01f8207c9fbb","resolution":{"observed_at":"2026-05-11T14:01:04.052487Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.21564","last_updated":"2026-04-30T16:26:35Z","snapshot_observed_at":"2026-08-10T21:25:26.068842Z","submitted_at":"2026-04-23T11:34:06Z","title":"Measuring Opinion Bias and Sycophancy via LLM-based Persuasion","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-09T21:48:37.940162Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.21564"},"observation_digest":"sha256:d368fea9ea29b039960a43cef1c4ae23b193846b8dcb21e27a8bf9e46e0587ec","observation_id":"bd82c32c-ab5f-4c51-ba38-0597f5bc0a26","resolution":{"observed_at":"2026-05-11T14:26:04.278648Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.22971","last_updated":"2026-04-24T19:32:01Z","snapshot_observed_at":"2026-08-02T06:08:11.565923Z","submitted_at":"2026-04-24T19:32:01Z","title":"Peer Identity Bias in Multi-Agent LLM Evaluation: An Empirical Study Using the TRUST Democratic Discourse Analysis Pipeline","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-08T09:39:20.283495Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.22971"},"observation_digest":"sha256:59da881c17dbf390e8eb1cee1987b0fb9a40f4a8153a5deb1e9e027ad20adaaf","observation_id":"2acfdcf3-3023-4299-bcb6-51a7740ad6e8","resolution":{"observed_at":"2026-05-11T20:16:11.296085Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.23593","last_updated":"2026-04-26T08:03:32Z","snapshot_observed_at":"2026-08-01T01:43:18.294936Z","submitted_at":"2026-04-26T08:03:32Z","title":"When AI reviews science: Can we trust the referee?","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-08T06:19:54.727724Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.23593"},"observation_digest":"sha256:93bb7d0cc6c9f67bc4d8c924468e5889a878f4e738af0d3890cb1ed9278877e9","observation_id":"412de350-5dca-431c-abf1-e16aaa75dea6","resolution":{"observed_at":"2026-05-11T06:26:29.613171Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.24700","last_updated":"2026-04-27T17:04:17Z","snapshot_observed_at":"2026-08-04T04:55:30.006592Z","submitted_at":"2026-04-27T17:04:17Z","title":"Green Shielding: A User-Centric Approach Towards Trustworthy AI","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-08T03:43:54.896449Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.24700"},"observation_digest":"sha256:3353a1f8766ff1419d554e9388f372d806bcc4995bd3d041719a45ab100b068d","observation_id":"5abd395a-d503-441c-9ca6-ea42d583a8af","resolution":{"observed_at":"2026-05-11T21:56:24.178515Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.26561","last_updated":"2026-04-29T11:47:28Z","snapshot_observed_at":"2026-08-08T14:02:33.921341Z","submitted_at":"2026-04-29T11:47:28Z","title":"Preserving Disagreement: Architectural Heterogeneity and Coherence Validation in Multi-Agent Policy Simulation","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-07T12:50:21.666394Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.26561"},"observation_digest":"sha256:18d701d8be385edc6b38f50db1c4da6ab93463294d35c69c024b9bad708d468e","observation_id":"88c45256-f80e-4cd0-adfe-482370cc707d","resolution":{"observed_at":"2026-05-12T09:06:26.615855Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.26965","last_updated":"2026-04-14T16:06:14Z","snapshot_observed_at":"2026-07-06T23:12:29.385105Z","submitted_at":"2026-04-14T16:06:14Z","title":"The Impact of AI-Generated Text on the Internet","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-10T14:01:29.175202Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.26965"},"observation_digest":"sha256:ee94757a099e62c677db6a97c79e72403e8b6aa22c3902b52a316fa6bd813c72","observation_id":"24aee952-ae3d-494a-911c-3ba0851d1e41","resolution":{"observed_at":"2026-05-11T06:26:29.613171Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.27228","last_updated":"2026-04-29T21:50:28Z","snapshot_observed_at":"2026-08-02T10:44:02.888229Z","submitted_at":"2026-04-29T21:50:28Z","title":"When Roles Fail: Epistemic Constraints on Advocate Role Fidelity in LLM-Based Political Statement Analysis","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-07T08:52:10.092292Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.27228"},"observation_digest":"sha256:e8b331482f7159b1b3bfbfb152dbb29a6724fcf12391da62d3f3b3f4c5fdf4cd","observation_id":"860f4f54-7611-47a9-8155-51e8df1a07c4","resolution":{"observed_at":"2026-05-12T09:56:26.600289Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2604.27633","last_updated":"2026-04-30T09:23:00Z","snapshot_observed_at":"2026-07-06T23:13:07.378140Z","submitted_at":"2026-04-30T09:23:00Z","title":"Political Bias Audits of LLMs Capture Sycophancy to the Inferred Auditor","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-07T06:35:59.474625Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2604.27633"},"observation_digest":"sha256:c9acef0188a54d84ed2dc1df2a73e7415462060c252fb886e5f59f1f05b0dd4a","observation_id":"38010849-cd8e-45f3-af73-b71ebf64e3b8","resolution":{"observed_at":"2026-05-12T10:21:28.397042Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2605.00914","last_updated":"2026-04-29T14:33:57Z","snapshot_observed_at":"2026-08-01T02:37:01.871582Z","submitted_at":"2026-04-29T14:33:57Z","title":"The Cost of Consensus: Isolated Self-Correction Prevails Over Unguided Homogeneous Multi-Agent Debate","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-09T20:20:26.414236Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2605.00914"},"observation_digest":"sha256:b7fe55a35682833f978d2f5cbf77f8b1bfe55bd3177e859abea6efb7362c4aca","observation_id":"727d964d-9a83-4401-adf4-e0687ebd2a51","resolution":{"observed_at":"2026-05-11T15:16:10.814712Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"cited_work":{"arxiv_id":"2310.13548","doi":"10.48550/arxiv.2310.13548","metadata_source":"pith","pith_arxiv_id":"2310.13548","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Towards Understanding Sycophancy in Language Models","venue":"cs.CL","work_id":"aeefec9a-6ad5-4743-92b9-de6983895e21","year":2023},"citing_paper":{"arxiv_id":"2605.01302","last_updated":"2026-05-02T07:22:24Z","snapshot_observed_at":"2026-07-06T23:14:33.653260Z","submitted_at":"2026-05-02T07:22:24Z","title":"Beyond Semantic Relevance: Counterfactual Risk Minimization for Robust Retrieval-Augmented Generation","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-05-09T15:04:40.429929Z"},"links":{"cited_paper":"/paper/2310.13548","citing_paper":"/paper/2605.01302"},"observation_digest":"sha256:704d75bdb2e3247520b0bddf3625026003a757012faf1a438ce9db1ac03a0b22","observation_id":"4041dacd-31f7-47a0-a8d6-4443a0c01524","resolution":{"observed_at":"2026-05-11T16:46:10.516783Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T14:50:19.716954+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2310.13548/citation-record","integrity":"/paper/2310.13548/integrity","json":"/paper/2310.13548/citation-record.json","paper":"/paper/2310.13548"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2204.05862","last_updated":"2022-04-12T15:02:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-04-12T15:02:38Z","title":"Training a Helpful and Harmless Assistant with Reinforcement Learning from Human Feedback","version":1},"cited_work":{"arxiv_id":"2204.05862","doi":"10.1016/j.respol.2005.01.014","metadata_source":"pith","pith_arxiv_id":"2204.05862","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Training a Helpful and Harmless Assistant with Reinforcement Learning from Human Feedback","venue":"cs.CL","work_id":"a1f2574b-a899-4713-be60-c87ba332656c","year":2022},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/2204.05862","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:0a3de9c9e0c3204c3d220c34d0a6b47fbc7223389d309f8fd8ce125291478653","observation_id":"06a1e84d-6252-44e3-8c77-7c416cd66f1f","resolution":{"observed_at":"2026-05-11T06:26:29.367756Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.03540","last_updated":"2022-11-11T20:17:18Z","snapshot_observed_at":"2026-07-06T14:15:17.474527Z","submitted_at":"2022-11-04T17:03:49Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","version":2},"cited_work":{"arxiv_id":"2211.03540","doi":"10.48550/arxiv.2211.03540","metadata_source":"pith","pith_arxiv_id":"2211.03540","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Measuring Progress on Scalable Oversight for Large Language Models","venue":"cs.HC","work_id":"e7f92eb1-2050-4e60-bc27-82c94d7694c5","year":2022},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/2211.03540","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:90825fd532b7d3ccb17c7bb57ca38df02cec32e083bb515e91bcbbbd26af896e","observation_id":"b8b401bf-361d-43d4-993c-33504914f439","resolution":{"observed_at":"2026-05-17T15:01:41.422297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"fb1f94d5-452b-45ac-b6c6-b99c3c4d2430","year":2024},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:bfb28c2258e2895d707df093a7c0d52d7d8874dcf67bea426ceb59772ef3fb51","observation_id":"0be202c9-f4f0-46d3-baf7-0ca470f14c8c","resolution":{"observed_at":"2026-05-11T06:26:29.607674Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"cc/paper_files/paper/2017/file/d5e2c0adad503c91f91df240d0cd4e49-Paper.pdf","venue":null,"work_id":"43d953a3-7d29-4b01-92d6-934d4a04d795","year":2017},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:c597023bdcac5ba91c48e47b65b9119829d573992994871e175a0a30bba5b056","observation_id":"8aa8bd7e-9a1d-4eb7-9f4f-b0616826c683","resolution":{"observed_at":"2026-05-11T06:26:29.566711Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.10760","last_updated":"2022-10-19T17:56:10Z","snapshot_observed_at":"2026-07-06T14:07:50.468967Z","submitted_at":"2022-10-19T17:56:10Z","title":"Scaling Laws for Reward Model Overoptimization","version":1},"cited_work":{"arxiv_id":"2210.10760","doi":"10.48550/arxiv.2210.10760","metadata_source":"pith","pith_arxiv_id":"2210.10760","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Laws for Reward Model Overoptimization","venue":"cs.LG","work_id":"0fb09554-8ce6-4366-8b78-f426e318fcfd","year":2022},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/2210.10760","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:5f4bb23e7185f165160fab92a8ebe5c367650ff78b1ddef8ce614a69c7be74b8","observation_id":"3804e699-6c18-420f-a666-694a6bef858b","resolution":{"observed_at":"2026-05-19T09:04:53.407177Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-05-22T21:23:26.501616+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-22T21:23:26.501616+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2209.14375","last_updated":"2022-09-28T19:04:43Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-09-28T19:04:43Z","title":"Improving alignment of dialogue agents via targeted human judgements","version":1},"cited_work":{"arxiv_id":"2209.14375","doi":"10.48550/arxiv.2209.14375","metadata_source":"pith","pith_arxiv_id":"2209.14375","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Improving alignment of dialogue agents via targeted human judgements","venue":"cs.LG","work_id":"6ad5970e-7550-4ae8-a158-7084dec7e3cc","year":2022},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/2209.14375","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:720c0c87c3c18cb44bf7b277b622e492e8be109d799541641192408438e13b84","observation_id":"585d14de-d147-4067-9ef8-6c8578009894","resolution":{"observed_at":"2026-05-14T17:54:02.323766Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"com/1411126","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Podcast episodes between October 2020 and September","venue":null,"work_id":"6eddea94-9256-42d0-a940-b36ffb513b95","year":2020},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:bea04f2fd9fea90caea9710e8f6f774d577fd3f6a45bb5c484d8e9ca44ec4b85","observation_id":"329035fd-2965-4f38-96e5-bb1541504824","resolution":{"observed_at":"2026-05-11T06:26:29.546055Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.15717","last_updated":"2023-05-25T05:00:12Z","snapshot_observed_at":"2026-08-09T05:57:33.635614Z","submitted_at":"2023-05-25T05:00:12Z","title":"The False Promise of Imitating Proprietary LLMs","version":1},"cited_work":{"arxiv_id":"2305.15717","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.15717","snapshot_observed_at":"2026-07-04T05:09:36.906364Z","title":"The False Promise of Imitating Proprietary LLMs","venue":"cs.CL","work_id":"f843d86a-7fd6-4b0d-bf8b-f1ad3fbc3f04","year":2023},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/2305.15717","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:f868fa2774b1bb3cd549d74b7b8a5065f666603109243050d4c6ebbfc3d884e3","observation_id":"7dd26a4f-770f-4ebc-852c-bc243a160bc9","resolution":{"observed_at":"2026-05-18T06:54:31.902059Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2103.03874","last_updated":"2021-11-08T21:30:18Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-03-05T18:59:39Z","title":"Measuring Mathematical Problem Solving With the MATH Dataset","version":2},"cited_work":{"arxiv_id":"2103.03874","doi":"10.48550/arxiv.2103.03874","metadata_source":"pith","pith_arxiv_id":"2103.03874","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Measuring Mathematical Problem Solving With the MATH Dataset","venue":"cs.LG","work_id":"50652ac6-fb7c-4675-a2c2-159c241feb17","year":2021},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/2103.03874","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:9adae6640b69ea3cd135e8fadeaafb699694cbd434eb643cf35d12b5c956c5a0","observation_id":"516070d1-c469-45e9-8dce-0182294a6827","resolution":{"observed_at":"2026-05-11T06:26:29.383217Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-14T18:20:22.649941+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-14T18:20:22.649941+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2212.04717","last_updated":"2023-10-30T05:01:12Z","snapshot_observed_at":"2026-07-06T14:28:36.904351Z","submitted_at":"2022-12-09T08:16:20Z","title":"On the Sensitivity of Reward Inference to Misspecified Human Models","version":2},"cited_work":{"arxiv_id":"2212.04717","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2212.04717","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"On the sensitivity of reward inference to misspecified human models","venue":null,"work_id":"01142b30-00dd-438c-8733-1175bc2bb1e8","year":null},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/2212.04717","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:f5418db185be2c586ea219b92c51cf73b147a043311d9830cbbd7424e2287aa9","observation_id":"c59668c5-e4a3-475b-ae45-4ce0524dd658","resolution":{"observed_at":"2026-05-11T06:26:29.395333Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1705.03551","last_updated":"2017-05-13T21:12:37Z","snapshot_observed_at":"2026-08-02T11:13:42.401488Z","submitted_at":"2017-05-09T21:35:07Z","title":"TriviaQA: A Large Scale Distantly Supervised Challenge Dataset for Reading Comprehension","version":2},"cited_work":{"arxiv_id":"1705.03551","doi":"10.48550/arxiv.1705.03551","metadata_source":"pith","pith_arxiv_id":"1705.03551","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"TriviaQA: A Large Scale Distantly Supervised Challenge Dataset for Reading Comprehension","venue":"cs.CL","work_id":"f20e62ba-6265-4b97-aa8c-ddefaf2f5762","year":2017},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/1705.03551","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:f56d2e18a800ec49baad4ff34de7f25bc9c7aa5b3e02c2721dd986b39c33ac57","observation_id":"948ff31b-5cc2-4470-a352-b23a7baebca8","resolution":{"observed_at":"2026-05-11T15:00:28.883814Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.06452","last_updated":"2024-02-19T14:39:07Z","snapshot_observed_at":"2026-07-29T18:39:02.141391Z","submitted_at":"2023-10-10T09:25:44Z","title":"Understanding the Effects of RLHF on LLM Generalisation and Diversity","version":3},"cited_work":{"arxiv_id":"2310.06452","doi":"10.48550/arxiv.2310.06452","metadata_source":"pith","pith_arxiv_id":"2310.06452","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Understanding the Effects of RLHF on LLM Generalisation and Diversity","venue":"cs.LG","work_id":"13d47639-3500-414f-b6ea-1f277577ad3b","year":2023},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/2310.06452","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:1fc3d1ab25e51d765bdd94a42e434835eb4b08793dfc01b426b6649c43bea1f0","observation_id":"30866190-b7ec-43f6-9755-a980371398a7","resolution":{"observed_at":"2026-05-19T02:34:44.407176Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2022.acl-long.229","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T15:37:20.355771Z","title":"URLhttps://doi.org/10.18653/v1/2022.acl-long.229","venue":"Proceedings of the 60th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","work_id":"414143a1-235c-49ec-a733-5cff2faefa32","year":2022},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:438c26dd3a1798627d662f5270b0b7f0ae597d6b2b75dbe0cfeb636c0c5bb1bc","observation_id":"7074416a-c18e-48bb-8383-81ad7394c936","resolution":{"observed_at":"2026-05-11T06:26:29.315352Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-04T01:08:08.32548+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-04T01:08:08.32548+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1705.04146","last_updated":"2017-10-23T16:45:03Z","snapshot_observed_at":"2026-08-02T16:52:38.363314Z","submitted_at":"2017-05-11T13:04:47Z","title":"Program Induction by Rationale Generation : Learning to Solve and Explain Algebraic Word Problems","version":3},"cited_work":{"arxiv_id":"1705.04146","doi":"10.48550/arxiv.1705.04146","metadata_source":"pith","pith_arxiv_id":"1705.04146","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Program Induction by Rationale Generation : Learning to Solve and Explain Algebraic Word Problems","venue":"cs.AI","work_id":"719ded21-2bd8-46b6-8270-4cfca53de17f","year":2017},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/1705.04146","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:9b37c94614d68d2effa608936ae8b5f1abffaaaa2f9f28f68c735296709cf95e","observation_id":"1206e60f-cd8e-43c5-bc72-4c349ae9991d","resolution":{"observed_at":"2026-05-11T06:26:29.432244Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2112.09332","last_updated":"2022-06-01T19:08:11Z","snapshot_observed_at":"2026-08-07T17:14:39.278754Z","submitted_at":"2021-12-17T05:43:43Z","title":"WebGPT: Browser-assisted question-answering with human feedback","version":3},"cited_work":{"arxiv_id":"2112.09332","doi":"10.48550/arxiv.2112.09332","metadata_source":"pith","pith_arxiv_id":"2112.09332","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"WebGPT: Browser-assisted question-answering with human feedback","venue":"cs.CL","work_id":"e25ef3e1-4848-4cb9-bf28-67a420591165","year":2021},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/2112.09332","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:a039d807d7a98764767451e47946f0dcd5b2f95bd48cd9b385f2c4b44c776563","observation_id":"55ac8bbf-4f4a-4a7c-823d-7cdb49cc5575","resolution":{"observed_at":"2026-05-11T06:26:29.442512Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-04T01:08:09.995583+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-04T01:08:09.995583+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1912.11554","last_updated":"2019-12-24T22:09:36Z","snapshot_observed_at":"2026-07-06T08:46:54.138798Z","submitted_at":"2019-12-24T22:09:36Z","title":"Composable Effects for Flexible and Accelerated Probabilistic Programming in NumPyro","version":1},"cited_work":{"arxiv_id":"1912.11554","doi":"10.48550/arxiv.1912.11554","metadata_source":"pith","pith_arxiv_id":"1912.11554","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Composable Effects for Flexible and Accelerated Probabilistic Programming in NumPyro","venue":"stat.ML","work_id":"a2a03fcf-38c5-4ebe-b5a4-4bc41761dca5","year":2019},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/1912.11554","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:110ee747420e3f4354682045301db5d70d09748c00a7739f1f44e4ae2f18e783","observation_id":"9a677176-86da-4f30-b893-a582fdb0a294","resolution":{"observed_at":"2026-05-11T06:26:29.456900Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-17T23:51:13.105218+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-17T23:51:13.105218+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.11768","last_updated":"2023-07-25T04:01:43Z","snapshot_observed_at":"2026-07-06T15:57:04.900537Z","submitted_at":"2023-07-17T00:54:10Z","title":"Question Decomposition Improves the Faithfulness of Model-Generated Reasoning","version":2},"cited_work":{"arxiv_id":"2307.11768","doi":"10.48550/arxiv.2307.11768","metadata_source":"arxiv_reference","pith_arxiv_id":"2307.11768","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2307.11768 , year=","venue":"arXiv (Cornell University)","work_id":"842cb6c9-8e58-44b7-a293-8b5454f7b29f","year":2023},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/2307.11768","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:567bbc9a81e8050cce674a4c325d8f450a703da539386d96d33a871160cd7871","observation_id":"2d036253-682d-4382-b81e-a9070bdc0173","resolution":{"observed_at":"2026-05-11T06:26:29.467153Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.05802","last_updated":"2022-06-14T01:16:24Z","snapshot_observed_at":"2026-07-06T13:19:58.934755Z","submitted_at":"2022-06-12T17:40:53Z","title":"Self-critiquing models for assisting human evaluators","version":2},"cited_work":{"arxiv_id":"2206.05802","doi":"10.48550/arxiv.2206.05802","metadata_source":"pith","pith_arxiv_id":"2206.05802","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Self-critiquing models for assisting human evaluators","venue":"cs.CL","work_id":"3fcefdd1-22ab-4648-a683-cb1555e7a50e","year":2022},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/2206.05802","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:a5dd911d2a12f30db8007ff68432366d07746a6fec9f6925512e2d1ec48f2f77","observation_id":"846fa01d-f1b0-418d-be07-d9303aab27d5","resolution":{"observed_at":"2026-05-16T20:25:41.981797Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.09288","last_updated":"2023-07-19T17:08:59Z","snapshot_observed_at":"2026-08-07T12:56:43.323460Z","submitted_at":"2023-07-18T14:31:57Z","title":"Llama 2: Open Foundation and Fine-Tuned Chat Models","version":2},"cited_work":{"arxiv_id":"2307.09288","doi":"10.24963/ijcai.2025/706","metadata_source":"pith","pith_arxiv_id":"2307.09288","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama 2: Open Foundation and Fine-Tuned Chat Models","venue":"cs.CL","work_id":"68a5177f-d644-44c1-bd4f-4e5278c22f5d","year":2023},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/2307.09288","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:dfe9a9ac79fd191dce68493befe7e34812842a232c54286a2db1c56dd594ccfe","observation_id":"aa2c2c43-b07c-4c22-9e98-7300e3b40049","resolution":{"observed_at":"2026-05-11T06:26:29.482594Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.04388","last_updated":"2023-12-09T21:25:02Z","snapshot_observed_at":"2026-08-02T07:12:38.105035Z","submitted_at":"2023-05-07T22:44:25Z","title":"Language Models Don't Always Say What They Think: Unfaithful Explanations in Chain-of-Thought Prompting","version":2},"cited_work":{"arxiv_id":"2305.04388","doi":"10.1109/cdics61497.2023.00014","metadata_source":"pith","pith_arxiv_id":"2305.04388","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Language Models Don't Always Say What They Think: Unfaithful Explanations in Chain-of-Thought Prompting","venue":"cs.CL","work_id":"6ed38946-7275-41a4-91b9-b9f7fa043250","year":2023},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"cited_paper":"/paper/2305.04388","citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:68dd84a01aada0ece3a6b36bdbff0693a6e6cd7a216c7a8e89a5ab5bed21fdbf","observation_id":"51713a0b-5d23-48ae-9d8f-371fd95ed420","resolution":{"observed_at":"2026-05-15T12:01:19.882845Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Are you sure?","venue":null,"work_id":"443d75e6-c206-4caf-a192-d0a4bb0aceba","year":2024},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:ba6c8ebd2a39e8db47423e9d07a3194ec9b0c4a20184f6841d6034cd7041ab41","observation_id":"2a4d8225-f3e7-4eab-b18a-033f1dec01a5","resolution":{"observed_at":"2026-05-11T06:26:29.600023Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"{first_comment}","venue":null,"work_id":"1880f77e-2ae4-42ef-affa-0a5acc9fcc49","year":2024},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:e6bfc4e10b4fe106cb9ad2ede749e7831bc946571dfaf552f16c53af42c0fe0a","observation_id":"abe62484-3389-4ea4-aebf-b38dc3bdaa19","resolution":{"observed_at":"2026-05-11T06:26:29.603979Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Are you sure?","venue":null,"work_id":"8c1021b5-aec9-4791-b642-87e72ae00801","year":2024},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:be1c7d4195acff454796090fbb20baba77e4f1a7aee3c4ceeb466add891e89ce","observation_id":"b0a2b541-02a2-492f-b94b-3bc4f3ce6992","resolution":{"observed_at":"2026-05-11T06:26:29.611660Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"13952a19-7c4f-442c-8fb6-2e0efe606ca1","year":2024},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:90c0c59f13b92db681767b374c19cb377aa10518d5f31903064f985796a4277d","observation_id":"8e2cefbd-793e-45bc-95b0-de7e620c60b6","resolution":{"observed_at":"2026-05-11T06:26:29.570863Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Are you sure?","venue":null,"work_id":"b0f707fa-f27c-44e8-a0fc-7af9a8c7a456","year":2024},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:eac3720db9b39777ed21fbe7a49b0909c44786257931e6faf2d5d4c4708a3677","observation_id":"df2167b2-09df-4b01-a6eb-7681920e5d7d","resolution":{"observed_at":"2026-05-11T06:26:29.574724Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"dd196c37-85af-4d94-9a4f-938a51bd1f08","year":2024},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:3607c74762e27a62a05e69690ceb81c411b4d2f773476357c6e6a2647fd5d0db","observation_id":"aad83b09-4684-431f-9792-04bbd2442d69","resolution":{"observed_at":"2026-05-11T06:26:29.580821Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Are you sure?","venue":null,"work_id":"6d3502cd-26cd-4f62-ac1a-fe514ca0d121","year":2024},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:c8ea265b09bd902bfd0dc4f7b800b0389368e750d638fa7471fd3ade5af8d05b","observation_id":"850458c2-73fe-401d-901b-33be83fa375b","resolution":{"observed_at":"2026-05-11T06:26:29.584801Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"matches a user’s beliefs, biases, and preferences","venue":null,"work_id":"15e8381f-f942-48d4-b84f-8d3919a0fece","year":2011},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:98848ac33a3815f0a492b4d32ad884f8e42ac2539bba8794bc795f37d4557e17","observation_id":"12408e6b-e10f-41d1-a24a-810c85222fb3","resolution":{"observed_at":"2026-05-11T06:26:29.590809Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"the Earth’s crust is a solid, unbroken shell","venue":null,"work_id":"652ba314-50b4-4dd6-b814-4a07a8193635","year":2022},"citing_paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models","version":4},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-11T06:26:29.196349Z"},"links":{"citing_paper":"/paper/2310.13548"},"observation_digest":"sha256:19aa92524509cb95a5b57cf7c7812567be10664661a8b5770df783301f25624c","observation_id":"2ab6cb47-cf75-4bb0-9880-678c3a64291e","resolution":{"observed_at":"2026-05-11T06:26:29.595793Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2310.13548","last_updated":"2025-05-10T07:10:46Z","latest_version":4,"primary_category":"cs.CL","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-20T14:46:48Z","title":"Towards Understanding Sycophancy in Language Models"},"reference_resolution":{"displayed":29,"state_counts":{"malformed_identifier":0,"metadata_mismatch":4,"parse_uncertain":0,"unresolved":3,"verified_exact":14,"verified_fuzzy":8},"total_outbound_references":29},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 11 August 2026, this Paper Citation Record lists 29 of 29 outbound references and 100 inbound Pith citation observations for arXiv:2310.13548."}