{"as_of":"2026-08-09T03:34:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:eb77a4f5aade3734ba072e4a92f651a2f619df225a90723b19bc4aaa6b5c7766","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":22,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":22,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":22,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":22,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T11:26:01.891765Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":9,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2405.14782","last_updated":"2026-05-31T00:04:33Z","snapshot_observed_at":"2026-08-07T18:21:12.072228Z","submitted_at":"2024-05-23T16:50:49Z","title":"Lessons from the Trenches on Reproducible Evaluation of Language Models","version":2},"reference_index":255,"source":"arxiv_source","source_observed_at":"2026-05-16T18:44:49.519995Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2405.14782"},"observation_digest":"sha256:dcabcd6e0fff40a0c6bb74c06043961dbd451ad96d1f832e9b0b6b60469f0a36","observation_id":"4027fdd0-17f8-4fbb-868a-7ff070657a16","resolution":{"observed_at":"2026-05-16T18:44:49.863860Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2407.21772","last_updated":"2024-08-04T22:13:39Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-07-31T17:48:14Z","title":"ShieldGemma: Generative AI Content Moderation Based on Gemma","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-20T13:17:39.444002Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2407.21772"},"observation_digest":"sha256:5d85b33f66272f35aab4e33059a909478cb252efe4d7e072dc74c64af3ed3850","observation_id":"7dd4cb82-3d90-4cf0-9336-09fe5bc730bc","resolution":{"observed_at":"2026-05-20T13:17:39.478203Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2410.20791","last_updated":"2026-04-06T19:29:25Z","snapshot_observed_at":"2026-08-03T01:41:19.546733Z","submitted_at":"2024-10-28T07:16:00Z","title":"From Cool Demos to Production-Ready FMware: Core Challenges and a Technology Roadmap","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-23T19:07:21.016824Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2410.20791"},"observation_digest":"sha256:17a9cbd026bc3c91c9201aa2167332ed7d7a44a75de6aa2899f69adee1005648","observation_id":"75fc38b4-ebb7-48a1-9090-e55fc7015d55","resolution":{"observed_at":"2026-05-23T19:08:20.821090Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2411.15594","last_updated":"2025-10-19T10:32:43Z","snapshot_observed_at":"2026-08-02T10:23:50.881300Z","submitted_at":"2024-11-23T16:03:35Z","title":"A Survey on LLM-as-a-Judge","version":6},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-23T17:33:13.394338Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2411.15594"},"observation_digest":"sha256:c803c86c0dcd378e1cc56071133cfd1015308a37fd5147cd0b3fc7c9f53d8a1a","observation_id":"d797e876-a248-4128-84ec-4e6e856441d3","resolution":{"observed_at":"2026-05-23T17:35:44.180203Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2412.05579","last_updated":"2024-12-10T05:49:12Z","snapshot_observed_at":"2026-07-31T01:42:39.468673Z","submitted_at":"2024-12-07T08:07:24Z","title":"LLMs-as-Judges: A Comprehensive Survey on LLM-based Evaluation Methods","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-11T23:08:34.312466Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2412.05579"},"observation_digest":"sha256:f47bd0a86abcdf2d72a631ee361e539f5406f9210d1eacc2db8533da01542540","observation_id":"7099f953-cf2c-41a1-b988-756f77c873d8","resolution":{"observed_at":"2026-05-11T23:08:36.091045Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-07T11:26:01.891765Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.02592","last_updated":"2025-06-03T08:12:47Z","snapshot_observed_at":"2026-08-07T15:19:01.070855Z","submitted_at":"2025-06-03T08:12:47Z","title":"Beyond the Surface: Measuring Self-Preference in LLM Judgments","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T11:26:01.891765Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2506.02592"},"observation_digest":"sha256:5af3d334d9c4a4e0983eddc1cddb01eabc8d62a76cfd2074bf7fff47b7b2944a","observation_id":"2a704e88-0a57-490d-aa55-a99138c00bb1","resolution":{"observed_at":"2026-08-07T11:26:01.891765Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-07T05:34:26.647558Z","title":"Humans or llms as the judge? a study on judgement biases","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07673","last_updated":"2026-07-21T11:15:43Z","snapshot_observed_at":"2026-08-09T03:05:20.348569Z","submitted_at":"2025-06-09T11:50:41Z","title":"How Benchmark Prediction from Fewer Data Misses the Mark","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T05:34:26.647558Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2506.07673"},"observation_digest":"sha256:c2788f1ec0672db36a2b9152b75032805d594544f7f8d084a97bb4089d383a1c","observation_id":"0696c877-6cb1-450c-898d-563b3cd76d5b","resolution":{"observed_at":"2026-08-07T05:34:26.647558Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2507.22359","last_updated":"2026-04-14T11:47:19Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-07-30T03:50:46Z","title":"League of LLMs: A Benchmark-Free Paradigm for Mutual Evaluation of Large Language Models","version":4},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-05-19T03:17:06.457421Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2507.22359"},"observation_digest":"sha256:3e10f723affaa4d280b3d0c4f110de8727db2092c7ee17e2563161ae322d292e","observation_id":"3212f0dd-80e7-48af-93e1-3ae1de774f2c","resolution":{"observed_at":"2026-05-19T03:22:01.345621Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T23:04:56.860420Z","title":", author Chen, S","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.05929","last_updated":"2025-09-10T12:09:14Z","snapshot_observed_at":"2026-08-05T23:04:47.746625Z","submitted_at":"2025-08-08T01:40:10Z","title":"Towards Reliable Generative AI-Driven Scaffolding: Reducing Hallucinations and Enhancing Quality in Self-Regulated Learning Support","version":2},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-05T23:04:56.860420Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2508.05929"},"observation_digest":"sha256:d22d00c5df45d94c2f67a9bcd9cf06730718e72962714f709b21b9eac2f91052","observation_id":"adde8bc6-e443-4c2b-9b0a-7bb1bc90f9fa","resolution":{"observed_at":"2026-08-05T23:04:56.860420Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T21:55:56.852022Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.07805","last_updated":"2025-08-11T09:45:02Z","snapshot_observed_at":"2026-08-07T10:05:08.818969Z","submitted_at":"2025-08-11T09:45:02Z","title":"Can You Trick the Grader? Adversarial Persuasion of LLM Judges","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-05T21:55:56.852022Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2508.07805"},"observation_digest":"sha256:744352d65cffbcbd160fac9e7fd12e98ea729988f44f148198f5a70969002655","observation_id":"1133009d-cbdc-4171-9003-eefcab3affe3","resolution":{"observed_at":"2026-08-05T21:55:56.852022Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2511.20284","last_updated":"2026-05-05T13:11:45Z","snapshot_observed_at":"2026-08-03T04:18:43.042605Z","submitted_at":"2025-11-25T13:11:23Z","title":"Can LLMs Make (Personalized) Access Control Decisions?","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-05-17T05:33:18.457421Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2511.20284"},"observation_digest":"sha256:2baf0bf91cba8930052fb878c6cdc5441ed65fedce62b47b78d2e6159e572e46","observation_id":"1e3af1f4-8c56-4cf6-8c82-b198eb344214","resolution":{"observed_at":"2026-05-17T05:34:05.015121Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2603.03332","last_updated":"2026-04-16T23:41:41Z","snapshot_observed_at":"2026-08-04T02:21:24.392691Z","submitted_at":"2026-02-11T03:11:30Z","title":"Fragile Thoughts: How Large Language Models Handle Chain-of-Thought Perturbations","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-16T03:43:18.987241Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2603.03332"},"observation_digest":"sha256:d25e5bf62c90992dff8123e46066da515189af083ae2d455b87eed0fa0bdb756","observation_id":"b0c0797a-c760-4f0d-a1b4-dd183d554350","resolution":{"observed_at":"2026-05-16T03:47:15.241556Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2604.02359","last_updated":"2026-03-20T04:31:03Z","snapshot_observed_at":"2026-07-31T07:47:28.480497Z","submitted_at":"2026-03-20T04:31:03Z","title":"Using LLM-as-a-Judge/Jury to Advance Scalable, Clinically-Validated Safety Evaluations of Model Responses to Users Demonstrating Psychosis","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-15T09:06:33.531027Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2604.02359"},"observation_digest":"sha256:bc2be0c0dd40d9c8a3fda76c07d80e33103c81480f343d1f1c70171845a5891b","observation_id":"bdc06a1f-3152-4c53-b325-23af5498a4cf","resolution":{"observed_at":"2026-05-15T09:09:53.360123Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2604.09791","last_updated":"2026-04-10T18:13:09Z","snapshot_observed_at":"2026-07-06T22:58:34.504325Z","submitted_at":"2026-04-10T18:13:09Z","title":"Pioneer Agent: Continual Improvement of Small Language Models in Production","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-05-10T17:48:40.520740Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2604.09791"},"observation_digest":"sha256:fc9bc753fcbbc4e93d053e83e2a2c7ed6c14f4bdb44da7f5ea352bdaf01c65ac","observation_id":"e11eef20-585c-4cd1-9f30-cc82530e31ea","resolution":{"observed_at":"2026-05-11T06:05:57.610523Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2604.18835","last_updated":"2026-04-20T20:59:25Z","snapshot_observed_at":"2026-08-03T02:44:21.086343Z","submitted_at":"2026-04-20T20:59:25Z","title":"Semantic Needles in Document Haystacks: Sensitivity Testing of LLM-as-a-Judge Similarity Scoring","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-05-10T04:31:53.825854Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2604.18835"},"observation_digest":"sha256:0bdc990505109435d9139bd3e6059f25a1556b58034cd0b1c47cd89cf4eb6478","observation_id":"55cad9a5-2ce2-4824-b350-a07331f019f6","resolution":{"observed_at":"2026-05-11T11:51:03.617526Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2604.21769","last_updated":"2026-04-23T15:28:32Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-23T15:28:32Z","title":"Who Defines \"Best\"? Towards Interactive, User-Defined Evaluation of LLM Leaderboards","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-09T21:58:05.584559Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2604.21769"},"observation_digest":"sha256:2080afcb83a2068b30d212ac5ce480b52178bd4942328c97cc704227618f5ad7","observation_id":"fa516d12-dbb4-4f71-b29a-ef868771e54e","resolution":{"observed_at":"2026-05-11T14:21:07.359059Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2604.23178","last_updated":"2026-06-24T13:27:28Z","snapshot_observed_at":"2026-08-08T12:19:05.889262Z","submitted_at":"2026-04-25T07:18:30Z","title":"Judging the Judges: A Systematic Evaluation of Bias Mitigation Strategies in LLM-as-a-Judge Pipelines","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-08T08:14:18.535385Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2604.23178"},"observation_digest":"sha256:7ff82ee9cf346099985ec38551dbdbc920c7be0e03b78d69f147277bbc66faaa","observation_id":"d52acf4d-da07-433f-a3f5-212bb7de8ab5","resolution":{"observed_at":"2026-05-11T20:41:14.136347Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2604.27132","last_updated":"2026-04-29T19:32:58Z","snapshot_observed_at":"2026-07-06T23:12:38.453388Z","submitted_at":"2026-04-29T19:32:58Z","title":"TRUST: A Framework for Decentralized AI Service v.0.1","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-05-07T08:22:14.239443Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2604.27132"},"observation_digest":"sha256:083700ad2c43668801d0a2a1574435fe4ec63589bc65374c97563e053dc6d9e8","observation_id":"2d1153fc-2934-442f-9d40-a63336ac8edf","resolution":{"observed_at":"2026-05-12T10:01:28.726224Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2605.18805","last_updated":"2026-05-11T18:55:32Z","snapshot_observed_at":"2026-08-01T14:27:48.374407Z","submitted_at":"2026-05-11T18:55:32Z","title":"RecoAtlas: From Semantic Plausibility to Set-Level Utility in LLM Recommendation Agents","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-20T22:27:16.974169Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2605.18805"},"observation_digest":"sha256:b08b02f4b53b4018884dd5a64f5c1380a5e5b3f00a0b8472659888e31917d4b1","observation_id":"4fca8ff2-f0e2-4b94-a225-46b6c1dc4946","resolution":{"observed_at":"2026-05-20T22:29:09.209885Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2606.11635","last_updated":"2026-06-10T03:56:07Z","snapshot_observed_at":"2026-08-08T17:34:26.394532Z","submitted_at":"2026-06-10T03:56:07Z","title":"Are LLMs Bad at Moral Reasoning?","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-27T08:20:24.251540Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2606.11635"},"observation_digest":"sha256:1f529aacdd111f162b0ea1137e8a6d5428e1cffca355c645761e6ab57bf8fa91","observation_id":"a9aa0af7-0d17-4bf2-942b-87235affd6c8","resolution":{"observed_at":"2026-07-03T13:18:12.640230Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2606.12422","last_updated":"2026-05-08T16:32:43Z","snapshot_observed_at":"2026-07-30T06:43:27.382598Z","submitted_at":"2026-05-08T16:32:43Z","title":"Creating and Evaluating K-12 GenAI Assessment Graders Through Context Engineering","version":1},"reference_index":236,"source":"arxiv_source","source_observed_at":"2026-06-30T22:54:03.054871Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2606.12422"},"observation_digest":"sha256:4c570215fd9cd20a33367e5a36ee21b927e1ae2536968500e7bbc18529760472","observation_id":"302284f3-6fdf-4743-b37b-05ebf3e098ae","resolution":{"observed_at":"2026-06-30T22:55:06.221097Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases","version":5},"cited_work":{"arxiv_id":"2402.10669","doi":"10.48550/arxiv.2402.10669","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.10669","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or llms as the judge? a study on judgement biases","venue":"arXiv (Cornell University)","work_id":"f26248cc-fd46-4eb4-94cf-a455fa482305","year":2024},"citing_paper":{"arxiv_id":"2606.24428","last_updated":"2026-06-23T11:05:05Z","snapshot_observed_at":"2026-08-05T04:36:12.703717Z","submitted_at":"2026-06-23T11:05:05Z","title":"Escaping the Self-Confirmation Trap: An Execute-Distill-Verify Paradigm for Agentic Experience Learning","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-06-25T23:49:38.932474Z"},"links":{"cited_paper":"/paper/2402.10669","citing_paper":"/paper/2606.24428"},"observation_digest":"sha256:6b0083f47ae963833c540a156b50dc118b216ec04a854ac7baee472f0f34b8d1","observation_id":"182049fe-ebd1-4305-9170-b82122fce22a","resolution":{"observed_at":"2026-07-04T17:20:00.097000Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2402.10669/citation-record","integrity":"/paper/2402.10669/integrity","json":"/paper/2402.10669/citation-record.json","paper":"/paper/2402.10669"},"outbound":[],"paper":{"arxiv_id":"2402.10669","last_updated":"2024-09-26T03:16:52Z","latest_version":5,"primary_category":"cs.CL","snapshot_observed_at":"2026-07-06T17:31:08.762955Z","submitted_at":"2024-02-16T13:21:06Z","title":"Humans or LLMs as the Judge? A Study on Judgement Biases"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 22 inbound Pith citation observations for arXiv:2402.10669."}