{"as_of":"2026-08-06T06:11:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:738db22ef336e5d6c62c8a9e92e2ae618a8d6780832709b7e69cf979ba4f647c","coverage":[{"denominator":55,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":55,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-10T19:38:11.595077Z","state":"measured"},{"denominator":56,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":56,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-29T05:36:21.058156Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-06-29T05:43:08.668932Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"cited_work":{"arxiv_id":"2604.05593","doi":null,"metadata_source":"pith","pith_arxiv_id":"2604.05593","snapshot_observed_at":"2026-06-29T05:43:08.668932Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","venue":"cs.AI","work_id":"481def1a-ffe3-4028-be2a-02d50fcbded1","year":2026},"citing_paper":{"arxiv_id":"2605.29928","last_updated":"2026-06-03T17:39:58Z","snapshot_observed_at":"2026-08-03T20:43:08.163978Z","submitted_at":"2026-05-28T13:41:43Z","title":"Label Over Logic? How Source Cues Bias Human Fallacy Judgments More Than LLMs","version":2},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-06-29T05:36:21.058156Z"},"links":{"cited_paper":"/paper/2604.05593","citing_paper":"/paper/2605.29928"},"observation_digest":"sha256:c534913cd9dadd3451d79a5697b10eb10ea9da64df024b13ea372bc124d653af","observation_id":"376b07e2-fcd9-4555-b82e-6f56eeb7a457","resolution":{"observed_at":"2026-06-29T05:43:08.670468Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2604.05593/citation-record","integrity":"/paper/2604.05593/integrity","json":"/paper/2604.05593/citation-record.json","paper":"/paper/2604.05593"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"619fa536-aee3-4692-8a9e-1488d8b402f2","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:06385eb279ec75f41c24e56dc63276c40aaef6504ba8fd5fd0c62d8d4e794c67","observation_id":"5a8acef9-8e2b-4f2b-a601-843d8d83fa04","resolution":{"observed_at":"2026-05-16T04:30:37.652023Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"7429ef87-187e-48d2-9ed6-841d4945aabb","year":2006},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:9df40ce40152519b5b51ade19eb2a8cc960f69220bceb70d4f3fcfc901353ed9","observation_id":"fbc64fda-7fd1-4891-bc00-bc92861188fe","resolution":{"observed_at":"2026-05-16T04:30:37.663685Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"53156670-4237-4bc8-8ba5-c211abf7c41f","year":2025},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:6bcd2f5274335417429c6e1eb94b6d0ee8d345a4a80d42805b790a705467e6d3","observation_id":"0f2971c2-a22c-4237-afd8-bf5a00c56cb2","resolution":{"observed_at":"2026-05-16T04:30:37.655760Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Cacioppo, Louis G","venue":null,"work_id":"f0cc763b-ac2c-493b-97f6-1ad7e9cf7ebf","year":2016},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:33def59719c7901582eeca13f9693713cd5aa7e8d8acdb4aba72d3186880509c","observation_id":"08946746-33d5-4b39-9ed2-eed6361133b7","resolution":{"observed_at":"2026-05-16T04:30:37.637626Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2024.emnlp-main.474","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Humans or LLM s as the Judge? A Study on Judgement Bias","venue":null,"work_id":"15d44b6d-5659-44dd-9508-d9df529b958b","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:d57506a9136004499de3a1fa40038039eccf4a2bcdb81e6287596b96b9d7544d","observation_id":"bf718d33-67ed-448e-98b1-47bf3966f37b","resolution":{"observed_at":"2026-05-10T19:40:45.627560Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-15T23:50:20.543814+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-15T23:50:20.543814+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1073/pnas.2412015122","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Proceedings of the National Academy of Scienceshttps://doi.org/10.1073/pnas.2412015122 (2025) doi:10.1073/pnas.2412015122","venue":"Proceedings of the National Academy of Sciences","work_id":"77ba5838-73de-4391-8ef7-326fff5fa74e","year":2025},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:2a106247ffc5ec6200a947238d85bc7a746a0df7e8a712fd20c8ae3f2f4d4325","observation_id":"305a6179-b5ad-4a44-8370-51fe04218837","resolution":{"observed_at":"2026-05-10T19:40:45.624821Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T14:49:07.16447+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T14:49:07.16447+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"7528.367188","doi":"10.1145/3637528.3671880","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"In: Proceedings of the 30th ACM SIGKDD Conference on Knowledge Discovery and Data Mining","venue":null,"work_id":"4ba97c0e-445f-48fc-96e0-8d08c21fce83","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:1fe2d8566db1ac75d689ee682b1ab0ac68b3643e34e28dfab6156de1208ab911","observation_id":"552f7a16-69f6-4ba6-a009-d0645ffc7510","resolution":{"observed_at":"2026-05-10T19:40:45.639719Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"304cf294-1e58-40ab-9fac-cb00df1e554f","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:779aa9cc1ff17fc51c1d29eb81632076c4447dcf4ab33419fa506473b49b6e12","observation_id":"f13d996f-bf89-42ad-b3ea-39c1ba4ac376","resolution":{"observed_at":"2026-05-16T04:30:37.648702Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2024.findings-emnlp.739","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Cognitive Bias in Decision-Making with LLM s","venue":null,"work_id":"b7a75642-bb73-42d7-89aa-e7479cce8197","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:fb2a685fdce7da562205ebca0433f50bbc09cc775f4da6f5672509965ceac63e","observation_id":"aa8216d6-e3c8-4ff3-8e62-994a6fcf227a","resolution":{"observed_at":"2026-05-10T19:40:45.634912Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"3905.365075","doi":"10.1145/3613905.3650750","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"We Need Structured Output","venue":null,"work_id":"c454260b-f502-44bd-8c38-7417c57db341","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:50c8ef03f0acc2d374385de3fe440460b83d68b82642fa53364e687b869641e6","observation_id":"9771e6ba-845e-4c12-b2d6-a503fae7d33b","resolution":{"observed_at":"2026-05-10T19:40:45.632186Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-05-24T00:23:41.798365+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-24T00:23:41.798365+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"38c7ea3d-5847-4df7-8c1e-3b4ff25461f9","year":2007},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:adc181a219cb2953acd5b63c2f42c9b2b729c26cb290856e66684c57bed43824","observation_id":"cf8cad75-8f11-4885-a9c5-461ee9a0fefb","resolution":{"observed_at":"2026-05-16T04:30:37.642170Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"92925a5d-26e6-4032-905b-22cf25888b07","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:48677b2aac77d8879ba464afaf6163bdc8993cc2f56043e0d8fc3a4e46d17b50","observation_id":"6f2ae0d9-334f-441a-b87f-c5eed7a82274","resolution":{"observed_at":"2026-05-16T04:30:37.659841Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"a2ea8703-42ed-43f0-87a8-dacc0a6d9aca","year":1992},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:1b4d7fec209a6c4e2c47b87f4040ac7f76c0b549952cb836c22aad46fd909e05","observation_id":"55e87f8e-ea62-40da-8b71-d0126673ed3e","resolution":{"observed_at":"2026-05-16T04:30:37.645438Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15594","last_updated":"2025-10-19T10:32:43Z","snapshot_observed_at":"2026-08-02T10:23:50.881300Z","submitted_at":"2024-11-23T16:03:35Z","title":"A Survey on LLM-as-a-Judge","version":6},"cited_work":{"arxiv_id":"2411.15594","doi":"10.1016/j.xinn.2025.101253","metadata_source":"pith","pith_arxiv_id":"2411.15594","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Survey on LLM-as-a-Judge","venue":"cs.CL","work_id":"2676656a-67bd-4ad5-bad6-cb6f5fcdbfbe","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"cited_paper":"/paper/2411.15594","citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:29cf776b3ebafb82e472ed8e8b48c13906021e934f898128203d113141598de7","observation_id":"db4b10cf-1fbe-42fc-ae2d-a1756478861e","resolution":{"observed_at":"2026-05-10T22:40:51.786196Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-10T18:20:01.818871+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-10T18:20:01.818871+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2025.findings-emnlp.1361","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rating Roulette: Self-Inconsistency in","venue":null,"work_id":"bea8f652-6b9b-42cf-b680-cd47b2d6f12c","year":2025},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:3681d756c9359d01fdce60b9cf5b14b6f8ad0d45c469895701ba5e5b72c9e182","observation_id":"bf075786-5d17-432b-869a-6fcad4375361","resolution":{"observed_at":"2026-05-10T19:40:45.660492Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"5ae0f71a-4330-41f7-943d-7f721f80fa60","year":2012},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:85ed5be4b84d3b7f45dee91b76f2ce74b183904c0b2346c8649cd6b3b64e85a5","observation_id":"50601554-e28d-432e-96e1-580e2c34ff1c","resolution":{"observed_at":"2026-05-16T04:30:37.660667Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1007/978-1-4419-9863-7","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":null,"venue":null,"work_id":"a4f38e67-fe72-40c0-982a-31bed3cc9efd","year":2013},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:85f03f1d9e7c855a8feb42181796b436b3796dddc6a87a521b6435b458cf6dd1","observation_id":"f4f4b314-965b-4a69-8b55-9fcd77f57253","resolution":{"observed_at":"2026-05-10T19:40:45.677197Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"0605.330046","doi":"10.1145/3290605.3300469","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"InProceedings of the 2019 CHI Conference on Human Factors in Computing Systems","venue":null,"work_id":"cdc31745-84c5-44b8-a438-8dd8886756bb","year":2019},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:b10ed3f8692f88fd90fb879ada8091fbea1a1779e1500b2f024be0876603e244","observation_id":"b5c7d5c1-8204-4339-9c72-88fb52b3e571","resolution":{"observed_at":"2026-05-10T19:40:45.681187Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11191","last_updated":"2024-06-18T08:18:33Z","snapshot_observed_at":"2026-08-03T11:31:07.805515Z","submitted_at":"2024-06-17T03:52:51Z","title":"A Survey on Human Preference Learning for Large Language Models","version":2},"cited_work":{"arxiv_id":"2406.11191","doi":"10.48550/arxiv.2406.11191","metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11191","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A survey on human preference learning for large language models","venue":"arXiv (Cornell University)","work_id":"6c9f1057-69b4-4eb7-9479-b114b8611feb","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"cited_paper":"/paper/2406.11191","citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:01b0ee56ecaf765c63c54f3767fd0b93567f20c10012263f881d5538b5af8afc","observation_id":"a4db0292-b962-4748-b391-63151e5eddc8","resolution":{"observed_at":"2026-05-10T22:40:51.769687Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Johnson, Jennifer E","venue":null,"work_id":"409b2336-33b1-4632-ac6c-864cbcc370c8","year":2015},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:210396d8fa69254bfe5dd78c9b4e0b158a6b2acfda742ef70f0aaf495c6a78ab","observation_id":"fd4b2506-0a61-42f4-aa5a-f64a21745a56","resolution":{"observed_at":"2026-05-16T04:30:37.649394Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"61f75253-5bda-4f6f-aa44-50d126aae16c","year":1980},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:a5a660c4293c110cbfaed1282bebd5f685a944afc7a855bb89ab8b6aac512b8c","observation_id":"4c15ba4b-d167-4603-b233-3e8697d7e006","resolution":{"observed_at":"2026-05-16T04:30:37.664282Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1162/tacl.a.58","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":null,"venue":"Transactions of the Association for Computational Linguistics","work_id":"5a5ba1fc-044f-4169-a7d1-163a58a5cb4d","year":2025},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:1b4a6d5c3996bf05d20f9dd21bd3fca62ff5cdb92bb12859e79281b58ccc8e65","observation_id":"c201a1e6-efd2-4a2b-aa98-f8ad1342f647","resolution":{"observed_at":"2026-05-10T19:40:45.663277Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1073/pnas.2415697122","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"URL https: //www.pnas.org/doi/abs/10.1073/pnas.2415 697122","venue":"Proceedings of the National Academy of Sciences","work_id":"29c77e68-689d-4b61-b8b2-1cba6ddf24b5","year":2025},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:de8bb5f27fc5cca292afa1036d08c76343ebca2967bb98a186e9201fcb4a5e07","observation_id":"8d11da68-f84b-42d3-9592-8b7e910a030b","resolution":{"observed_at":"2026-05-10T19:40:45.647583Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2025.emnlp-main.138","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T14:48:32.225027Z","title":"From generation to judgment: Opportunities and challenges of LLM-as-a-judge","venue":"Proceedings of the 2025 Conference on Empirical Methods in Natural Language Processing","work_id":"d1491133-c760-4410-a6df-a93780897bef","year":2025},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:2c469a6ec48f5e792b7b4a05b1b0446401d6ae236f249b8cd1b215807bfae415","observation_id":"824cc18f-9be0-4592-a031-9de30e033a9d","resolution":{"observed_at":"2026-05-10T19:40:45.670863Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.05579","last_updated":"2024-12-10T05:49:12Z","snapshot_observed_at":"2026-07-31T01:42:39.468673Z","submitted_at":"2024-12-07T08:07:24Z","title":"LLMs-as-Judges: A Comprehensive Survey on LLM-based Evaluation Methods","version":2},"cited_work":{"arxiv_id":"2412.05579","doi":"10.48550/arxiv.2412.05579","metadata_source":"pith","pith_arxiv_id":"2412.05579","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"LLMs-as-Judges: A Comprehensive Survey on LLM-based Evaluation Methods","venue":"cs.CL","work_id":"377b190f-39ed-4db3-9747-433802e5303c","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"cited_paper":"/paper/2412.05579","citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:1c6e8571899b93218b6675eac31da418c0b8909ab5a93f1913e19a7225547b91","observation_id":"2ba0eb17-e0dd-46fd-9695-7ac96219d4d6","resolution":{"observed_at":"2026-05-11T23:08:37.431680Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-03T15:08:45.340409+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-03T15:08:45.340409+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.22316","last_updated":"2026-02-03T11:44:29Z","snapshot_observed_at":"2026-07-06T21:48:45.170162Z","submitted_at":"2025-06-27T15:25:23Z","title":"Evaluating Scoring Bias in LLM-as-a-Judge","version":4},"cited_work":{"arxiv_id":"2506.22316","doi":"10.48550/arxiv.2506.22316","metadata_source":"pith","pith_arxiv_id":"2506.22316","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Evaluating scoring bias in llm-as-a-judge","venue":"cs.CL","work_id":"29bdd127-b356-442d-b893-10e4c6ed8d4c","year":2025},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"cited_paper":"/paper/2506.22316","citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:23e2f48ce88ccd6816c639acce69bcade1bd3e48daca306314c9b39580e491e8","observation_id":"8102e053-c053-4b67-bd3d-d6dc4f24cfd6","resolution":{"observed_at":"2026-05-22T02:03:47.800935Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2506.09443","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T14:17:10.144588Z","title":"LLMs Cannot Reliably Judge (Yet?): A Comprehensive Assessment on the Robustness of LLM-as-a-Judge","venue":null,"work_id":"bd49bdd6-b2b5-4c7e-ad0d-cc07261f2b50","year":2025},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:798b4593394771e5248e4defa7521f6dd60405172ad65513a7ef674644c1cb80","observation_id":"37ec9b1c-1c4c-444f-9ef0-5ce82bf4230c","resolution":{"observed_at":"2026-05-10T22:40:51.670602Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"1146.353318","doi":"10.1145/3531146.3533182","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Designing for Responsible Trust in AI Systems: A Communication Perspective","venue":"2022 ACM Conference on Fairness, Accountability, and Transparency","work_id":"52ae2e2e-bc6d-4537-b33f-4080c7729078","year":2022},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:ede51c1feb67847fcf33a9074c81af25e279228f5bfaf1b8e7674a4db0f62c48","observation_id":"83b4cb02-b702-4301-965f-acc244a000e3","resolution":{"observed_at":"2026-05-10T19:40:45.654523Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.16634","last_updated":"2023-05-23T22:12:16Z","snapshot_observed_at":"2026-08-02T04:02:36.848064Z","submitted_at":"2023-03-29T12:46:54Z","title":"G-Eval: NLG Evaluation using GPT-4 with Better Human Alignment","version":3},"cited_work":{"arxiv_id":"2303.16634","doi":"10.48550/arxiv.2303.16634","metadata_source":"pith","pith_arxiv_id":"2303.16634","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"G-Eval: NLG Evaluation using GPT-4 with Better Human Alignment","venue":"cs.CL","work_id":"6bc7f654-ec90-4d95-b023-416717a28f45","year":2023},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"cited_paper":"/paper/2303.16634","citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:c7fc99cf51981399d3aaf70ad7ed14355d948d9aaf6757100cb389ddcb91e8c4","observation_id":"840b0dea-dd41-4f33-b352-41ec181f9447","resolution":{"observed_at":"2026-05-12T22:55:50.972702Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-05-20T14:22:05.615976+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T14:22:05.615976+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.3390/journalmedia5020046","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":null,"venue":"Journalism and Media","work_id":"80620f6d-3da1-4af6-b13f-76833066d1af","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:89b137bba0a2692dac74eaa73602dd519ca516e4ab10db038ddaad2937b01ccc","observation_id":"fd366746-7250-4011-809d-00549e858da4","resolution":{"observed_at":"2026-05-10T19:40:45.650505Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2509.26072","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"InProceedings of the AAAI Conference on Artificial Intelligence, volume 38, pages 4171–4179","venue":null,"work_id":"94d0aeb4-f032-4240-ba7e-d62a57c3b219","year":2025},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:10c7d9c9996dccb34a4a09d8cc9a04ab7105d2ba3a11835e1f65d327364215ea","observation_id":"ab3ffc62-1552-4f01-8234-af9be4012561","resolution":{"observed_at":"2026-05-10T22:40:51.731097Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":"2303.08774","doi":"10.1002/tea.20265","metadata_source":"pith","pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"GPT-4 Technical Report","venue":"cs.CL","work_id":"b928e041-6991-4c08-8c81-0359e4097c7b","year":2023},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:f485b98b6ae534e51156b3a15834f46eb1cd6168b9c3992ef53dda1e9bb415f3","observation_id":"1d431efc-21fa-468b-bc09-d3da4aae071d","resolution":{"observed_at":"2026-05-10T22:40:51.706199Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T05:23:20.885972Z","title":null,"venue":null,"work_id":"f51bef7e-e7ae-4d3d-be8b-be8d15c4cfbc","year":2022},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:00de0bea84bd6a721a25a17d86934b653804ab343dc19f62e594cb589e69c3a9","observation_id":"5ddbabef-7c17-4039-bb25-c78e7c0ae631","resolution":{"observed_at":"2026-05-16T04:30:37.656709Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"ca62865a-e34e-450a-ad3d-0c5bcbeaa7cd","year":2019},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:e55a582a48efcfdc00b06941ad8332b499b282573756afd6e38e9c07edbb7bc8","observation_id":"0c9b5d3e-e8c1-4c3b-8c83-c515b4f10e6f","resolution":{"observed_at":"2026-05-16T04:30:37.652761Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"0bf64d52-99e8-4739-93df-6741d4f39024","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:9787f2bc5943d3001aa6dd559a811dbd6363b725eb8b0264ecdadbf6ca2f3f49","observation_id":"f3936e39-9b4d-481a-b4b7-67d57aa04e70","resolution":{"observed_at":"2026-05-16T04:30:37.674263Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T04:03:19.144903Z","title":null,"venue":null,"work_id":"c6c5a4c0-c246-4a81-aaf0-914b33694904","year":2006},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:66702c7641c7a408e8605ef25616021cdd7a79d52e49ca37c5cb09aa16f8ddce","observation_id":"aa976e1d-07a5-41c1-be22-8e31410a2c53","resolution":{"observed_at":"2026-05-16T04:30:37.645923Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Rowley, Frances C","venue":null,"work_id":"c371ecad-6823-4378-ba68-ff3e39d1b99f","year":2015},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:2c524a89c500136e1fcd76b2a42f9eab2a760c49c2348f3e62dd34f0cc2d3683","observation_id":"2ef216c8-9d43-425e-bef8-14fa8868509c","resolution":{"observed_at":"2026-05-16T04:30:37.667522Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"3013.359399","doi":"10.1145/3593013.3593990","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Ali, Angèle Christin, Andrew Smart, and Riitta Katila","venue":null,"work_id":"aea877c6-5781-432f-a83b-11deee26f5f9","year":2023},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:bf2520cea7d2ddda263b3e3c943dc8b32e8094279ece5bbec651bde7147bee2d","observation_id":"d02b100b-3dd7-4709-b705-b94181f5ab37","resolution":{"observed_at":"2026-05-10T19:40:45.667455Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.12509","last_updated":"2025-02-18T15:24:25Z","snapshot_observed_at":"2026-08-01T23:28:04.240794Z","submitted_at":"2024-12-17T03:37:31Z","title":"Can You Trust LLM Judgments? Reliability of LLM-as-a-Judge","version":2},"cited_work":{"arxiv_id":"2412.12509","doi":"10.48550/arxiv.2412.12509","metadata_source":"arxiv_reference","pith_arxiv_id":"2412.12509","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Can you trust llm judgments? reliability of llm-as- a-judge","venue":"arXiv (Cornell University)","work_id":"fbc020bc-197e-43ce-92fd-3dca29743d27","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"cited_paper":"/paper/2412.12509","citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:1e61cfc1a12f92588f586534dd7a6d4283455e9c658eb8765693d90e4d512f7d","observation_id":"b2143566-bc8f-46b0-babc-aadb24da9869","resolution":{"observed_at":"2026-05-10T22:40:51.745869Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2025.emnlp-main.569","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Analyzing Uncertainty of LLM -as-a-Judge: Interval Evaluations with Conformal Prediction","venue":null,"work_id":"67d1009a-74c1-481b-b8b7-7bc2d74c6a6a","year":2025},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:a38eab0a52855981d50238f86c6df7f9acaa3f6d7add9e2980d399bc6e57c698","observation_id":"2066a129-7a5d-489e-a368-b95074777215","resolution":{"observed_at":"2026-05-10T19:40:45.657596Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2406.07791","doi":"10.48550/arxiv.2406.07791","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Judging the judges: A systematic study of position bias in llm-as-a-judge","venue":"arXiv (Cornell University)","work_id":"11308458-d080-4667-b088-70cc9c815627","year":2025},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:5c655ba1b0e177319abc9ec9802ae1e6ddea806396ef63c15e6ce06cc3bef836","observation_id":"d573bc3d-09e6-463b-bc25-74df479dc81f","resolution":{"observed_at":"2026-05-10T22:40:51.699055Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.06709","last_updated":"2025-08-08T21:22:12Z","snapshot_observed_at":"2026-08-05T22:40:39.830161Z","submitted_at":"2025-08-08T21:22:12Z","title":"Play Favorites: A Statistical Method to Measure Self-Bias in LLM-as-a-Judge","version":1},"cited_work":{"arxiv_id":"2508.06709","doi":"10.48550/arxiv.2508.06709","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.06709","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"URLhttps://arxiv.org/abs/2508.06709.2508.06709","venue":"ArXiv.org","work_id":"03aa7daa-6b25-43d7-9f74-a0a4973e37ef","year":2025},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"cited_paper":"/paper/2508.06709","citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:72cce09d52dc81b34c03abcba39c24c74c0705b178de279c3ac944c3dc36443c","observation_id":"4d2d3d9e-46ad-4bc4-8ec8-5f727af6279c","resolution":{"observed_at":"2026-05-10T22:40:51.666192Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-05-23T16:27:30.742825+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T16:27:30.742825+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1162/tacl_a_00685","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-02T07:06:43.464196Z","title":"doi: 10.1162/tacl_a_00685","venue":"Transactions of the Association for Computational Linguistics","work_id":"9bef856a-be77-4c6b-bb0c-3127550c747c","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:5f365ce15e8c446507065357e6a8dd91b32e9b70a9a2201a123dc79bbde10175","observation_id":"9be947f1-4a2d-4551-b3cc-b779460aa926","resolution":{"observed_at":"2026-05-10T19:40:45.674149Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-05-23T02:23:23.651111+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T02:23:23.651111+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.09946","last_updated":"2025-04-18T02:05:38Z","snapshot_observed_at":"2026-08-01T15:42:51.877844Z","submitted_at":"2025-04-14T07:14:27Z","title":"Assessing Judging Bias in Large Reasoning Models: An Empirical Study","version":2},"cited_work":{"arxiv_id":"2504.09946","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.09946","snapshot_observed_at":"2026-07-02T13:46:58.798974Z","title":"Assessing judging bias in large reasoning models: An empirical study.arXiv preprint arXiv:2504.09946, 2025a","venue":null,"work_id":"c9624367-78fa-4f5b-84a1-604a4ae46c17","year":2025},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"cited_paper":"/paper/2504.09946","citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:8b3570cb299e588da5d74c92b8fad2cd845687796ca4d855d313acef01830c72","observation_id":"c37bc6d8-22b6-47f0-9ab2-6df4471574f2","resolution":{"observed_at":"2026-05-10T22:40:51.715805Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2509.21117","doi":"10.48550/arxiv.2509.21117","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Trust- judge: Inconsistencies of LLM-as-a-judge and how to alleviate them.arXiv preprint arXiv:2509.21117,","venue":"ArXiv.org","work_id":"10d40448-0d0f-4f32-bc22-b2d999802f0d","year":2025},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:9f3400c0dedad76ee40cac60b2d23acd0b2b82ad9d749e417b2c0904a6d349d1","observation_id":"f3917525-caf0-4018-abc8-d69acf93800c","resolution":{"observed_at":"2026-05-10T22:40:51.710379Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21819","last_updated":"2025-06-21T08:39:06Z","snapshot_observed_at":"2026-08-04T18:30:37.623897Z","submitted_at":"2024-10-29T07:42:18Z","title":"Self-Preference Bias in LLM-as-a-Judge","version":2},"cited_work":{"arxiv_id":"2410.21819","doi":"10.48550/arxiv.2410.21819","metadata_source":"pith","pith_arxiv_id":"2410.21819","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Self-Preference Bias in LLM-as-a-Judge","venue":"cs.CL","work_id":"b98a1723-3323-42b3-a741-8b6785806d2d","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"cited_paper":"/paper/2410.21819","citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:5ed8a862bbf1497de93d28dfc76890e8378fc36f46310f786fb70ad2bb856ad6","observation_id":"e231de27-cbaa-40e7-b90a-c2de4c084f23","resolution":{"observed_at":"2026-05-15T14:46:32.590484Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/d19-1002","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Attention is not not explanation","venue":null,"work_id":"24ebdc1e-3edd-483e-8a1a-78c231cbb355","year":2019},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:62d1b783e49f3aae1836ae92aaa08b79ed5595ee1a6586b91c9daf3e65338ce4","observation_id":"16700c6f-f671-4db3-b04e-fe4bb98e3a23","resolution":{"observed_at":"2026-05-10T19:40:45.644803Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Torr, Bernard Ghanem, and Guohao Li","venue":null,"work_id":"de538698-c5f6-4d74-9d04-b3b8cc6d90b1","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:9afd68571da76b12907320c3b65fdb7dcb3fea6dfb296cd0a0df004257d6fa99","observation_id":"ca39f562-e034-4905-adc6-2f78af4da0f4","resolution":{"observed_at":"2026-05-16T04:30:37.642434Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.06581","last_updated":"2022-06-15T16:07:12Z","snapshot_observed_at":"2026-08-06T02:23:14.214916Z","submitted_at":"2022-06-14T03:49:03Z","title":"CHQ-Summ: A Dataset for Consumer Healthcare Question Summarization","version":2},"cited_work":{"arxiv_id":"2206.06581","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2206.06581","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"17fed372-3c76-47a0-897a-5da8d255b51d","year":2022},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"cited_paper":"/paper/2206.06581","citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:9e06246bcf3ce7640a8896b2a015a000c824736c24d192e4dea022078869247f","observation_id":"5ff2e9ce-e592-4a0a-877e-c351d3384938","resolution":{"observed_at":"2026-05-10T22:40:51.694620Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2505.17100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"88bf8b50-b055-4ef9-b302-6200ee72e79c","year":2025},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:991fc40e032eee7544f53761d4f794c256f951ebd448afdc3afa8e2e3c170778","observation_id":"8769697d-4564-4c24-9fbe-20e191c5a429","resolution":{"observed_at":"2026-05-10T22:40:51.759792Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.02736","last_updated":"2024-10-04T03:57:47Z","snapshot_observed_at":"2026-08-01T08:21:19.528254Z","submitted_at":"2024-10-03T17:53:30Z","title":"Justice or Prejudice? Quantifying Biases in LLM-as-a-Judge","version":2},"cited_work":{"arxiv_id":"2410.02736","doi":"10.48550/arxiv.2410.02736","metadata_source":"pith","pith_arxiv_id":"2410.02736","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Justice or Prejudice? Quantifying Biases in LLM-as-a-Judge","venue":"cs.CL","work_id":"226fa5a3-d2b0-46fc-b73d-cfb5ca7b6a4c","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"cited_paper":"/paper/2410.02736","citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:6e96948d4d65fd5209e323fc1f89a83b8f394468a6b90dbda33b3cf50161849d","observation_id":"ece987b0-a795-4e6f-a376-41d4c4f320ad","resolution":{"observed_at":"2026-05-15T20:00:24.757706Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:47.157032+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:47.157032+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1073/pnas.2319112121","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":", date =","venue":"Proceedings of the National Academy of Sciences","work_id":"bb48a64a-29c1-4159-b13c-d48c8cb0a079","year":2024},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:49fae3fda03881f64c34a3ccfaf7ae5d25c22c879b920e25b757bd73db136d20","observation_id":"796a9868-9c95-4778-80af-8e52c33ad9da","resolution":{"observed_at":"2026-05-10T19:40:45.642098Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.05685","last_updated":"2023-12-24T02:01:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-06-09T05:55:52Z","title":"Judging LLM-as-a-Judge with MT-Bench and Chatbot Arena","version":4},"cited_work":{"arxiv_id":"2306.05685","doi":"10.1109/4235.797969","metadata_source":"pith","pith_arxiv_id":"2306.05685","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Judging LLM-as-a-Judge with MT-Bench and Chatbot Arena","venue":"cs.CL","work_id":"d0c30cd7-81e1-4159-a87f-f6adca77ff08","year":2023},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"cited_paper":"/paper/2306.05685","citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:1a87a08caacab37a540b3e85ef67d410b879f36aacff6e57773618821a3a0189","observation_id":"aee7455e-b161-4f5d-b82c-55b366002df0","resolution":{"observed_at":"2026-05-10T22:40:51.686239Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T10:56:11.938856Z","title":"online\" 'onlinestring :=","venue":null,"work_id":"38c07273-5071-4bbb-b2af-067af11becc7","year":null},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:e47ded59cc0a4676c74987e846ecb369751370595d882c9be03ba3f31e0c7171","observation_id":"598d470d-9110-4cbb-8512-7b008e8e7461","resolution":{"observed_at":"2026-05-16T04:30:37.671038Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T10:56:11.963345Z","title":"write newline","venue":null,"work_id":"b6ace3ad-bd44-4a95-86ee-c7b73bd4cda2","year":null},"citing_paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-05-10T19:38:11.595077Z"},"links":{"citing_paper":"/paper/2604.05593"},"observation_digest":"sha256:d90a8596f1f87e83e0c353fbf7c82059f1ad8a99357017d691a59cc1d56d5ea6","observation_id":"289700b1-1d06-46c5-bcaf-1ebc888d6758","resolution":{"observed_at":"2026-05-16T04:30:37.677639Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2604.05593","last_updated":"2026-04-07T08:43:30Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-07-06T22:54:16.913706Z","submitted_at":"2026-04-07T08:43:30Z","title":"Label Effects: Shared Heuristic Reliance in Trust Assessment by Humans and LLM-as-a-Judge"},"reference_resolution":{"displayed":55,"state_counts":{"malformed_identifier":0,"metadata_mismatch":4,"parse_uncertain":0,"unresolved":13,"verified_exact":32,"verified_fuzzy":6},"total_outbound_references":55},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 6 August 2026, this Paper Citation Record lists 55 of 55 outbound references and 1 inbound Pith citation observation for arXiv:2604.05593."}