{"as_of":"2026-08-06T13:18:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:c5a00e7b35c81eb2a7a5b9316255d2a0610a6164ad1818a7c01555d817c2b3af","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":15,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":15,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-06T06:34:29.942622+00:00","state":"measured"},{"denominator":15,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":15,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T06:17:17.719091Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge","version":2},"cited_work":{"arxiv_id":"2502.00561","doi":"10.48550/arxiv.2502.00561","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.00561","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"ACM Transactions on Computer-Human Interaction, 27(5)","venue":"ArXiv.org","work_id":"1079b5c0-2a24-4af7-982f-03e9cd0441fe","year":2025},"citing_paper":{"arxiv_id":"2509.14528","last_updated":"2026-05-03T23:02:14Z","snapshot_observed_at":"2026-07-06T22:30:11.897367Z","submitted_at":"2025-09-18T01:51:29Z","title":"Why Johnny Can't Use Agents: Industry Aspirations vs. User Realities with AI Agents","version":2},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-05-18T16:55:47.922639Z"},"links":{"cited_paper":"/paper/2502.00561","citing_paper":"/paper/2509.14528"},"observation_digest":"sha256:29a548bcf9c83cfbad877de5d4c83551445e154c0300f622c377cd4c118fef28","observation_id":"c69343e9-7ea8-4ff1-ba37-df93a101e6d9","resolution":{"observed_at":"2026-05-18T16:56:38.062223Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge","version":2},"cited_work":{"arxiv_id":"2502.00561","doi":"10.48550/arxiv.2502.00561","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.00561","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"ACM Transactions on Computer-Human Interaction, 27(5)","venue":"ArXiv.org","work_id":"1079b5c0-2a24-4af7-982f-03e9cd0441fe","year":2025},"citing_paper":{"arxiv_id":"2602.00065","last_updated":"2026-04-28T08:46:23Z","snapshot_observed_at":"2026-08-06T10:24:27.916413Z","submitted_at":"2026-01-20T12:55:10Z","title":"Responsible Evaluation of AI for Mental Health","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-16T12:50:33.070425Z"},"links":{"cited_paper":"/paper/2502.00561","citing_paper":"/paper/2602.00065"},"observation_digest":"sha256:a8430f1a4cc027aef6b39888fe55066259bcd9b899fb4e56ae5e6a15e416ff31","observation_id":"3a33c585-df94-4a14-a6c2-f96c21ae8e7d","resolution":{"observed_at":"2026-05-16T12:50:54.931535Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge","version":2},"cited_work":{"arxiv_id":"2502.00561","doi":"10.48550/arxiv.2502.00561","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.00561","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"ACM Transactions on Computer-Human Interaction, 27(5)","venue":"ArXiv.org","work_id":"1079b5c0-2a24-4af7-982f-03e9cd0441fe","year":2025},"citing_paper":{"arxiv_id":"2603.06811","last_updated":"2026-05-08T14:52:11Z","snapshot_observed_at":"2026-08-05T19:05:35.953649Z","submitted_at":"2026-03-06T19:17:50Z","title":"Making AI Evaluation Deployment Relevant Through Context Specification","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-15T14:46:09.944168Z"},"links":{"cited_paper":"/paper/2502.00561","citing_paper":"/paper/2603.06811"},"observation_digest":"sha256:8248399d0fe77e9565815af43fe18818d3cea40df26e4e6762be1b7322a722c2","observation_id":"729587ef-4445-41e6-a3ae-e208359ae64a","resolution":{"observed_at":"2026-05-15T14:50:05.159501Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.00561","snapshot_observed_at":"2026-07-13T21:26:43.149095Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.20510","last_updated":"2026-06-23T04:29:10Z","snapshot_observed_at":"2026-08-04T06:04:31.773147Z","submitted_at":"2026-03-20T21:24:28Z","title":"Grounded Chess Reasoning in Language Models via Master Distillation","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-07-13T21:26:43.149095Z"},"links":{"cited_paper":"/paper/2502.00561","citing_paper":"/paper/2603.20510"},"observation_digest":"sha256:499241dd7e3f50f29c4fd64ec0c8a7834cebbcf2a7ac537bf051efe04c07bcbd","observation_id":"70749d48-bbd7-4928-bf76-ad85e797bbcd","resolution":{"observed_at":"2026-07-13T21:26:43.149095Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge","version":2},"cited_work":{"arxiv_id":"2502.00561","doi":"10.48550/arxiv.2502.00561","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.00561","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"ACM Transactions on Computer-Human Interaction, 27(5)","venue":"ArXiv.org","work_id":"1079b5c0-2a24-4af7-982f-03e9cd0441fe","year":2025},"citing_paper":{"arxiv_id":"2604.03238","last_updated":"2026-05-29T23:39:37Z","snapshot_observed_at":"2026-08-03T05:55:14.155931Z","submitted_at":"2026-01-31T21:51:17Z","title":"RLHF May Not Reflect Genuine Preferences","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-16T08:39:08.880486Z"},"links":{"cited_paper":"/paper/2502.00561","citing_paper":"/paper/2604.03238"},"observation_digest":"sha256:cd8d2d9897be0de78a6be403e97865b6c55cef878c92c86e6c969b9f530a9a0f","observation_id":"0191584d-4385-4212-ae50-19f6e3560e0f","resolution":{"observed_at":"2026-05-16T08:40:46.318423Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge","version":2},"cited_work":{"arxiv_id":"2502.00561","doi":"10.48550/arxiv.2502.00561","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.00561","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"ACM Transactions on Computer-Human Interaction, 27(5)","venue":"ArXiv.org","work_id":"1079b5c0-2a24-4af7-982f-03e9cd0441fe","year":2025},"citing_paper":{"arxiv_id":"2604.07591","last_updated":"2026-04-08T20:49:03Z","snapshot_observed_at":"2026-07-06T22:55:46.307881Z","submitted_at":"2026-04-08T20:49:03Z","title":"From Ground Truth to Measurement: A Statistical Framework for Human Labeling","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-10T17:09:43.892161Z"},"links":{"cited_paper":"/paper/2502.00561","citing_paper":"/paper/2604.07591"},"observation_digest":"sha256:95ed67d741e8a5acb81c5bb8c772c6f3dc77c9ab3fe9f8f767b55e63a7d3613d","observation_id":"e1116046-03f9-40e2-9144-9a0d98caa63a","resolution":{"observed_at":"2026-05-11T07:30:58.092181Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge","version":2},"cited_work":{"arxiv_id":"2502.00561","doi":"10.48550/arxiv.2502.00561","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.00561","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"ACM Transactions on Computer-Human Interaction, 27(5)","venue":"ArXiv.org","work_id":"1079b5c0-2a24-4af7-982f-03e9cd0441fe","year":2025},"citing_paper":{"arxiv_id":"2604.14266","last_updated":"2026-04-15T17:21:56Z","snapshot_observed_at":"2026-07-30T12:52:59.090472Z","submitted_at":"2026-04-15T17:21:56Z","title":"\"I Just Don't Want My Work Being Fed Into The AI Blender\": Queer Artists on Refusing and Resisting Generative AI","version":1},"reference_index":138,"source":"pdf_text","source_observed_at":"2026-05-10T12:21:20.533180Z"},"links":{"cited_paper":"/paper/2502.00561","citing_paper":"/paper/2604.14266"},"observation_digest":"sha256:36b2e23a6c8f5037deb3a1f753401c6fce49336d94dd1189189e86fa5fd6a8ae","observation_id":"88589ab5-5cf6-4cf8-a6b4-3a499773718a","resolution":{"observed_at":"2026-05-10T12:25:22.668278Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge","version":2},"cited_work":{"arxiv_id":"2502.00561","doi":"10.48550/arxiv.2502.00561","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.00561","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"ACM Transactions on Computer-Human Interaction, 27(5)","venue":"ArXiv.org","work_id":"1079b5c0-2a24-4af7-982f-03e9cd0441fe","year":2025},"citing_paper":{"arxiv_id":"2605.07986","last_updated":"2026-05-08T16:44:55Z","snapshot_observed_at":"2026-07-06T23:20:15.696237Z","submitted_at":"2026-05-08T16:44:55Z","title":"Towards Apples to Apples for AI Evaluations: From Real-World Use Cases to Evaluation Scenarios","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-05-11T02:54:52.984994Z"},"links":{"cited_paper":"/paper/2502.00561","citing_paper":"/paper/2605.07986"},"observation_digest":"sha256:41b783249aa7706e3fa1b1bf2be0df9aceecc76272b41bc5ed518fb464ee4070","observation_id":"fd24ef72-532d-402a-bfae-f5c653595ac9","resolution":{"observed_at":"2026-05-11T02:55:53.066086Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge","version":2},"cited_work":{"arxiv_id":"2502.00561","doi":"10.48550/arxiv.2502.00561","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.00561","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"ACM Transactions on Computer-Human Interaction, 27(5)","venue":"ArXiv.org","work_id":"1079b5c0-2a24-4af7-982f-03e9cd0441fe","year":2025},"citing_paper":{"arxiv_id":"2606.02293","last_updated":"2026-06-01T14:16:13Z","snapshot_observed_at":"2026-07-06T23:42:44.661608Z","submitted_at":"2026-06-01T14:16:13Z","title":"AI as a Tool for Simulation-Based Experiments in Literary Studies","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-06-28T14:54:14.909995Z"},"links":{"cited_paper":"/paper/2502.00561","citing_paper":"/paper/2606.02293"},"observation_digest":"sha256:36dfd2203904ccb9a002d9d97d2f27bbbef0b9f3b616590a79327dce54ba377f","observation_id":"a1357ba7-1d60-473f-bc96-e9f837330683","resolution":{"observed_at":"2026-06-28T15:02:18.979706Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge","version":2},"cited_work":{"arxiv_id":"2502.00561","doi":"10.48550/arxiv.2502.00561","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.00561","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"ACM Transactions on Computer-Human Interaction, 27(5)","venue":"ArXiv.org","work_id":"1079b5c0-2a24-4af7-982f-03e9cd0441fe","year":2025},"citing_paper":{"arxiv_id":"2606.11018","last_updated":"2026-07-29T19:31:42Z","snapshot_observed_at":"2026-08-06T10:55:02.059560Z","submitted_at":"2026-06-09T15:55:55Z","title":"Measuring Human Value Expression in Social Media Texts: Calibrated LLM Annotation and Encoder Transfer","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-06-27T13:13:19.690367Z"},"links":{"cited_paper":"/paper/2502.00561","citing_paper":"/paper/2606.11018"},"observation_digest":"sha256:d85127b1cbdaf17b42a1f172faf92b1a39b6c9aa9d553eb339e2eab5061f4421","observation_id":"9562a4fc-d8f1-4cda-ba93-c52ece092f06","resolution":{"observed_at":"2026-07-03T05:27:40.352337Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge","version":2},"cited_work":{"arxiv_id":"2502.00561","doi":"10.48550/arxiv.2502.00561","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.00561","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"ACM Transactions on Computer-Human Interaction, 27(5)","venue":"ArXiv.org","work_id":"1079b5c0-2a24-4af7-982f-03e9cd0441fe","year":2025},"citing_paper":{"arxiv_id":"2606.17799","last_updated":"2026-07-18T22:14:13Z","snapshot_observed_at":"2026-08-02T11:05:59.500924Z","submitted_at":"2026-06-16T11:21:01Z","title":"Position: Coding Benchmarks Are Misaligned with Agentic Software Engineering","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-06-26T23:48:58.497927Z"},"links":{"cited_paper":"/paper/2502.00561","citing_paper":"/paper/2606.17799"},"observation_digest":"sha256:32545463e499040b7c9636b33f669cdf700bb36b6b63774a5aa0ecb4ec31ae0f","observation_id":"7948401f-e3a0-43b7-9be0-a9cd98d0453b","resolution":{"observed_at":"2026-07-03T22:08:58.949519Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.00561","snapshot_observed_at":"2026-08-02T11:06:03.560871Z","title":"Feder Cooper, Angelina Wang, Chad Atalla, Solon Barocas, Su Lin Blodgett, Alexandra Chouldechova, Emily Corvi, P","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.17799","last_updated":"2026-07-18T22:14:13Z","snapshot_observed_at":"2026-08-02T11:05:59.500924Z","submitted_at":"2026-06-16T11:21:01Z","title":"Position: Coding Benchmarks Are Misaligned with Agentic Software Engineering","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-02T11:06:03.560871Z"},"links":{"cited_paper":"/paper/2502.00561","citing_paper":"/paper/2606.17799"},"observation_digest":"sha256:b51f4750ef1ccfc61c287181eab753ebc9d3ecf5fd23c7625680208f7ebce1f9","observation_id":"2f29de4b-0b44-496a-8c61-ed0646ca3c24","resolution":{"observed_at":"2026-08-02T11:06:03.560871Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.00561","snapshot_observed_at":"2026-08-02T07:44:03.248723Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.09306","last_updated":"2026-07-30T11:40:45Z","snapshot_observed_at":"2026-08-02T23:59:13.473449Z","submitted_at":"2026-07-10T11:39:40Z","title":"Exposure is not manifestation: measurement target and output resolution jointly determine which behavioural-faithfulness evaluator wins","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-02T07:44:03.248723Z"},"links":{"cited_paper":"/paper/2502.00561","citing_paper":"/paper/2607.09306"},"observation_digest":"sha256:84147dc16c3eae20785899f691b938738f6c30ab23057d944c0bda2f014ecffc","observation_id":"13e8e9b3-298f-419d-a664-93ccfc27f6f8","resolution":{"observed_at":"2026-08-02T07:44:03.248723Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.00561","snapshot_observed_at":"2026-07-31T01:29:53.636089Z","title":"Feder Cooper, Angelina Wang, Chad Atalla, Solon Barocas, Su Lin Blodgett, Alexandra Chouldechova, Emily Corvi, P","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.25097","last_updated":"2026-07-27T21:54:19Z","snapshot_observed_at":"2026-08-01T00:05:08.188877Z","submitted_at":"2026-07-27T21:54:19Z","title":"On the Convergent Validity of Offline Evaluation Designs for Recommender Systems","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-07-31T01:29:53.636089Z"},"links":{"cited_paper":"/paper/2502.00561","citing_paper":"/paper/2607.25097"},"observation_digest":"sha256:ee8caf6dbfb1438dfeca7251bd4a5ecb78bbb46a4af4a8273e44dfcce1031d2a","observation_id":"077493da-6a1e-4c2a-a038-729646409e34","resolution":{"observed_at":"2026-07-31T01:29:53.636089Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.00561","snapshot_observed_at":"2026-08-04T06:17:17.719091Z","title":"Health and Quality of Life Outcomes, 5(63)","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2608.02491","last_updated":"2026-08-05T04:54:08Z","snapshot_observed_at":"2026-08-06T12:30:41.336478Z","submitted_at":"2026-08-03T16:55:50Z","title":"Long-term Measurements: Towards a Longitudinal Understanding of Human-AI Interactions","version":1},"reference_index":2007,"source":"pdf_text","source_observed_at":"2026-08-04T06:17:17.719091Z"},"links":{"cited_paper":"/paper/2502.00561","citing_paper":"/paper/2608.02491"},"observation_digest":"sha256:8bbf9ac05faf8162696340b8932ac3107392d121c30def3a345dc3f83fd8d213","observation_id":"eab46b33-0a67-4492-931c-1842b1c368b1","resolution":{"observed_at":"2026-08-04T06:17:17.719091Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2502.00561/citation-record","integrity":"/paper/2502.00561/integrity","json":"/paper/2502.00561/citation-record.json","paper":"/paper/2502.00561"},"outbound":[],"paper":{"arxiv_id":"2502.00561","last_updated":"2025-06-06T22:15:14Z","latest_version":2,"primary_category":"cs.CY","snapshot_observed_at":"2026-07-06T20:29:36.154633Z","submitted_at":"2025-02-01T21:09:51Z","title":"Position: Evaluating Generative AI Systems Is a Social Science Measurement Challenge"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"thesis":"As of 6 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 15 inbound Pith citation observations for arXiv:2502.00561."}