{"as_of":"2026-08-10T17:06:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:fe5638eb59ff8f6c23fbcdd4670ae9850efbace5068796ae5967131347d94ebd","coverage":[{"denominator":39,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":39,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T10:44:41.145054Z","state":"measured"},{"denominator":39,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":39,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2502.02896/citation-record","integrity":"/paper/2502.02896/integrity","json":"/paper/2502.02896/citation-record.json","paper":"/paper/2502.02896"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:40.967981Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:40.967981Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:740bd4e52e718af19f527762e8732a59c777948d5b548ae373145577a9c1709b","observation_id":"c02d07c5-c1bc-4984-8422-a20e7cc26d16","resolution":{"observed_at":"2026-08-09T10:44:40.967981Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:40.973429Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:40.973429Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:6d4890b6b3761fdc1f7170dd5e024b37484baaf2bfdbcef3fce02ca641dafaf2","observation_id":"e7081593-2b30-49f0-9104-a2e77789caf2","resolution":{"observed_at":"2026-08-09T10:44:40.973429Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:40.978142Z","title":"Koutsiana, J","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:40.978142Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:3124d206024848cfa84a134fe3641f8fecaeda56c7981546982b86f8e9ccc9d6","observation_id":"1f20b524-c4c3-435e-965e-b1520a7ec2a4","resolution":{"observed_at":"2026-08-09T10:44:40.978142Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.928560Z","title":null,"venue":null,"work_id":"e2ff0cdc-7730-4a5b-8c70-e072a76df9f9","year":2023},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:40.982854Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:804cdfb07340a1f925edee7e61bda77bd9828b53148c6047529286da30a3c6f3","observation_id":"4b311f49-816b-4cd7-9484-79b28e3931e6","resolution":{"observed_at":"2026-08-09T10:44:41.933813Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.05232","last_updated":"2024-11-19T12:42:45Z","snapshot_observed_at":"2026-07-06T16:45:07.733095Z","submitted_at":"2023-11-09T09:25:37Z","title":"A Survey on Hallucination in Large Language Models: Principles, Taxonomy, Challenges, and Open Questions","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.05232","snapshot_observed_at":"2026-08-09T10:44:40.987851Z","title":"Huang, W","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:40.987851Z"},"links":{"cited_paper":"/paper/2311.05232","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:365f044a16f3b87817ec45fcccc88ad834a4e54bec32cb108427c057cbc6e81b","observation_id":"c6bc7544-0ae3-4e2e-b219-9c4cfd4d64f6","resolution":{"observed_at":"2026-08-09T10:44:40.987851Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.913674Z","title":"Mickus, E","venue":null,"work_id":"8b786fe3-bf88-4a74-8fab-b07fc12813f2","year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:40.993396Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:6701741532c0340999d62f995bf270b0c37bbeb552a9218b75b76c97e481ff6b","observation_id":"4a1a8053-efa9-4829-9497-0f6a199a8725","resolution":{"observed_at":"2026-08-09T10:44:41.918875Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17468","last_updated":"2024-07-24T17:59:05Z","snapshot_observed_at":"2026-07-06T18:51:28.182977Z","submitted_at":"2024-07-24T17:59:05Z","title":"WildHallucinations: Evaluating Long-form Factuality in LLMs with Real-World Entity Queries","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17468","snapshot_observed_at":"2026-08-09T10:44:40.999318Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:40.999318Z"},"links":{"cited_paper":"/paper/2407.17468","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:691de20452988919f9b1afa5fa2c889e730d2b06a7952514ed06bf8bd41524d5","observation_id":"65ef6303-f4f0-4fe3-9d56-80217d083c21","resolution":{"observed_at":"2026-08-09T10:44:40.999318Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.02243","last_updated":"2025-02-17T11:09:58Z","snapshot_observed_at":"2026-08-10T16:06:17.326698Z","submitted_at":"2024-02-03T19:19:34Z","title":"Language Writ Large: LLMs, ChatGPT, Grounding, Meaning and Understanding","version":2},"cited_work":{"arxiv_id":"2402.02243","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.02243","snapshot_observed_at":"2026-08-09T10:44:41.466841Z","title":"Language Writ Large: LLMs, ChatGPT, Grounding, Meaning and Understanding","venue":"cs.CL","work_id":"da5e65a9-63b0-462c-aebb-524bb3fc765e","year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.005227Z"},"links":{"cited_paper":"/paper/2402.02243","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:62eed887475a18f4e4eed2285a842714f227fb26cbdb000f33688158ce2140b5","observation_id":"210837a3-5568-4ad4-99f8-a8b503349d7d","resolution":{"observed_at":"2026-08-09T10:44:41.473023Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.10997","last_updated":"2024-03-27T09:16:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-18T07:47:33Z","title":"Retrieval-Augmented Generation for Large Language Models: A Survey","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.10997","snapshot_observed_at":"2026-08-09T10:44:41.010402Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.010402Z"},"links":{"cited_paper":"/paper/2312.10997","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:5cdf178044285dc1c5b031af56135e4ac60af4c41d82ef2b1c30a92caab0de9b","observation_id":"74ac46a2-0ad4-433a-8a05-19e091f9047c","resolution":{"observed_at":"2026-08-09T10:44:41.010402Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.898617Z","title":null,"venue":null,"work_id":"7dd110b8-2a60-403d-8353-9ee95a3b44b6","year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.015461Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:42f23e90b37b09a84212cd6b79a3500d1013550caa7886f21a2e17deaa478e5b","observation_id":"262e9596-d95c-439e-95e5-7bca45c13c30","resolution":{"observed_at":"2026-08-09T10:44:41.903803Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1909.01066","last_updated":"2019-09-04T09:33:20Z","snapshot_observed_at":"2026-07-06T08:18:37.267833Z","submitted_at":"2019-09-03T11:11:08Z","title":"Language Models as Knowledge Bases?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1909.01066","snapshot_observed_at":"2026-08-09T10:44:41.020316Z","title":"Petroni, T","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.020316Z"},"links":{"cited_paper":"/paper/1909.01066","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:8132bb3978ea4773db7602eed1ea5a974b6eb0fb233bf85fc75009a369ecd703","observation_id":"0bc3fd8d-c396-4d4e-b1b1-7877b3797bff","resolution":{"observed_at":"2026-08-09T10:44:41.020316Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.883145Z","title":null,"venue":null,"work_id":"19a8fef0-5098-4398-bd08-6e15ce7380b9","year":2022},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.025454Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:0409b737ad14b2d0dd34cf612ee179398f9f02408475ebb3e06346b1ef280703","observation_id":"5e48f049-bb9f-4863-977e-e89609d0b630","resolution":{"observed_at":"2026-08-09T10:44:41.888553Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14251","last_updated":"2023-10-11T05:27:50Z","snapshot_observed_at":"2026-08-01T16:29:58.693560Z","submitted_at":"2023-05-23T17:06:00Z","title":"FActScore: Fine-grained Atomic Evaluation of Factual Precision in Long Form Text Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14251","snapshot_observed_at":"2026-08-09T10:44:41.029692Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.029692Z"},"links":{"cited_paper":"/paper/2305.14251","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:67991a099d7213cba4aa0d9d1b635e78cd2e0762a6df7cb191b7efbebd37c3c1","observation_id":"35564da3-153b-4a32-9e5b-423c364ba9ea","resolution":{"observed_at":"2026-08-09T10:44:41.029692Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.869805Z","title":"Plunkett, T","venue":null,"work_id":"bb732e4c-bfcd-4513-862a-1d481fe7052b","year":2023},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.033847Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:46f1313c5e1cbf608e345c6f9ef391b3a8d8b07e201efbf3870f51d891307ab3","observation_id":"1223efc6-f851-4b4d-826c-94df6c0cd338","resolution":{"observed_at":"2026-08-09T10:44:41.873649Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.856992Z","title":"Plunkett, T","venue":null,"work_id":"9cceb0ba-50eb-40c5-a313-00692bc907f2","year":2013},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.037790Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:5e2033130fc80f40ba7b1c8a9a1982127b1efed92b5c8e73d15491f4fe1d3d4a","observation_id":"5d5f1297-6aff-4c50-9f7b-14cba7969d0e","resolution":{"observed_at":"2026-08-09T10:44:41.861110Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.843925Z","title":null,"venue":null,"work_id":"f5903014-28d2-4837-b897-5a9755b3b244","year":2023},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.041602Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:5acfca7e3d9763d1ab6ac818e3e3d2f8661aaaf050db18bcec25cdffd0cf6e1e","observation_id":"dcf8159e-ec49-4a2d-b95f-adb79cbab499","resolution":{"observed_at":"2026-08-09T10:44:41.847818Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.830662Z","title":"Paulheim, Knowledge graph refinement: A survey of approaches and evaluation methods, Semantic web 8 (2017) 489–508","venue":null,"work_id":"a085aeae-4688-4ddb-88ad-e8592eb88e59","year":2017},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.045373Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:3608a3514df184e1c8fa1a56d4b7fd5406040890d9f90685d504f8b104a61fb0","observation_id":"115790db-21f6-4c9e-9380-a55bcd7c4eda","resolution":{"observed_at":"2026-08-09T10:44:41.835129Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.03749","last_updated":"2024-11-02T01:36:11Z","snapshot_observed_at":"2026-07-06T16:57:55.137830Z","submitted_at":"2023-12-01T01:58:16Z","title":"Conceptual Engineering Using Large Language Models","version":2},"cited_work":{"arxiv_id":"2312.03749","doi":null,"metadata_source":"pith","pith_arxiv_id":"2312.03749","snapshot_observed_at":"2026-08-09T10:44:41.398835Z","title":"Conceptual Engineering Using Large Language Models","venue":"cs.CL","work_id":"6bc26b0e-6ac2-45e3-b9d1-7a97ef221e42","year":2023},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.049154Z"},"links":{"cited_paper":"/paper/2312.03749","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:f8969f6feee63579c5d8e4a6efa0a3efdc8a147b632b18d100731b23fe6602dc","observation_id":"aea1492c-2fcd-421b-8348-d4456da17f11","resolution":{"observed_at":"2026-08-09T10:44:41.407098Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.816805Z","title":"Khatri, C","venue":null,"work_id":"96a259ba-4b4e-4f15-afca-c313172d76c6","year":2010},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.053192Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:094a497d604cb52ba30b0b2955186820a2fff5945deb7da27c8ca079284aa8a4","observation_id":"a16f1638-79ec-408c-a502-9015ee746ee2","resolution":{"observed_at":"2026-08-09T10:44:41.821559Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.801192Z","title":null,"venue":null,"work_id":"605d4405-489e-4062-9e87-37f429fbe7de","year":2016},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.056990Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:0edddc507ecd4b5602091beb7831f0ac5e4dda8284e9b0cc275b287811f4ca1d","observation_id":"7754fa78-ed36-4b53-8b1c-bfef38c369b9","resolution":{"observed_at":"2026-08-09T10:44:41.806745Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.03345","last_updated":"2024-05-06T10:51:09Z","snapshot_observed_at":"2026-07-06T18:10:20.988634Z","submitted_at":"2024-05-06T10:51:09Z","title":"FAIR 2.0: Extending the FAIR Guiding Principles to Address Semantic Interoperability","version":1},"cited_work":{"arxiv_id":"2405.03345","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.03345","snapshot_observed_at":"2026-08-09T10:44:41.375706Z","title":"FAIR 2.0: Extending the FAIR Guiding Principles to Address Semantic Interoperability","venue":"cs.DB","work_id":"769b048b-b5a9-4c39-b2ae-ce9b5da2dd55","year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.061584Z"},"links":{"cited_paper":"/paper/2405.03345","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:b6cf20e76b4a12c9d2f33b347223b3150bf5135c3d92a6a125a06cb32886d9c8","observation_id":"58b07540-9f08-49df-95c3-7957f0dc2b55","resolution":{"observed_at":"2026-08-09T10:44:41.382356Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.784534Z","title":"Elsahar, P","venue":null,"work_id":"0966bd56-4784-476b-8ce3-941550ff1d94","year":2018},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.066069Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:c0bd984103671bfce54bba46b58b258acb907a6f77d1361e1870b8ef826173ae","observation_id":"10b97c26-7a43-4478-a969-70c262d4757b","resolution":{"observed_at":"2026-08-09T10:44:41.790628Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.768294Z","title":"Kojima, S","venue":null,"work_id":"5ec39047-bdd7-459f-9353-c388bd50d898","year":2022},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.070405Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:8fdfe7ad6edbd4f3803fbfea7854019f9e934814b41f6b463e8d9fec2d926e3d","observation_id":"2899adbd-6ce3-4629-b788-84263596e522","resolution":{"observed_at":"2026-08-09T10:44:41.773787Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.17000","last_updated":"2024-04-25T19:44:46Z","snapshot_observed_at":"2026-07-06T18:05:51.085146Z","submitted_at":"2024-04-25T19:44:46Z","title":"Evaluating Class Membership Relations in Knowledge Graphs using Large Language Models","version":1},"cited_work":{"arxiv_id":"2404.17000","doi":null,"metadata_source":"pith","pith_arxiv_id":"2404.17000","snapshot_observed_at":"2026-08-09T10:44:41.350108Z","title":"Evaluating Class Membership Relations in Knowledge Graphs using Large Language Models","venue":"cs.CL","work_id":"c6153b13-0252-40e6-a7fe-a1d4b0effd98","year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.074731Z"},"links":{"cited_paper":"/paper/2404.17000","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:804c5375bb76fe046997f4a84c44be95c759670db73ad2ab434255c04b6a9718","observation_id":"c0a761df-ce65-4946-936d-7d4d60a26015","resolution":{"observed_at":"2026-08-09T10:44:41.358514Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.01937","last_updated":"2023-05-03T07:28:50Z","snapshot_observed_at":"2026-08-02T17:07:32.299595Z","submitted_at":"2023-05-03T07:28:50Z","title":"Can Large Language Models Be an Alternative to Human Evaluations?","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.01937","snapshot_observed_at":"2026-08-09T10:44:41.079145Z","title":"Chiang, H.-y","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.079145Z"},"links":{"cited_paper":"/paper/2305.01937","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:3099a0c8a768576a36e9c4ba00dda924f53fdbc3d04060214f96506ab453a7e4","observation_id":"2044898d-6f81-4aa9-bd60-dcbe72f24209","resolution":{"observed_at":"2026-08-09T10:44:41.079145Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.18403","last_updated":"2025-06-02T11:31:19Z","snapshot_observed_at":"2026-07-06T18:37:21.973331Z","submitted_at":"2024-06-26T14:56:13Z","title":"LLMs instead of Human Judges? A Large Scale Empirical Study across 20 NLP Evaluation Tasks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.18403","snapshot_observed_at":"2026-08-09T10:44:41.083648Z","title":"Bavaresco, R","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.083648Z"},"links":{"cited_paper":"/paper/2406.18403","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:662952adc30de33ec51f14ab2d325680d141c7a523d0684182abe900bcef39e0","observation_id":"2c92bdb5-2795-40cf-8fef-06fc07aecfda","resolution":{"observed_at":"2026-08-09T10:44:41.083648Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12624","last_updated":"2025-08-18T00:07:23Z","snapshot_observed_at":"2026-08-01T14:58:17.402553Z","submitted_at":"2024-06-18T13:49:54Z","title":"Judging the Judges: Evaluating Alignment and Vulnerabilities in LLMs-as-Judges","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12624","snapshot_observed_at":"2026-08-09T10:44:41.088451Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.088451Z"},"links":{"cited_paper":"/paper/2406.12624","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:d27d8786dbecf8481c7ea5f4ca5ce876d6d9ccc46c280dbf1c889cd3a6ca5190","observation_id":"880d571c-3ab7-4bc6-8fb6-19a2c5ab27e3","resolution":{"observed_at":"2026-08-09T10:44:41.088451Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.03732","last_updated":"2024-04-04T18:01:21Z","snapshot_observed_at":"2026-08-05T12:28:35.188792Z","submitted_at":"2024-04-04T18:01:21Z","title":"SHROOM-INDElab at SemEval-2024 Task 6: Zero- and Few-Shot LLM-Based Classification for Hallucination Detection","version":1},"cited_work":{"arxiv_id":"2404.03732","doi":"10.48550/arxiv.2404.03732","metadata_source":"pith","pith_arxiv_id":"2404.03732","snapshot_observed_at":"2026-08-10T05:30:23.456663Z","title":"SHROOM-INDElab at SemEval-2024 Task 6: Zero- and Few-Shot LLM-Based Classification for Hallucination Detection","venue":"cs.CL","work_id":"831ccf2a-051b-4738-a7f6-0e79629135b4","year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.093526Z"},"links":{"cited_paper":"/paper/2404.03732","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:735b4996a6991b906f17a70022f6746538c8c381c0162fe13f29eb90c5dd7bb7","observation_id":"3e4b5eb8-cd73-499a-91e0-a84eaaccf2c7","resolution":{"observed_at":"2026-08-09T10:44:41.188893Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.751730Z","title":null,"venue":null,"work_id":"9a1040b9-9bcb-4003-988f-a4382dfb60a7","year":2020},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.099593Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:79de951d54a6ce00d7c1100d33064bc722d21d1cad5993f77e879189e8825fc8","observation_id":"fadf6e18-da0e-4dd9-8aa3-4b714d7a6fd0","resolution":{"observed_at":"2026-08-09T10:44:41.757669Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.05576","last_updated":"2024-03-04T14:24:10Z","snapshot_observed_at":"2026-07-06T16:04:50.826084Z","submitted_at":"2023-08-10T13:39:40Z","title":"Do Language Models' Words Refer?","version":3},"cited_work":{"arxiv_id":"2308.05576","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.05576","snapshot_observed_at":"2026-08-09T10:44:41.283017Z","title":"Do Language Models' Words Refer?","venue":"cs.CL","work_id":"034afc6f-7b0f-4586-b323-ccfb3906df0e","year":2023},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.104145Z"},"links":{"cited_paper":"/paper/2308.05576","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:1b9461b3af8decdefa1c0e737f8fa00f729ecd9cf290b178c70d901a5c860613","observation_id":"e3c0eb50-c8d2-4773-b658-300f26ff2440","resolution":{"observed_at":"2026-08-09T10:44:41.288932Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04854","last_updated":"2024-06-03T17:01:06Z","snapshot_observed_at":"2026-07-06T17:13:29.467366Z","submitted_at":"2024-01-10T00:05:45Z","title":"Are Language Models More Like Libraries or Like Librarians? Bibliotechnism, the Novel Reference Problem, and the Attitudes of LLMs","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.04854","snapshot_observed_at":"2026-08-09T10:44:41.109110Z","title":"Lederman, K","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.109110Z"},"links":{"cited_paper":"/paper/2401.04854","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:0662ff4a350a58e2b182ae4b58d5914fbadd79c6263a328bc83de17a152b7070","observation_id":"8c5be06f-a1a6-461b-b91b-d31a612b9d8b","resolution":{"observed_at":"2026-08-09T10:44:41.109110Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.734573Z","title":null,"venue":null,"work_id":"d9150c68-d632-469d-a1ad-0c839e08f7cb","year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.113762Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:daa36ebec01a2ee080c476f40f5c3cd9ed5a468b7451ad984770f968e6fc6b16","observation_id":"cd10a721-29ea-475e-a296-577cbf42e9db","resolution":{"observed_at":"2026-08-09T10:44:41.740636Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.00159","last_updated":"2024-05-31T19:43:20Z","snapshot_observed_at":"2026-07-06T18:23:35.715237Z","submitted_at":"2024-05-31T19:43:20Z","title":"On the referential capacity of language models: An internalist rejoinder to Mandelkern & Linzen","version":1},"cited_work":{"arxiv_id":"2406.00159","doi":null,"metadata_source":"pith","pith_arxiv_id":"2406.00159","snapshot_observed_at":"2026-08-09T10:44:41.241281Z","title":"On the referential capacity of language models: An internalist rejoinder to Mandelkern & Linzen","venue":"cs.CL","work_id":"ab4154bf-7a75-44fc-9953-5463bfc8527c","year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.118138Z"},"links":{"cited_paper":"/paper/2406.00159","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:e96e64292f1143f81a593801166271f6b018ad1a66d3aae05cb78a749940f6e5","observation_id":"80dada28-09cd-4bed-991d-26d8474f11ab","resolution":{"observed_at":"2026-08-09T10:44:41.250029Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.716327Z","title":"Grindrod, Large language models and linguistic intentionality, Synthese 204 (2024) 71","venue":null,"work_id":"73e2a6e4-1fde-4316-9d97-b66bf8674240","year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.122847Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:10ea026e1ecd24acaacd3cc121d5e94ba88b8b9db29daa16c10472433342d5cc","observation_id":"f2d0ac24-53e6-4e65-8631-e23c60b58723","resolution":{"observed_at":"2026-08-09T10:44:41.722490Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.700061Z","title":"Berto, Topics of thought: The logic of knowledge, belief, imagination, Oxford University Press, 2022","venue":null,"work_id":"61027ad9-41af-4807-ad2b-5756362a9d56","year":2022},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.127253Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:43d744b72f051f025eaffb50a0805579d818d10dc08a805ab337c81f24ef65a4","observation_id":"f76fc0fc-9a92-496f-a0ee-56ed4520bd15","resolution":{"observed_at":"2026-08-09T10:44:41.706031Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.682281Z","title":"Hawke, Theories of aboutness, Australasian Journal of Philosophy 96 (2018) 697–723","venue":null,"work_id":"dcd1a39b-e354-4d9e-aed6-2ed01470a0ac","year":2018},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.131619Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:bbc3a2107388a42903fefa4d3ad38b18153577786f3540454e642d4751a6b89b","observation_id":"58cd10ed-76c6-42cd-ba3a-aa838abebad8","resolution":{"observed_at":"2026-08-09T10:44:41.689149Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.666434Z","title":"Hawke, L","venue":null,"work_id":"06713f31-8339-48ee-9f3f-fb99cff2e9ae","year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.136271Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:8efd1bd2d5a159207a2dbd0c29bb645e665ad3a55e193bbabd7b3fcd64f2d0e8","observation_id":"4e098e24-39a8-407e-9f3d-0918b3d110ca","resolution":{"observed_at":"2026-08-09T10:44:41.671857Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.21030","last_updated":"2025-03-14T16:14:16Z","snapshot_observed_at":"2026-07-06T18:23:27.760658Z","submitted_at":"2024-05-31T17:21:52Z","title":"Standards for Belief Representations in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.21030","snapshot_observed_at":"2026-08-09T10:44:41.140689Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.140689Z"},"links":{"cited_paper":"/paper/2405.21030","citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:6395def119612c7f2a9d97ff8443b4dd4858a6ff4de323074ba9139a1c8ccd52","observation_id":"a7064b6a-cf26-48ef-bd1f-3039b5af2f78","resolution":{"observed_at":"2026-08-09T10:44:41.140689Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:44:41.638636Z","title":"Harding, Operationalising representation in natural language processing, The British Journal for the Philosophy of Science (2023)","venue":null,"work_id":"c4782f11-033b-41a8-bc32-a8b6b19c2e35","year":2023},"citing_paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-09T10:44:41.145054Z"},"links":{"citing_paper":"/paper/2502.02896"},"observation_digest":"sha256:42285193284a9e8f51fc17a54a2eaeabb67590cc4620972e13a81206657e784e","observation_id":"5782b04b-2898-4993-941d-ee1a9e4b7846","resolution":{"observed_at":"2026-08-09T10:44:41.654092Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2502.02896","last_updated":"2025-02-05T05:37:26Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-09T10:39:32.754426Z","submitted_at":"2025-02-05T05:37:26Z","title":"A Benchmark for the Detection of Metalinguistic Disagreements between LLMs and Knowledge Graphs"},"reference_resolution":{"displayed":39,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":20,"verified_exact":7,"verified_fuzzy":12},"total_outbound_references":39},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 39 of 39 outbound references and 0 inbound Pith citation observations for arXiv:2502.02896."}