{"as_of":"2026-08-23T12:00:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:96f08912e124716ea17a4f8ec325031daae3e9f207a67c78c306fee864b9e823","coverage":[{"denominator":32,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":32,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-27T02:17:49.019550Z","state":"measured"},{"denominator":32,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":32,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-23T06:30:58.430688+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2606.17433/citation-record","integrity":"/paper/2606.17433/integrity","json":"/paper/2606.17433/citation-record.json","paper":"/paper/2606.17433"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2601.03267","last_updated":"2026-05-01T23:55:43Z","snapshot_observed_at":"2026-08-15T00:17:32.875866Z","submitted_at":"2025-12-19T07:05:38Z","title":"OpenAI GPT-5 System Card","version":2},"cited_work":{"arxiv_id":"2601.03267","doi":"10.48550/arxiv.2601.03267","metadata_source":"pith","pith_arxiv_id":"2601.03267","snapshot_observed_at":"2026-08-05T02:49:54.815029Z","title":"OpenAI GPT-5 System Card","venue":"cs.CL","work_id":"ca87689a-0d29-4476-b504-b65dbbb08af4","year":2025},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"cited_paper":"/paper/2601.03267","citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:ad65c0add4ee47fb70d2c1383d127af7f7982260e6e1d7a58ad425da624a95c3","observation_id":"e7c88955-a72f-4db0-8d07-a3e645d6b3db","resolution":{"observed_at":"2026-07-03T19:08:49.452425Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-22T05:08:14.623228+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-22T05:08:14.623228+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.11805","last_updated":"2025-05-09T21:04:06Z","snapshot_observed_at":"2026-08-20T18:27:04.837880Z","submitted_at":"2023-12-19T02:39:27Z","title":"Gemini: A Family of Highly Capable Multimodal Models","version":5},"cited_work":{"arxiv_id":"2312.11805","doi":"10.1038/nrn2888","metadata_source":"pith","pith_arxiv_id":"2312.11805","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Gemini: A Family of Highly Capable Multimodal Models","venue":"cs.CL","work_id":"83f7c85b-3f11-450f-ac0c-64d9745220b2","year":2023},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"cited_paper":"/paper/2312.11805","citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:71ee69dec559a774a4e3a98032ea38f7808d2fd04ac1ca06c1c5ef2fda870375","observation_id":"2352fad6-87d1-49a2-afc5-a399933be957","resolution":{"observed_at":"2026-07-03T19:08:49.455259Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Mmbench: Is your multi-modal model an all-around player?","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:ba5cb3fe409d5029f8d952c976936f0f279540010d52a89531290583a6d70e9e","observation_id":"6e71a3bb-398e-4b67-af80-64ccaddba30c","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Logic unseen: Revealing the logical blindspots of vision-language models,","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:26f5bb1b99f2df6f8baa7393579e81bec7cff203b7cfd5b557eb3733521ef84a","observation_id":"9518697d-fb1e-42b2-bc65-99c815a70e29","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Logicqa: Logical anomaly detection with vision language model generated questions,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:fadaf78dd6ba937b5bd571cc68da59baeba7cf9475681ad4bd1986969238bc90","observation_id":"024aac2b-4c29-4c12-ad0d-1e380a56fe4f","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Vision-language models can’t see the obvious,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:6cb6679b2f676966b5976a074cd35d1b5cc3d082ca2692069c54135ab2fcbb23","observation_id":"18e3e7d3-17eb-458d-a5bd-8c63d12939f0","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2603.16952","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T19:08:49.459838Z","title":"Embodied foundation models at the edge: A survey of deployment constraints and mitigation strategies","venue":null,"work_id":"9bf76be7-0f63-4840-a14a-ee36f8facb47","year":2026},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:0e2c07d39ea1f8ebf60b72159de30e5ec83f3a23b7c7fec76c575fa265ffb329","observation_id":"580bf83f-033b-4fb4-934d-ab4dd165aedb","resolution":{"observed_at":"2026-07-03T19:08:49.462296Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Spotting the unexpected (stu): A 3d lidar dataset for anomaly seg- mentation in autonomous driving,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:bdec92e8d09e91d437bb35cf0e408a4ee76c355ff5e91f2107d02c3144417b77","observation_id":"56151c8b-6c55-48ae-a08d-06d93b8aca14","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"gpt-5-nano,","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:ae9eaf21feb761595b23b59495f6f905ce26f2f8d7fe565f31bc6e6b5c66979a","observation_id":"64714918-5420-4d4b-af52-d24e86b8046b","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Plovad: Prompting vision- language models for open vocabulary video anomaly detection,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:081486f497e563b22211c0643530fb58edd05aca7e88265c6bd1d217ebee614e","observation_id":"e32ae8f5-f3f0-46de-ad93-e1f3eecfa732","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Video anomaly detection in 10 years: A survey and outlook,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:db1ef6f7212ce26924a05feb7ed3eeb61a4be5d84945e3c68e4135a06e4fb7e9","observation_id":"3139b6cb-2882-49e0-829f-74c2cfef2601","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Nesylad: A neuro-symbolic approach for unsupervised logi- cal anomaly detection,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:7aeceff61340a83b82b070649ed4df1e7c216cc1f7bdb033832e94d587128b5a","observation_id":"99169d5e-3807-49f6-8a7e-83fc13ee3dc8","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Towards training-free anomaly detection with vision and language foundation models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:2c7f2ce1a52db14a7269a9941fd3ac440b6d87122bcaf05bbf48802be84a49b1","observation_id":"f0548e74-f3f6-4908-aa3e-00d2ed815df1","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Madclip: few-shot medical anomaly detection with clip,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:0e072fa7560f36611c3dcc3bbc14c9687182bcb511a08977337c00a5d3e3d128","observation_id":"84e2162b-e41e-44bf-8b41-9120be79a244","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Adapting visual-language models for generalizable anomaly detection in medical images,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:b770b0d6894d4807ace40081c9f3f3f4ddfb5b82e04b7656698a929edf92225d","observation_id":"ef1b1475-44c8-458f-90be-70476f42d503","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Smarthome-bench: A comprehensive benchmark for video anomaly detection in smart homes using multi-modal large language models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:a3eaf342d8018b8cb23a4cccf09955758c429295913a32bc7c24d7ae8d090b1e","observation_id":"41da329b-f357-45b3-b1b1-9f8f27b2dc62","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Sequential keypoint density estimator: an overlooked baseline of skeleton-based video anomaly detection,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:49a06a87db52651776730fd31baaed803e52d87bdd88846d044f752f285b97fc","observation_id":"2e567bae-7aed-42a5-9d77-8cfde723b73d","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Vane-bench: Video anomaly evaluation benchmark for conversational lmms,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:1a760bc2aca7a1e546941ec4d5a8eeba45a3f4cba6c801a907bd037272997a79","observation_id":"8c9bfdc5-9621-401c-82cf-510865703ded","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Logicad: Explainable anomaly detection via vlm-based text feature extraction,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:d2410497f571f1ea902a63a84e06c587be767b4a00737f9a0fc57cf087848640","observation_id":"39e50aaa-6b8f-40cf-b0b1-5a9ac6c7a4fd","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Gpt-4v-ad: Exploring grounding potential of vqa-oriented gpt- 4v for zero-shot anomaly detection,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:c21c7504cc9a70ab912c7932cdfe956fa4159f5702f779ce11bd59109f27842e","observation_id":"6c85a7a3-8a7e-4921-a783-97385cda46f6","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Vbench: Comprehensive benchmark suite for video generative models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:23df0c2c3b8936b9c8b5b3e7ab4d9346a8233bfe9aceb042884b8a5d42b19fdf","observation_id":"c85c487f-3064-4c53-875c-366fa315c325","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Multi-rag: A multimodal retrieval- augmented generation system for adaptive video understanding,","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:bbea0f0c1a4ae8c3c79662361d5b0750a6352a6ee90a1fea31c868dcf703b424","observation_id":"f7022d8e-730d-435d-b8f6-b5eae83f4fae","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"Cave: Detecting and explaining commonsense anomalies in visual environments,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:90f2c6d4edfbc0cac97b457e214ca100718234a0dab4bb9b3109eb6073eedd19","observation_id":"e77f55b4-1968-4786-a9ed-460eb3683ace","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.31410","last_updated":"2026-05-29T15:13:53Z","snapshot_observed_at":"2026-08-03T04:10:32.210591Z","submitted_at":"2026-05-29T15:13:53Z","title":"FAM-Bench: A Multimodal Benchmark for Condition-Aware Food-as-Medicine Reasoning","version":1},"cited_work":{"arxiv_id":"2605.31410","doi":null,"metadata_source":"pith","pith_arxiv_id":"2605.31410","snapshot_observed_at":"2026-07-03T19:08:49.456428Z","title":"FAM-Bench: A Multimodal Benchmark for Condition-Aware Food-as-Medicine Reasoning","venue":"cs.AI","work_id":"18969fd3-fce9-416d-98de-7215ff869404","year":2026},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"cited_paper":"/paper/2605.31410","citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:61aa82955f0c7fbd3e3306ecb2531ac9b2e8bdf9076f4682ce1a3f8bbd5cc22e","observation_id":"7bc9c847-16e8-4457-a7fb-c5f45f91b4dd","resolution":{"observed_at":"2026-07-03T19:08:49.458294Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"gpt-image-1,","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:ecc9d8496099fbb7e0315e31e9ba9da01ffb9aedfd45ec733a371c85c98ec3ce","observation_id":"ede7f848-765c-468d-a6f4-5ce1faad4303","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"gpt-5-mini,","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:60b3147e633cecdbabe5295b9e1aa97a82ec4f66f80581eb64586ce72c1cf517","observation_id":"09c33f6f-249e-40ad-be17-0d79d619206e","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"[Online]","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:a182f2278d18227c21e991b79efebdbc9c8206179e6a1a26a2d14b9155f2fdcb","observation_id":"8170f364-5ab1-4208-b5fa-26018fa8c427","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"claude-sonnet-4-6,","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:8b66cd8d9e0849c50badf9a53ddea4dcc0cf23a57d50e6e867799bbcd64633c9","observation_id":"36c4b20a-ecd4-4bd9-aab6-ac2c7b5dc88a","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"gemini-3-flash-preview,","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:1b51937d018f1e24c68fefdde2f68312e24a9d7f3db87fb4ee0916f92f561fc7","observation_id":"5053256f-2367-483e-ac65-3c93ea7a9f20","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T02:17:49.019550Z","title":"grok-4.1-fast-reasoning,","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:8da938fd5ae1b4a4d8aad08ab17d67b43682b791fe4416c61763bcb1c1b11d2d","observation_id":"ab91e0cb-6928-4988-8a30-d5158ece00ae","resolution":{"observed_at":"2026-06-27T02:17:49.019550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-08-13T17:20:44.002518Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":"2407.21783","doi":"10.1016/s0749-0720(15","metadata_source":"pith","pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"The Llama 3 Herd of Models","venue":"cs.AI","work_id":"1549a635-88af-4ac1-acfe-51ae7bb53345","year":2024},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:b35109a5d153053cfb8f4ef7f904d7ebbf81021213c8fcecdeaec430c5a6e3e0","observation_id":"0f9c97cb-a115-4dd5-83c4-0e80a6003762","resolution":{"observed_at":"2026-07-03T19:08:49.454964Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.21631","last_updated":"2025-11-27T12:16:54Z","snapshot_observed_at":"2026-08-17T13:26:10.378579Z","submitted_at":"2025-11-26T17:59:08Z","title":"Qwen3-VL Technical Report","version":2},"cited_work":{"arxiv_id":"2511.21631","doi":"10.1016/j.neunet.2025.107777","metadata_source":"pith","pith_arxiv_id":"2511.21631","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Qwen3-VL Technical Report","venue":"cs.CV","work_id":"1fe243aa-e3c0-4da6-b391-4cbcfc88d5c0","year":2025},"citing_paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-06-27T02:17:49.019550Z"},"links":{"cited_paper":"/paper/2511.21631","citing_paper":"/paper/2606.17433"},"observation_digest":"sha256:d1d5c226486c71a29464084b417dc2523ce235ed607f1123242be6d32101cdd7","observation_id":"03fb8593-2b9b-4203-9c38-f69e64f9143d","resolution":{"observed_at":"2026-07-03T19:08:49.452095Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2606.17433","last_updated":"2026-06-16T02:32:38Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-15T01:21:27.365355Z","submitted_at":"2026-06-16T02:32:38Z","title":"LADBench: A Benchmark for Logical Fault Detection in Images"},"reference_resolution":{"displayed":32,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":26,"verified_exact":6,"verified_fuzzy":0},"total_outbound_references":32},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"thesis":"As of 23 August 2026, this Paper Citation Record lists 32 of 32 outbound references and 0 inbound Pith citation observations for arXiv:2606.17433."}