{"as_of":"2026-08-21T02:28:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:607c1d6314e49815c94f0d2406903b693355551f1a48fd7ebe43a994e2538079","coverage":[{"denominator":34,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":34,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T20:47:45.582191Z","state":"measured"},{"denominator":35,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":35,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-20T06:33:59.587034+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-28T23:15:56.598968Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-06-29T00:02:50.465540Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"cited_work":{"arxiv_id":"2505.12000","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.12000","snapshot_observed_at":"2026-06-29T00:02:50.465540Z","title":"Iqbench: How\" smart”are vision-language models? a study with human iq tests.arXiv preprint arXiv:2505.12000, 2025","venue":null,"work_id":"bf353519-6bfd-4a43-8f5c-e8356d1ce11c","year":2025},"citing_paper":{"arxiv_id":"2606.00148","last_updated":"2026-05-29T03:20:05Z","snapshot_observed_at":"2026-07-06T23:40:51.363776Z","submitted_at":"2026-05-29T03:20:05Z","title":"StemBind: When MLLMs Get Lost Between Rules and Instances in Abstract Visual Reasoning","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-06-28T23:15:56.598968Z"},"links":{"cited_paper":"/paper/2505.12000","citing_paper":"/paper/2606.00148"},"observation_digest":"sha256:528727409022cefc106f10d122f0a356cfe2925af8b5ce5a00d6d30df6989657","observation_id":"dd01d045-2af0-481f-9a13-36c057b4febe","resolution":{"observed_at":"2026-06-29T00:02:50.467451Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2505.12000/citation-record","integrity":"/paper/2505.12000/integrity","json":"/paper/2505.12000/citation-record.json","paper":"/paper/2505.12000"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.465082Z","title":"Learning transferable visual models from natural language supervision","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.465082Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:e1d7cb76e6bc3fa3304c5229b7bc789bef20b98e6abb7cd914d71f07d03e23ed","observation_id":"811f8fb8-15dd-4217-aeea-4cad982d74e0","resolution":{"observed_at":"2026-08-15T20:47:45.465082Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.904398Z","title":"Silvar-med: A speech-driven visual language model for explainable abnormality detection in medical imaging","venue":null,"work_id":"5e53c549-61b5-4ccb-9e30-91a2bab0c5cd","year":2025},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.469568Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:f56d530ff0e60a688c643b8e659c235bb050be788ef5be11c8f445b444702af8","observation_id":"d77bdb45-6dab-4c54-89a6-81d2192f7efa","resolution":{"observed_at":"2026-08-15T20:47:45.908047Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.473156Z","title":"Language models are few-shot learners","venue":null,"work_id":null,"year":1901},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.473156Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:d1adf50691c8fdff363cd3b69c7d88c58e9dbb0abf0e73b859e38bd223a92766","observation_id":"8041657e-d0dd-48cc-a873-0ecc93f54265","resolution":{"observed_at":"2026-08-15T20:47:45.473156Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.476740Z","title":"Blip-2: Bootstrapping language-image pre-training with frozen image encoders and large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.476740Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:33d538ef8cf0d0fa23d851a431239b4cc0c1e1ad6cfcd97a379a46a47f7469b4","observation_id":"c0e5018d-5e8e-4302-b7c1-99de428175bb","resolution":{"observed_at":"2026-08-15T20:47:45.476740Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.880647Z","title":"Piergiovanni, Piotr Padlewski, Daniel Salz, et al","venue":null,"work_id":"2ea1c185-1883-4ed2-9436-ff3a196e51f6","year":2023},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.480518Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:ccc6537c55db20b8569ce33e0e2599384135b9e631b42103f3591122266844aa","observation_id":"594547d5-e0c1-4ca3-b48f-3463fb55783a","resolution":{"observed_at":"2026-08-15T20:47:45.884364Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.484196Z","title":"Flamingo: a visual language model for few-shot learning","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.484196Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:555d1607f14e905aa05fc978b995548a6e169ae99653e10334f92146c08d2f6f","observation_id":"03296a20-e40c-4cfe-8905-5d1b53716463","resolution":{"observed_at":"2026-08-15T20:47:45.484196Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.16609","last_updated":"2023-09-28T17:07:49Z","snapshot_observed_at":"2026-08-20T15:31:01.041088Z","submitted_at":"2023-09-28T17:07:49Z","title":"Qwen Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.16609","snapshot_observed_at":"2026-08-15T20:47:45.488096Z","title":"Qwen technical report","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.488096Z"},"links":{"cited_paper":"/paper/2309.16609","citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:94a2a1421663c1f5a2d2acd78aec753665494d4e4b09e51e795f6cf7eaabdcc2","observation_id":"99f706e8-1c5b-4438-bd6f-efd681ec3d22","resolution":{"observed_at":"2026-08-15T20:47:45.488096Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.864366Z","title":"Artificial general intelligence: Emergence and definition","venue":null,"work_id":"c35dc644-b5bb-46bf-a712-27a86c4def8f","year":2007},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.491924Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:39fcce8aba94b69410aaa07c65d04d882844fdf5acc2b2b82e0d36a43ab3e4b7","observation_id":"a46f1eaf-4b9f-4d8c-8209-8a11e2d04438","resolution":{"observed_at":"2026-08-15T20:47:45.867715Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.854449Z","title":"Intelligence: Its structure, growth and action, volume 35","venue":null,"work_id":"fcc7713f-b398-48ca-a596-2a2a9851b440","year":1987},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.495384Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:7311e62a53021c0e376bf129ba1e7f54f93e432ff0a18f68f634072972b83703","observation_id":"37d43c10-a53b-46ec-92f5-52f6252e1c0d","resolution":{"observed_at":"2026-08-15T20:47:45.858100Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1911.01547","last_updated":"2019-11-25T13:02:04Z","snapshot_observed_at":"2026-08-17T05:38:14.090895Z","submitted_at":"2019-11-05T00:31:38Z","title":"On the Measure of Intelligence","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1911.01547","snapshot_observed_at":"2026-08-15T20:47:45.498719Z","title":"On the measure of intelligence","venue":null,"work_id":null,"year":1911},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.498719Z"},"links":{"cited_paper":"/paper/1911.01547","citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:40ac22c3e0c0295e5d4af1be5697c848c55f60c3232bfc403eb22339ecaf25a6","observation_id":"7330ebee-1b6c-45c6-8b8d-5407e0d73b84","resolution":{"observed_at":"2026-08-15T20:47:45.498719Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.09110","last_updated":"2023-10-01T21:44:23Z","snapshot_observed_at":"2026-08-10T23:10:13.900680Z","submitted_at":"2022-11-16T18:51:34Z","title":"Holistic Evaluation of Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.09110","snapshot_observed_at":"2026-08-15T20:47:45.502670Z","title":"Holistic evaluation of language models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.502670Z"},"links":{"cited_paper":"/paper/2211.09110","citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:c13806488041bb4d7ede970c4b6db3553991c7faf11221716799255622226da2","observation_id":"1e1b4266-b683-466f-8c1b-d68a8833ed8f","resolution":{"observed_at":"2026-08-15T20:47:45.502670Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.04615","last_updated":"2023-06-12T17:51:15Z","snapshot_observed_at":"2026-07-06T13:19:12.109592Z","submitted_at":"2022-06-09T17:05:34Z","title":"Beyond the Imitation Game: Quantifying and extrapolating the capabilities of language models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.04615","snapshot_observed_at":"2026-08-15T20:47:45.506519Z","title":"Beyond the imitation game: Quantifying and extrapolating the capabilities of language models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.506519Z"},"links":{"cited_paper":"/paper/2206.04615","citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:b8e3665c1fdddede3ea5cfdbb603c67c06ea9ebfef0a7b593bd55c99e225a0aa","observation_id":"7b9c4d78-6026-49f4-b8dd-fa9de88d40b0","resolution":{"observed_at":"2026-08-15T20:47:45.506519Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.509978Z","title":"Vqa: Visual question answering","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.509978Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:9eddbc09d36e6c1b7c2fb7675b458eacfa7d326a68bad515c8b3869e1399b727","observation_id":"de88bab3-3ccd-4c71-86ef-93374ee426f2","resolution":{"observed_at":"2026-08-15T20:47:45.509978Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.513284Z","title":"Gqa: A new dataset for real-world visual reasoning and compositional question answering","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.513284Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:e21c2c39d6a3d60a35d869569e9d60d49ef34082e37b253c7952234e25108b90","observation_id":"da4de3db-bccc-4e43-bf16-e045d4e09501","resolution":{"observed_at":"2026-08-15T20:47:45.513284Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.516498Z","title":"Learn to explain: Multimodal reasoning via thought chains for science question answering","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.516498Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:24b4bf07ab8fc9287c523e3f7e30dc30cbf170fe52f787e9ceb16ee5398d7766","observation_id":"81fe1045-6514-4a9c-8af1-e6fd1afc3645","resolution":{"observed_at":"2026-08-15T20:47:45.516498Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.823345Z","title":"From recognition to cognition: Visual commonsense reasoning","venue":null,"work_id":"1d7d535d-b65a-4266-8610-c17ccd7160cb","year":2019},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.519905Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:d10cc34caf4784012dba8358241cee3d3465508662631670c9a15dc3e4296cd3","observation_id":"d3cea163-a80c-451a-81ef-b8655fc866e9","resolution":{"observed_at":"2026-08-15T20:47:45.827261Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.523268Z","title":"Mmmu: A massive multi-discipline multimodal understanding and reasoning benchmark for expert agi","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.523268Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:ce7cb783d99b8bc0a47ecbc91bf094b76b75d2a9867c2764e7cbda9adc34395b","observation_id":"b502e3ff-36cf-44ee-b2ee-3a7c433bbba0","resolution":{"observed_at":"2026-08-15T20:47:45.523268Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.806327Z","title":"Mathvista: Evaluating mathematical reasoning of foundation models in visual contexts","venue":null,"work_id":"fa34cc03-9354-404f-a641-c759f8291029","year":2023},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.526373Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:d4506b9e085ff0a97c63a721b28c2281c70c6488d5b429da2e02b5108ae8cd04","observation_id":"3109e0e3-f61f-4cc5-85f4-4513bcf5177e","resolution":{"observed_at":"2026-08-15T20:47:45.810067Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.795146Z","title":"ChartQA: A benchmark for question answering about charts with visual and logical reasoning","venue":null,"work_id":"07a4ac74-6708-4d86-a479-ffbf5c6e0269","year":2022},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.529805Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:87123b02d503a28b80b2eda0123c819e19e891ebd70db74cf2ed008bf6fbdb78","observation_id":"fde4e335-535a-475f-b677-55d86ef931e2","resolution":{"observed_at":"2026-08-15T20:47:45.799033Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.784045Z","title":"Image2struct: Benchmarking structure extraction for vision-language models","venue":null,"work_id":"58564df9-69a9-4e91-ab75-a8ee6bb96592","year":2024},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.533529Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:105e67c6ef14b0b4e205d93b7cd7062c70bef930e1b66b71fbb177ad4f0efdba","observation_id":"0ace1743-ffa9-47d5-a6c1-f614407f5286","resolution":{"observed_at":"2026-08-15T20:47:45.787712Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.773465Z","title":"Know what you don’t know: Unanswerable ques- tions for squad","venue":null,"work_id":"e3fc6678-2412-4a70-9e64-e7013407eb6d","year":2018},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.537218Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:42ce9a7dff530e729474c2628bcd5c76ee485dc97da70fe4924b13b5090ebff0","observation_id":"34880561-8475-42df-957a-01f9e98c4167","resolution":{"observed_at":"2026-08-15T20:47:45.777151Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.763609Z","title":"Don’t just assume; look and answer: Overcoming priors for visual question answering","venue":null,"work_id":"63ee4bee-3c0e-4408-948f-2d16248cdf56","year":2018},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.540619Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:3c3a814ae1b57c94e141d2d9949efb434038f076c954ae5b0507aaa6dbd4a4ad","observation_id":"188c0ed8-2a56-4f63-9a10-479e38ae9617","resolution":{"observed_at":"2026-08-15T20:47:45.767149Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.753267Z","title":"Shortcut learning in deep neural networks","venue":null,"work_id":"26121bfc-8e8d-4f3e-a50a-342975efa2f4","year":2020},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.544333Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:8b291a6de9ba78270a47e3b34bcc0fdcd3bb662afd5cf28c36de30e65d0cbfdf","observation_id":"88cd7512-f17f-4769-becb-0e51c6ce91d3","resolution":{"observed_at":"2026-08-15T20:47:45.756678Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.742403Z","title":"Evaluation of openai o1: Opportunities and challenges of agi","venue":null,"work_id":"23bd6ea8-9d63-43f5-b3f2-7b07ed3075b6","year":2024},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.547733Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:e1030d07d6593728569ae841b725c7a43d11f272526e20eab112b3c3693d5485","observation_id":"47568c9e-bf56-4862-98e6-b3af4fd5d795","resolution":{"observed_at":"2026-08-15T20:47:45.746048Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21276","last_updated":"2024-10-25T17:43:01Z","snapshot_observed_at":"2026-08-15T14:02:47.366139Z","submitted_at":"2024-10-25T17:43:01Z","title":"GPT-4o System Card","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21276","snapshot_observed_at":"2026-08-15T20:47:45.551167Z","title":"Gpt-4o system card","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.551167Z"},"links":{"cited_paper":"/paper/2410.21276","citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:4e4b45e933ca67b51b0fbeae6b7405b0226856c45c43e6b141d04b5f89ad0969","observation_id":"00a15185-b594-4829-8be1-c0282f3136d3","resolution":{"observed_at":"2026-08-15T20:47:45.551167Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.554787Z","title":"Claude 3.5 sonnet, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.554787Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:59b138644156616ba5ab2f4df80c299365b70e717df82561b9b0486927cd41c7","observation_id":"77c1c264-5ab0-41e1-be61-389508b5c746","resolution":{"observed_at":"2026-08-15T20:47:45.554787Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-08-15T12:33:55.451951Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-15T20:47:45.558280Z","title":"Deepseek-r1: Incentivizing reasoning capability in llms via reinforcement learning","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.558280Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:565c7afc1052edff6429ea0dca84a20d3cf62e3fd0a43f3b215445c618bba57c","observation_id":"7410d060-0c66-43b4-95b0-a7dd94dabf32","resolution":{"observed_at":"2026-08-15T20:47:45.558280Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.725101Z","title":"Start building with gemini 2.5 flash, 2025","venue":null,"work_id":"d5202dca-6aba-43c0-aa45-5b201d4b8590","year":2025},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.561908Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:129a0a8a26b67680350e13272213f3358900ee7d6ecddbba8b3c0c880f74f5c2","observation_id":"96ff7bfa-bbd2-49d9-9c3a-ce58f0f1dcbf","resolution":{"observed_at":"2026-08-15T20:47:45.728854Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.565196Z","title":"Grok 3 beta — the age of reasoning agents, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.565196Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:e55c99da5142b8daaca8de178bc6d607220ad692bed834e679aabbfbf687424f","observation_id":"9dce250c-1982-4376-9840-33251b7480bb","resolution":{"observed_at":"2026-08-15T20:47:45.565196Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.568599Z","title":"Mmlu-pro: A more robust and challenging multi-task language understanding benchmark","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.568599Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:464456a54c6dec2c97cc740a6a180e9c994fcb1fb4ad341ae0f256672cafa5c8","observation_id":"96df672c-96c7-4108-b2a2-0908fe471689","resolution":{"observed_at":"2026-08-15T20:47:45.568599Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.20330","last_updated":"2024-04-09T15:17:50Z","snapshot_observed_at":"2026-08-07T12:15:30.838846Z","submitted_at":"2024-03-29T17:59:34Z","title":"Are We on the Right Way for Evaluating Large Vision-Language Models?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.20330","snapshot_observed_at":"2026-08-15T20:47:45.571914Z","title":"Are we on the right way for evaluating large vision-language models? arXiv preprint arXiv:2403.20330, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.571914Z"},"links":{"cited_paper":"/paper/2403.20330","citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:db797aadf1ddb35efebaf737ede9913d4169105cbb9a82018a5d1b345a2ccf0d","observation_id":"47780940-cfd7-48aa-987a-d7048a25a80a","resolution":{"observed_at":"2026-08-15T20:47:45.571914Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:47:45.697921Z","title":"Generalized planning for the abstraction and reasoning corpus","venue":null,"work_id":"3d2425a8-34fe-40e5-bb67-44e426d4e6d2","year":2024},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.575531Z"},"links":{"citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:3271901d4575119afbb2213716f9e3bff5157f2a1361e5bda7948a9f4392abf0","observation_id":"0e14f998-51c4-42ed-bc77-fe91d8d76c25","resolution":{"observed_at":"2026-08-15T20:47:45.704055Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.11557","last_updated":"2025-03-14T16:26:11Z","snapshot_observed_at":"2026-08-16T12:49:57.148289Z","submitted_at":"2025-03-14T16:26:11Z","title":"VERIFY: A Benchmark of Visual Explanation and Reasoning for Investigating Multimodal Reasoning Fidelity","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.11557","snapshot_observed_at":"2026-08-15T20:47:45.578765Z","title":"Verify: A benchmark of visual explanation and reasoning for investigating multimodal reasoning fidelity","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.578765Z"},"links":{"cited_paper":"/paper/2503.11557","citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:6973f79809482de2fffe2ce73539306193c67283a5389465272d5372a32e6bdb","observation_id":"718017ec-4abe-4d7e-b368-34f1cbed5cca","resolution":{"observed_at":"2026-08-15T20:47:45.578765Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00698","last_updated":"2025-06-04T16:20:49Z","snapshot_observed_at":"2026-08-16T03:48:36.336234Z","submitted_at":"2025-02-02T07:12:03Z","title":"MM-IQ: Benchmarking Human-Like Abstraction and Reasoning in Multimodal Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.00698","snapshot_observed_at":"2026-08-15T20:47:45.582191Z","title":"Mm-iq: Benchmarking human-like abstraction and reasoning in multimodal models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-15T20:47:45.582191Z"},"links":{"cited_paper":"/paper/2502.00698","citing_paper":"/paper/2505.12000"},"observation_digest":"sha256:7cecf5a4fd87eca3e216702af6fff31ca692eed74ca557b8321ad62e3d74f54c","observation_id":"2a5de3af-4bc8-4fdb-a7c9-2cadce4c1340","resolution":{"observed_at":"2026-08-15T20:47:45.582191Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2505.12000","last_updated":"2025-05-17T13:24:08Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-18T20:20:45.683036Z","submitted_at":"2025-05-17T13:24:08Z","title":"IQBench: How \"Smart'' Are Vision-Language Models? A Study with Human IQ Tests"},"reference_resolution":{"displayed":34,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":20,"verified_exact":0,"verified_fuzzy":14},"total_outbound_references":34},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"thesis":"As of 21 August 2026, this Paper Citation Record lists 34 of 34 outbound references and 1 inbound Pith citation observation for arXiv:2505.12000."}