{"as_of":"2026-08-09T08:07:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:be7209725f878fd62ae2a34eaf3d7df887834c3db8c2dc867ac4150bd7bc0211","coverage":[{"denominator":41,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":41,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-02T20:36:07.066080Z","state":"measured"},{"denominator":41,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":41,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2602.22971/citation-record","integrity":"/paper/2602.22971/integrity","json":"/paper/2602.22971/citation-record.json","paper":"/paper/2602.22971"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2502.13923","last_updated":"2025-02-19T18:00:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-19T18:00:14Z","title":"Qwen2.5-VL Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.13923","snapshot_observed_at":"2026-08-02T20:36:03.456501Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:03.456501Z"},"links":{"cited_paper":"/paper/2502.13923","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:097cbb603bf2ea72859956d7b326f61be1b6090cb806230ef8cbe7c2fe01e3e2","observation_id":"b93cd7b7-a8b3-46fc-acd4-9c15cf99b37a","resolution":{"observed_at":"2026-08-02T20:36:03.456501Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:03.518398Z","title":null,"venue":null,"work_id":null,"year":1986},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:03.518398Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:7bc44d00069115edbbcfde30b340b31034f1f48b5bee4cc3cf6bd1a6ba10111a","observation_id":"b4dcd2fb-6552-4c4e-8d32-97a6f8c0c07f","resolution":{"observed_at":"2026-08-02T20:36:03.518398Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:03.599604Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:03.599604Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:df3b773bcdcc89dfa73d59886bc9f78eebe36e6a6be109c91f1534ae15b3ee6b","observation_id":"5bdee61c-8ba2-44a6-a99e-3737ee616e4b","resolution":{"observed_at":"2026-08-02T20:36:03.599604Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.15472","last_updated":"2025-05-22T03:34:41Z","snapshot_observed_at":"2026-08-09T02:19:48.115567Z","submitted_at":"2025-05-21T12:48:16Z","title":"PhysicsArena: The First Multimodal Physics Reasoning Benchmark Exploring Variable, Process, and Solution Dimensions","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.15472","snapshot_observed_at":"2026-08-02T20:36:03.681688Z","title":"Dai et al","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:03.681688Z"},"links":{"cited_paper":"/paper/2505.15472","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:029cc8aa2c85e76abf2236ac6c4676271dc423d97e2d43e021e07632b065dba3","observation_id":"e722644f-bc16-4322-ac22-4a40cd8a24e2","resolution":{"observed_at":"2026-08-02T20:36:03.681688Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.14739","last_updated":"2025-03-28T15:21:44Z","snapshot_observed_at":"2026-07-29T21:41:16.650188Z","submitted_at":"2025-02-20T17:05:58Z","title":"SuperGPQA: Scaling LLM Evaluation across 285 Graduate Disciplines","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.14739","snapshot_observed_at":"2026-08-02T20:36:03.734791Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:03.734791Z"},"links":{"cited_paper":"/paper/2502.14739","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:955f416ae1d20ebdac1b2213526a551d3861e37ebc45b0e599cd5fd779da8686","observation_id":"a9ed11c0-a727-47e8-b9b0-151b9527053e","resolution":{"observed_at":"2026-08-02T20:36:03.734791Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:03.812752Z","title":null,"venue":null,"work_id":null,"year":1972},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:03.812752Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:4abbf74737482e07000492d9e2fac73cbafdca006a775c62afb4963d48c85206","observation_id":"44f79c57-c5e5-4118-8260-109cfd2980b9","resolution":{"observed_at":"2026-08-02T20:36:03.812752Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.09988","last_updated":"2024-12-13T22:03:43Z","snapshot_observed_at":"2026-07-06T19:32:42.378146Z","submitted_at":"2024-10-13T20:09:41Z","title":"HARDMath: A Benchmark Dataset for Challenging Problems in Applied Mathematics","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.09988","snapshot_observed_at":"2026-08-02T20:36:03.885587Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:03.885587Z"},"links":{"cited_paper":"/paper/2410.09988","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:99b3036cbea4b77a27a4fddf7913c4f9ff42f61ea31d357742eac348e85031aa","observation_id":"0a4b12d1-cfb6-44a7-9720-fbd6ce1667f4","resolution":{"observed_at":"2026-08-02T20:36:03.885587Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:03.974504Z","title":null,"venue":null,"work_id":null,"year":1988},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:03.974504Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:eb5a7d1643123672bc05551696191674c703e8d8b09f49d5954cd4cf120d771f","observation_id":"e99ca7c3-b2a5-4678-9f16-42a32055d4f4","resolution":{"observed_at":"2026-08-02T20:36:03.974504Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:04.041093Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:04.041093Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:8e2f8b4c428b0efef4df9bc7fce07e54e0d945c40ba43f7cb3e62d3e47612d83","observation_id":"bcb0b92a-1bd2-4dde-a7b6-de25db5e9711","resolution":{"observed_at":"2026-08-02T20:36:04.041093Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.00063","last_updated":"2024-09-29T23:49:17Z","snapshot_observed_at":"2026-08-03T12:08:32.142306Z","submitted_at":"2024-09-29T23:49:17Z","title":"Uniform price auctions with pre-announced revenue targets: Evidence from China's SEOs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.00063","snapshot_observed_at":"2026-08-02T20:36:04.141571Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:04.141571Z"},"links":{"cited_paper":"/paper/2410.00063","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:f6a526d09d644268572df7a11ff1c1d8ed868493c9f41b63e08b296b4db4bd60","observation_id":"c642fb21-1872-4630-9e01-3cd78695b477","resolution":{"observed_at":"2026-08-02T20:36:04.141571Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-02T20:36:04.282898Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:04.282898Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:21979bd03423c13dc935f9a1e3797a4f279a96ac835d316472cdfefe2a3a7a08","observation_id":"9553ac3a-bc57-4d9c-a6b8-639ca6ccf701","resolution":{"observed_at":"2026-08-02T20:36:04.282898Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.05444","last_updated":"2025-01-09T18:55:52Z","snapshot_observed_at":"2026-08-07T17:40:59.052312Z","submitted_at":"2025-01-09T18:55:52Z","title":"Can MLLMs Reason in Multimodality? EMMA: An Enhanced MultiModal ReAsoning Benchmark","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.05444","snapshot_observed_at":"2026-08-02T20:36:04.457394Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:04.457394Z"},"links":{"cited_paper":"/paper/2501.05444","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:c7ca2a30f7eb31ba0d0b8ae906c860d012dcd2af08620c4a600ea9e9c8075d73","observation_id":"b8dcbd05-29aa-456e-a0fa-05bb7fc6ef92","resolution":{"observed_at":"2026-08-02T20:36:04.457394Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:04.626985Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:04.626985Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:5c86e911bdf05a5d1e55112e8e27021cdadb8e72e4119eca1892761e3984f9e9","observation_id":"7cf8735f-616a-4392-b2af-cadfc9771d67","resolution":{"observed_at":"2026-08-02T20:36:04.626985Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:04.707502Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:04.707502Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:d84807217bc6c34f5fc5ccf5c52fef7e171010cc95e12c8e61c5b5e66930e7ac","observation_id":"e3cab88a-a664-48e9-b69e-fbe4191a1cbc","resolution":{"observed_at":"2026-08-02T20:36:04.707502Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:04.816225Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:04.816225Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:26fb6f9aa719f17a99640d24db41fbf737a1444e33fe4753eb0b96aa3609e18c","observation_id":"1869e037-e6f9-4501-92b4-9d5e5c255d8c","resolution":{"observed_at":"2026-08-02T20:36:04.816225Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2207.05221","last_updated":"2022-11-21T16:38:35Z","snapshot_observed_at":"2026-08-06T08:34:11.887259Z","submitted_at":"2022-07-11T22:59:39Z","title":"Language Models (Mostly) Know What They Know","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2207.05221","snapshot_observed_at":"2026-08-02T20:36:04.879932Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:04.879932Z"},"links":{"cited_paper":"/paper/2207.05221","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:c0dbdddf0cef110cd7c04d26f27a5282edc232d21b32efcbce8e64c0148c4133","observation_id":"8728fc90-ac8a-4eab-8b90-3807c3ed36f0","resolution":{"observed_at":"2026-08-02T20:36:04.879932Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:04.975187Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:04.975187Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:2cb7639e3574bd27f78cbf7e22788c1ca1d5cdd9095668e8fe06b79db026f0ac","observation_id":"2abfdf59-bdb5-452f-b4d4-37cddf0299f3","resolution":{"observed_at":"2026-08-02T20:36:04.975187Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:05.073625Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:05.073625Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:9c215ca2cbc0758803c76ed016cacb2c37289b0ee952c3c2a4a4ca2d431809bc","observation_id":"42c3cf1f-a0ef-4749-adce-24d26d725e20","resolution":{"observed_at":"2026-08-02T20:36:05.073625Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.11284","last_updated":"2025-01-20T05:44:01Z","snapshot_observed_at":"2026-08-06T13:30:45.600197Z","submitted_at":"2025-01-20T05:44:01Z","title":"RedStar: Does Scaling Long-CoT Data Unlock Better Slow-Reasoning Systems?","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.11284","snapshot_observed_at":"2026-08-02T20:36:05.177375Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:05.177375Z"},"links":{"cited_paper":"/paper/2501.11284","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:334cd6c4916da1d398f2d95faade4ed0c7a522d9932dd9d417e45b4926148e6b","observation_id":"68011b37-55ec-4367-a506-ff6573385c05","resolution":{"observed_at":"2026-08-02T20:36:05.177375Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.16634","last_updated":"2023-05-23T22:12:16Z","snapshot_observed_at":"2026-08-02T04:02:36.848064Z","submitted_at":"2023-03-29T12:46:54Z","title":"G-Eval: NLG Evaluation using GPT-4 with Better Human Alignment","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.16634","snapshot_observed_at":"2026-08-02T20:36:05.329680Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:05.329680Z"},"links":{"cited_paper":"/paper/2303.16634","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:883f25688e9209e79e590f673ee1ec5bba2700db378e00517799210507892800","observation_id":"6b0ddadd-d6b1-4f83-af87-6bb32fb0bef5","resolution":{"observed_at":"2026-08-02T20:36:05.329680Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:05.435663Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:05.435663Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:31b3f5686b6155eab0da19c25abfb77554a8cf30c64a261831a9214b469fa651","observation_id":"2c903030-e23e-4f95-bf68-899534f3a981","resolution":{"observed_at":"2026-08-02T20:36:05.435663Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:05.547251Z","title":"O’Gorman","venue":null,"work_id":null,"year":1993},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:05.547251Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:3cb12de0fd5d749a3e313274325e6cd83eb2494345aa43df2ff83b21d22df4c1","observation_id":"0fbe72ec-4589-49ad-9562-38047b97996c","resolution":{"observed_at":"2026-08-02T20:36:05.547251Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.14249","last_updated":"2026-02-20T04:23:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-24T05:27:46Z","title":"Humanity's Last Exam","version":10},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.14249","snapshot_observed_at":"2026-08-02T20:36:05.644586Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:05.644586Z"},"links":{"cited_paper":"/paper/2501.14249","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:5e59d2efa730b34fd56210c51c5b571ff9334695571bd7161d75a2d1853ac453","observation_id":"58ff00c2-1462-4f30-ab0d-8fbe86447e0c","resolution":{"observed_at":"2026-08-02T20:36:05.644586Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.15929","last_updated":"2025-02-21T20:56:34Z","snapshot_observed_at":"2026-08-07T17:57:31.077422Z","submitted_at":"2025-02-21T20:56:34Z","title":"Approximate Differential Privacy of the $\\ell_2$ Mechanism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.15929","snapshot_observed_at":"2026-08-02T20:36:05.729513Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:05.729513Z"},"links":{"cited_paper":"/paper/2502.15929","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:ac5f87e690de9448ba4d7a9bda8bab90da659e5ee3a52eaa945395130759132f","observation_id":"eec9d077-8ef3-4230-b425-b710656a10d7","resolution":{"observed_at":"2026-08-02T20:36:05.729513Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1145/3728725","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"2025.From Chart to QA Pairs: A Context-A ware Generation Framework for Chart-Containing Documents","venue":null,"work_id":"011b917c-cf69-40d6-9564-c36cabf69cde","year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:05.829330Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:a020873c963c5f45f03a2a57a5e694f37c93cd894b49f8f0c49c702198251660","observation_id":"63ef16db-1005-4f0a-801b-f1d0d3467d05","resolution":{"observed_at":"2026-08-02T20:39:35.135560Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:05.911860Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:05.911860Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:e15f77174184c71e3a713468542114e4d41982175d722e747fb3dcbdc6c51bba","observation_id":"e4da7aed-b833-48df-9024-865354294996","resolution":{"observed_at":"2026-08-02T20:36:05.911860Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:06.109922Z","title":null,"venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:06.109922Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:3d8d731e9db3a450c8a17f88d805c5c8db752e50296649d340a225bd85f72d11","observation_id":"13285e2c-d00c-4008-a430-f98533df6adf","resolution":{"observed_at":"2026-08-02T20:36:06.109922Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2203.11171","last_updated":"2023-03-07T17:57:37Z","snapshot_observed_at":"2026-07-06T12:50:22.773056Z","submitted_at":"2022-03-21T17:48:52Z","title":"Self-Consistency Improves Chain of Thought Reasoning in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2203.11171","snapshot_observed_at":"2026-08-02T20:36:06.193932Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:06.193932Z"},"links":{"cited_paper":"/paper/2203.11171","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:c155ebfac2b732900410e2e216aa1e31a4a44ca7175c7bd4ccd3a685c61a704d","observation_id":"0ce3803e-3cdf-44a6-813d-b33eb51ca2b4","resolution":{"observed_at":"2026-08-02T20:36:06.193932Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:06.280767Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:06.280767Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:823284fa2de473fd26c2f16cffd16edc16e11326b2e38e8fb819ea8e56ff24dd","observation_id":"99468108-72e7-4114-96a3-0b6f7142a8ec","resolution":{"observed_at":"2026-08-02T20:36:06.280767Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:06.315900Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:06.315900Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:efacda22182124d20aec3acf6a2a4b3039fc5e750ae2d4db82ae5009db815ba0","observation_id":"a711866c-4df5-41cd-add7-1f5f5fc8793f","resolution":{"observed_at":"2026-08-02T20:36:06.315900Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15114","last_updated":"2025-05-27T03:32:23Z","snapshot_observed_at":"2026-08-01T07:23:06.300790Z","submitted_at":"2024-11-22T18:30:46Z","title":"RE-Bench: Evaluating frontier AI R&D capabilities of language model agents against human experts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.15114","snapshot_observed_at":"2026-08-02T20:36:06.369914Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:06.369914Z"},"links":{"cited_paper":"/paper/2411.15114","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:6daa8e8f676c906589e9ca0c52409be23f0e336321d273a0ac63d39e52688d09","observation_id":"9126949c-9665-45a0-b476-7436e91deb07","resolution":{"observed_at":"2026-08-02T20:36:06.369914Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:06.429817Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:06.429817Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:4eabe8d22157c3759591c97c9e7b2c4a403f2e2865679b64f36772ac0927aaf5","observation_id":"781ac77b-6663-4f4a-820f-4405bc89ebc2","resolution":{"observed_at":"2026-08-02T20:36:06.429817Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00334","last_updated":"2025-06-03T07:13:03Z","snapshot_observed_at":"2026-07-06T20:29:26.524162Z","submitted_at":"2025-02-01T06:42:02Z","title":"UGPhysics: A Comprehensive Benchmark for Undergraduate Physics Reasoning with Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.00334","snapshot_observed_at":"2026-08-02T20:36:06.501254Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:06.501254Z"},"links":{"cited_paper":"/paper/2502.00334","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:53947de61d9bbaa9ea9341b0ffbd5c9e68c25ca9498d15f95254bca207bd972e","observation_id":"8990da40-1d4b-4f9b-b376-85f45e29373e","resolution":{"observed_at":"2026-08-02T20:36:06.501254Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12122","last_updated":"2024-09-18T16:45:37Z","snapshot_observed_at":"2026-07-06T19:17:41.512834Z","submitted_at":"2024-09-18T16:45:37Z","title":"Qwen2.5-Math Technical Report: Toward Mathematical Expert Model via Self-Improvement","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12122","snapshot_observed_at":"2026-08-02T20:36:06.560810Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:06.560810Z"},"links":{"cited_paper":"/paper/2409.12122","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:dae29be39cbddf0211e431068d16a398a423b63b4475c58ad996d10776c402f5","observation_id":"8e455da9-b0ee-432b-b59a-aad95c47fa0b","resolution":{"observed_at":"2026-08-02T20:36:06.560810Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:06.609031Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:06.609031Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:35794700ab1426c398e88674d01771033a6bfadbf78a718ec0b644c74d28e5ce","observation_id":"e93b4a60-cc03-43c2-80ec-dfbcb4a1446a","resolution":{"observed_at":"2026-08-02T20:36:06.609031Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:06.650050Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:06.650050Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:1855e238660fb1d69d08df61f0a3cc3b4a555fbfb73352bb30107da8d5a081c1","observation_id":"1f80385e-99cf-4a7f-a3bf-78ec598f5fc3","resolution":{"observed_at":"2026-08-02T20:36:06.650050Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.12054","last_updated":"2025-05-26T13:42:06Z","snapshot_observed_at":"2026-08-09T05:38:56.887390Z","submitted_at":"2025-02-17T17:24:14Z","title":"PhysReason: A Comprehensive Benchmark towards Physics-Based Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.12054","snapshot_observed_at":"2026-08-02T20:36:06.714065Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:06.714065Z"},"links":{"cited_paper":"/paper/2502.12054","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:3ccf882c5f229048a49b779a5e7e6c88310ff3ac4d3ffe1a7c9dfb842fb10674","observation_id":"2de830a1-1420-4e5c-807f-b3a24c0bb779","resolution":{"observed_at":"2026-08-02T20:36:06.714065Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13372","last_updated":"2024-06-27T22:44:48Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-03-20T08:08:54Z","title":"LlamaFactory: Unified Efficient Fine-Tuning of 100+ Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.13372","snapshot_observed_at":"2026-08-02T20:36:06.817627Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:06.817627Z"},"links":{"cited_paper":"/paper/2403.13372","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:85f00abb5633aaad14f9efe8dba1a03be5c883bde6cdec00b7c22c893b3ad3ee","observation_id":"359f3c9a-3551-43fa-81a9-c2aefd339c6d","resolution":{"observed_at":"2026-08-02T20:36:06.817627Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.06364","last_updated":"2023-09-18T14:23:02Z","snapshot_observed_at":"2026-08-07T12:51:50.861720Z","submitted_at":"2023-04-13T09:39:30Z","title":"AGIEval: A Human-Centric Benchmark for Evaluating Foundation Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.06364","snapshot_observed_at":"2026-08-02T20:36:06.954126Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:06.954126Z"},"links":{"cited_paper":"/paper/2304.06364","citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:b6f6c32e1b8588d2d7651c5466a10ffb0db0b80db1f139064f0def5e72376044","observation_id":"d0910937-4558-4948-90de-e21ad3005511","resolution":{"observed_at":"2026-08-02T20:36:06.954126Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T20:36:07.066080Z","title":"PhD-level","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:07.066080Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:47c1de122e0287c26dfcb91696217df7f6161b459adbff54f2b90cb4870b6bb6","observation_id":"6b79a6d5-589d-4658-992f-86888e6ac4c8","resolution":{"observed_at":"2026-08-02T20:36:07.066080Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1038/s41598-025-18622-6","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"https://doi.org/ 10.1038/s41598-025-18622-6","venue":"Scientific Reports","work_id":"597e3378-c08c-4a9e-a8b6-c0e60d7f697f","year":2025},"citing_paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy","version":2},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:05.981523Z"},"links":{"citing_paper":"/paper/2602.22971"},"observation_digest":"sha256:804befe2e5e167d1ef9cf4ce961f1a0c217ffb3da950a8134a9015b78463ced2","observation_id":"719e3fa9-de86-4c55-9829-8602e94955a1","resolution":{"observed_at":"2026-08-02T20:39:34.844758Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2602.22971","last_updated":"2026-05-29T07:44:51Z","latest_version":2,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-06T11:32:52.767125Z","submitted_at":"2026-02-26T13:08:56Z","title":"SPM-Bench: Benchmarking Large Language Models for Scanning Probe Microscopy"},"reference_resolution":{"displayed":41,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":39,"verified_exact":2,"verified_fuzzy":0},"total_outbound_references":41},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 41 of 41 outbound references and 0 inbound Pith citation observations for arXiv:2602.22971."}