{"as_of":"2026-08-10T04:30:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:14012830bf88ab8dfb46fc9dca329ce6cdc895e027fccead5ff8e642cad3ec4e","coverage":[{"denominator":14,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":14,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T23:41:25.883309Z","state":"measured"},{"denominator":19,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":19,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":5,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":5,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-01T07:22:25.349398Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-01T20:56:14.268938Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"cited_work":{"arxiv_id":"2502.08796","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.08796","snapshot_observed_at":"2026-07-01T20:56:14.268938Z","title":"arXiv preprint arXiv:2502.08796 , year=","venue":null,"work_id":"7611978f-87e0-4c4d-9d33-58a0653a4882","year":2025},"citing_paper":{"arxiv_id":"2602.10298","last_updated":"2026-04-26T12:37:24Z","snapshot_observed_at":"2026-08-03T20:30:47.690102Z","submitted_at":"2026-02-10T21:12:12Z","title":"On Emergent Social World Models -- Evidence for Functional Integration of Theory of Mind and Pragmatic Reasoning in Language Models","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-16T05:13:40.377973Z"},"links":{"cited_paper":"/paper/2502.08796","citing_paper":"/paper/2602.10298"},"observation_digest":"sha256:f6e31553901c3dcd57b55a80de1c20a24dd2493fca416f8a9cf18a6452f86c84","observation_id":"8582d296-503b-4c2f-9de2-1188a5a317e3","resolution":{"observed_at":"2026-05-16T05:17:22.486999Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"cited_work":{"arxiv_id":"2502.08796","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.08796","snapshot_observed_at":"2026-07-01T20:56:14.268938Z","title":"arXiv preprint arXiv:2502.08796 , year=","venue":null,"work_id":"7611978f-87e0-4c4d-9d33-58a0653a4882","year":2025},"citing_paper":{"arxiv_id":"2605.15205","last_updated":"2026-04-28T15:38:31Z","snapshot_observed_at":"2026-08-08T12:32:02.265510Z","submitted_at":"2026-04-28T15:38:31Z","title":"Does Theory of Mind Improvement Really Benefit Human-AI Interactions? Empirical Findings from Interactive Evaluations","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-05-19T17:55:04.458343Z"},"links":{"cited_paper":"/paper/2502.08796","citing_paper":"/paper/2605.15205"},"observation_digest":"sha256:48077c8680efaaafaef6372be27df4151c3d9c1e286aa759de7268eca05010d4","observation_id":"456439ea-fb2f-4fcf-842a-9d31ce88f359","resolution":{"observed_at":"2026-05-19T17:57:42.630853Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"cited_work":{"arxiv_id":"2502.08796","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.08796","snapshot_observed_at":"2026-07-01T20:56:14.268938Z","title":"arXiv preprint arXiv:2502.08796 , year=","venue":null,"work_id":"7611978f-87e0-4c4d-9d33-58a0653a4882","year":2025},"citing_paper":{"arxiv_id":"2606.01145","last_updated":"2026-07-06T16:09:36Z","snapshot_observed_at":"2026-07-30T15:24:45.114962Z","submitted_at":"2026-05-31T10:27:18Z","title":"Reasoning4Sciences: Bridging Reasoning Language Models to All Scientific Branches","version":2},"reference_index":244,"source":"pdf_text","source_observed_at":"2026-06-28T17:35:01.285534Z"},"links":{"cited_paper":"/paper/2502.08796","citing_paper":"/paper/2606.01145"},"observation_digest":"sha256:55e980f8680c99d6342b196e8bd00db889812d58ccc5c479b1a85ffaa83f8887","observation_id":"aa045270-4f7a-4551-8351-c8cd9393df86","resolution":{"observed_at":"2026-07-01T20:56:14.270428Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"cited_work":{"arxiv_id":"2502.08796","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.08796","snapshot_observed_at":"2026-07-01T20:56:14.268938Z","title":"arXiv preprint arXiv:2502.08796 , year=","venue":null,"work_id":"7611978f-87e0-4c4d-9d33-58a0653a4882","year":2025},"citing_paper":{"arxiv_id":"2606.01145","last_updated":"2026-07-06T16:09:36Z","snapshot_observed_at":"2026-07-30T15:24:45.114962Z","submitted_at":"2026-05-31T10:27:18Z","title":"Reasoning4Sciences: Bridging Reasoning Language Models to All Scientific Branches","version":3},"reference_index":260,"source":"pdf_text","source_observed_at":"2026-07-01T07:22:25.349398Z"},"links":{"cited_paper":"/paper/2502.08796","citing_paper":"/paper/2606.01145"},"observation_digest":"sha256:26970311b811d22bfd01a77ae093bfeaf1e2af242622d08ed4d76d74cc0ad3bd","observation_id":"99b0aac6-1f44-4960-9505-0bd796011684","resolution":{"observed_at":"2026-07-01T08:35:34.033269Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"cited_work":{"arxiv_id":"2502.08796","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.08796","snapshot_observed_at":"2026-07-01T20:56:14.268938Z","title":"arXiv preprint arXiv:2502.08796 , year=","venue":null,"work_id":"7611978f-87e0-4c4d-9d33-58a0653a4882","year":2025},"citing_paper":{"arxiv_id":"2606.31916","last_updated":"2026-06-30T16:22:12Z","snapshot_observed_at":"2026-08-02T10:24:05.379334Z","submitted_at":"2026-06-30T16:22:12Z","title":"Theory of Mind and Persuasion Beyond Conversation: Assessing the Capacity of LLMs to Induce Belief States via Planning and Action","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-07-01T05:40:54.002702Z"},"links":{"cited_paper":"/paper/2502.08796","citing_paper":"/paper/2606.31916"},"observation_digest":"sha256:0f511a9e551a88d4d674473a39d118832feba9fa519e0c801233629dd44648d5","observation_id":"bbbaa650-2d14-487b-89f2-25e2a1c38aea","resolution":{"observed_at":"2026-07-01T10:15:44.636103Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2502.08796/citation-record","integrity":"/paper/2502.08796/integrity","json":"/paper/2502.08796/citation-record.json","paper":"/paper/2502.08796"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"1612.01175","last_updated":"2017-03-31T16:36:53Z","snapshot_observed_at":"2026-07-06T05:21:18.233606Z","submitted_at":"2016-12-04T20:45:42Z","title":"Who is Mistaken?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1612.01175","snapshot_observed_at":"2026-08-07T23:41:25.839197Z","title":"Susanne A","venue":null,"work_id":null,"year":1986},"citing_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T23:41:25.839197Z"},"links":{"cited_paper":"/paper/1612.01175","citing_paper":"/paper/2502.08796"},"observation_digest":"sha256:8d78aba9b9629ff3658e91e8ebb0e00635f01f50a952c920d761559534f17962","observation_id":"ca079f7b-05f2-4fdc-8ea2-f058cecdcc6a","resolution":{"observed_at":"2026-08-07T23:41:25.839197Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T23:41:26.059898Z","title":"Rabeeh Karimi Mahabadi, Yonatan Belinkov, and James Henderson","venue":null,"work_id":"1fdce4ef-d0bf-4545-aa58-2bbd53024fed","year":2020},"citing_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T23:41:25.844911Z"},"links":{"citing_paper":"/paper/2502.08796"},"observation_digest":"sha256:cb71f24a7031b8e39378a8d17797e0fec318771a69005ae48f55590f7203ed43","observation_id":"d69feb74-5de3-4ce2-b52f-40cca0141306","resolution":{"observed_at":"2026-08-07T23:41:26.065083Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T23:41:25.994729Z","title":null,"venue":null,"work_id":"446a439c-3208-4841-a7f9-016656be46c2","year":2021},"citing_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T23:41:25.863301Z"},"links":{"citing_paper":"/paper/2502.08796"},"observation_digest":"sha256:66f393b02672a6f4ee8cad49a8f5fd4c294b2d6512491d3763083274586233ca","observation_id":"5766bd49-031c-4d77-8ea7-165b610321ad","resolution":{"observed_at":"2026-08-07T23:41:25.999510Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T23:41:25.978180Z","title":"In Proceedings of the 2023 Conference on Empirical Methods in Natural Language Processing","venue":null,"work_id":"fd6fc6d3-d631-4d28-afb2-650dd9ef17b0","year":2023},"citing_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T23:41:25.868031Z"},"links":{"citing_paper":"/paper/2502.08796"},"observation_digest":"sha256:1d315cd7484db0a26adbaeed0d81ef415c93137f578735e473148eeaf1c6cbe8","observation_id":"88b625b0-be49-42b9-ae4d-e2ac0208b656","resolution":{"observed_at":"2026-08-07T23:41:25.983246Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T23:41:25.961794Z","title":"john thinks that mary thinks that","venue":null,"work_id":"45d2c655-8fea-41c6-b0c0-3594e34048ab","year":2006},"citing_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T23:41:25.872585Z"},"links":{"citing_paper":"/paper/2502.08796"},"observation_digest":"sha256:608fbdfc0c6af32f36cc6e329793985ed05c86d28d3b2f247f1f600ab9495953","observation_id":"a654277e-92ad-4a7d-91fe-b5aea22fbc2b","resolution":{"observed_at":"2026-08-07T23:41:25.967583Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.10076","last_updated":"2023-10-16T05:19:02Z","snapshot_observed_at":"2026-08-05T00:07:01.149187Z","submitted_at":"2023-10-16T05:19:02Z","title":"Verbosity Bias in Preference Labeling by Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.10076","snapshot_observed_at":"2026-08-07T23:41:25.877634Z","title":"I’m sorry to hear that","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T23:41:25.877634Z"},"links":{"cited_paper":"/paper/2310.10076","citing_paper":"/paper/2502.08796"},"observation_digest":"sha256:14b25cf9a78644b1aa56c29e0b6e9721a96072d0eccad4e2111969cdb177d764","observation_id":"28d1d3ea-edd4-448e-b1a2-29d55363add4","resolution":{"observed_at":"2026-08-07T23:41:25.877634Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T23:41:25.945311Z","title":null,"venue":null,"work_id":"53c97517-7fb4-4283-88b5-6c3f33073315","year":2021},"citing_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T23:41:25.883309Z"},"links":{"citing_paper":"/paper/2502.08796"},"observation_digest":"sha256:1b491d17c9417e5f8c0363a81e46d66755dc12c5bd301622dfd7e73b66204117","observation_id":"b5da36a1-8bcf-4b9b-8e64-88d99e5bbb27","resolution":{"observed_at":"2026-08-07T23:41:25.950863Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T23:41:26.107662Z","title":"theory of mind","venue":null,"work_id":"f78c7a3a-02cd-4bae-9513-87a6ca1bfaa6","year":1999},"citing_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"reference_index":1985,"source":"pdf_text","source_observed_at":"2026-08-07T23:41:25.824224Z"},"links":{"citing_paper":"/paper/2502.08796"},"observation_digest":"sha256:f5a2e900cdf7edf14a73fea623e6e60b55534143edfb105ef38c10ecd8dbc7ff","observation_id":"635658df-82c3-4eb8-8c6f-aabba11bd492","resolution":{"observed_at":"2026-08-07T23:41:26.112518Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T23:41:26.026954Z","title":"Annals of the New York Academy of Sciences, 1167:103–114","venue":null,"work_id":"294ebb75-4a12-45a4-b502-faaa6ae6c065","year":2024},"citing_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"reference_index":2009,"source":"pdf_text","source_observed_at":"2026-08-07T23:41:25.854400Z"},"links":{"citing_paper":"/paper/2502.08796"},"observation_digest":"sha256:4a4228edffe7b166e8e6c4da6f19941721e081ca7d72f51195a7a53d7243f6f2","observation_id":"1de39887-0f2e-45f7-a99a-a8408a21541d","resolution":{"observed_at":"2026-08-07T23:41:26.032633Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T23:41:26.010532Z","title":null,"venue":null,"work_id":"899f72b6-d58b-401a-9b19-5bc4129e5a68","year":2019},"citing_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"reference_index":2019,"source":"pdf_text","source_observed_at":"2026-08-07T23:41:25.858851Z"},"links":{"citing_paper":"/paper/2502.08796"},"observation_digest":"sha256:b241ecb0bf31fef30164342df74aebb06f73a16a9bb309ad33a0097cc9462af5","observation_id":"8c4842f1-6ab3-4250-a4bd-5a747f616e94","resolution":{"observed_at":"2026-08-07T23:41:26.015459Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T23:41:26.076156Z","title":"Beaudoin C., Leblanc É., Gagner C., and Beauchamp M","venue":null,"work_id":"d0b07ce8-4a38-4b34-abfc-7a20a6aaa915","year":2020},"citing_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-07T23:41:25.833946Z"},"links":{"citing_paper":"/paper/2502.08796"},"observation_digest":"sha256:8a79dc40e2aa2eeb2385f87bd3689cc7b1e72d6c9f9dbfd6272f21767f0312cc","observation_id":"5042e66b-99d6-4b14-9f8e-259563d7306c","resolution":{"observed_at":"2026-08-07T23:41:26.081555Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T23:41:26.122533Z","title":"In Proceedings of the 2021 Conference on Empirical Methods in Natural Language Processing, pages 1112–1125","venue":null,"work_id":"7b8a24b4-8ea0-4539-985f-b1cdd1038308","year":2021},"citing_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-07T23:41:25.818762Z"},"links":{"citing_paper":"/paper/2502.08796"},"observation_digest":"sha256:d0597fc4ccbaa1580846e69706257c0374851826b06a6009e04bed10f4ea16e9","observation_id":"86702ea1-7808-4808-a233-abcec26de164","resolution":{"observed_at":"2026-08-07T23:41:26.127387Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T23:41:26.044173Z","title":"Jan-Christoph Klie, Bonnie Webber, and Iryna Gurevych","venue":null,"work_id":"f475675e-97e0-47a8-a4ce-90f74e040e7f","year":2022},"citing_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-07T23:41:25.850013Z"},"links":{"citing_paper":"/paper/2502.08796"},"observation_digest":"sha256:a1a854df23b3a5e7027bcef87288bf616d2a75bfe244956286e0618ad2d0e00c","observation_id":"32b8cf54-dc9d-4f69-a9f4-6afaa8b9fc0b","resolution":{"observed_at":"2026-08-07T23:41:26.049508Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T23:41:26.092266Z","title":"Tolga Bolukbasi, Kai-Wei Chang, James Zou, Venkatesh Saligrama, and Adam Kalai","venue":null,"work_id":"a2e7367e-57a5-4641-9242-7484c2402eae","year":2016},"citing_paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-07T23:41:25.828970Z"},"links":{"citing_paper":"/paper/2502.08796"},"observation_digest":"sha256:105f9d284c3e3408f33cf851dd8632b5fa74c87f6b98c767008eb160f7149fc5","observation_id":"4cf67cfb-9aaa-4181-be55-e6a80067ecc9","resolution":{"observed_at":"2026-08-07T23:41:26.097403Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2502.08796","last_updated":"2025-02-12T21:19:30Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T23:36:07.953495Z","submitted_at":"2025-02-12T21:19:30Z","title":"A Systematic Review on the Evaluation of Large Language Models in Theory of Mind Tasks"},"reference_resolution":{"displayed":14,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":5,"verified_exact":0,"verified_fuzzy":9},"total_outbound_references":14},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 14 of 14 outbound references and 5 inbound Pith citation observations for arXiv:2502.08796."}