{"as_of":"2026-08-07T23:34:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:d5f6dff71fcc1ef29cccb408e674687c810cd30152fd01ca648775d46804c7d1","coverage":[{"denominator":45,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":45,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T05:09:24.020605Z","state":"measured"},{"denominator":46,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":46,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-27T19:36:57.231932Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T21:37:25.289403Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"cited_work":{"arxiv_id":"2506.08691","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.08691","snapshot_observed_at":"2026-07-02T21:37:25.289403Z","title":null,"venue":null,"work_id":"859f9f97-29f9-4d62-ae4a-f69ebf5dcf02","year":2025},"citing_paper":{"arxiv_id":"2606.08231","last_updated":"2026-06-06T15:39:29Z","snapshot_observed_at":"2026-08-07T16:58:05.550639Z","submitted_at":"2026-06-06T15:39:29Z","title":"Test-Time Scaling in Multimodal Foundation Models: A Comprehensive Survey of Generation and Reasoning","version":1},"reference_index":109,"source":"arxiv_source","source_observed_at":"2026-06-27T19:36:57.231932Z"},"links":{"cited_paper":"/paper/2506.08691","citing_paper":"/paper/2606.08231"},"observation_digest":"sha256:03e7621f0102627a49ce856ede9ac176db4a08b661dddeb5dfc350191b2102a4","observation_id":"cc0b75d9-5201-44ca-bc33-ab92b9cfcfb7","resolution":{"observed_at":"2026-07-02T21:37:25.290672Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.08691/citation-record","integrity":"/paper/2506.08691/integrity","json":"/paper/2506.08691/citation-record.json","paper":"/paper/2506.08691"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2405.16473","last_updated":"2024-05-26T07:56:30Z","snapshot_observed_at":"2026-07-06T18:20:05.707144Z","submitted_at":"2024-05-26T07:56:30Z","title":"M$^3$CoT: A Novel Benchmark for Multi-Domain Multi-step Multi-modal Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16473","snapshot_observed_at":"2026-08-07T05:09:18.196208Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:18.196208Z"},"links":{"cited_paper":"/paper/2405.16473","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:8acbfc40951660e1fee2bbce86578b56647734788ab56710983e98688a297e68","observation_id":"23ace000-230b-4d70-a2f6-a22c80d74a04","resolution":{"observed_at":"2026-08-07T05:09:18.196208Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.00855","last_updated":"2024-10-30T14:45:00Z","snapshot_observed_at":"2026-08-07T12:15:46.504578Z","submitted_at":"2024-10-30T14:45:00Z","title":"Vision-Language Models Can Self-Improve Reasoning via Reflection","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.00855","snapshot_observed_at":"2026-08-07T05:09:18.339751Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:18.339751Z"},"links":{"cited_paper":"/paper/2411.00855","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:9183690ac7e6743415ece8426a9a554094bb3fd4d0e81077bc47b6980b0ce69b","observation_id":"7c80bb66-02a8-4420-a3c7-08d9ebd843db","resolution":{"observed_at":"2026-08-07T05:09:18.339751Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:26.778172Z","title":null,"venue":null,"work_id":"475fa8f2-4ba3-4957-b36c-5070a55b6265","year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:18.463000Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:d5d882613ae79f2912efaf5ee0e0f88b58a6579e4cef387d1989325517e865b2","observation_id":"a88b7e8b-8b0b-4109-b718-e29bf56f8044","resolution":{"observed_at":"2026-08-07T05:09:26.869391Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.17179","last_updated":"2024-02-09T00:13:46Z","snapshot_observed_at":"2026-07-06T16:25:25.534843Z","submitted_at":"2023-09-29T12:20:19Z","title":"Alphazero-like Tree-Search can Guide Large Language Model Decoding and Training","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.17179","snapshot_observed_at":"2026-08-07T05:09:18.616592Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:18.616592Z"},"links":{"cited_paper":"/paper/2309.17179","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:f74b134ee8a76554450eca3910948c484ca43f0f94b8a9d79f6938f12c034b32","observation_id":"4234ff1d-a7b8-47fb-b9ec-4092ced0f56f","resolution":{"observed_at":"2026-08-07T05:09:18.616592Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:26.590332Z","title":null,"venue":null,"work_id":"fe5c0430-4980-4956-8043-fb32655c60aa","year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:18.772129Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:d528adf6ad7c4dbeb208c3214493699e9484ac4af2d7cadce77006cb58b8a12f","observation_id":"1ff5ce33-eb00-477a-9bfe-dedaac308c02","resolution":{"observed_at":"2026-08-07T05:09:26.688387Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.05237","last_updated":"2025-06-04T10:07:57Z","snapshot_observed_at":"2026-08-03T04:20:15.211547Z","submitted_at":"2024-12-06T18:14:24Z","title":"MAmmoTH-VL: Eliciting Multimodal Reasoning with Instruction Tuning at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.05237","snapshot_observed_at":"2026-08-07T05:09:18.908261Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:18.908261Z"},"links":{"cited_paper":"/paper/2412.05237","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:3f4d616efdc7b6369798bf3ecfeec97a31e5da4c8ab2d6b74cf3b17835a92eeb","observation_id":"545bdbd5-21ae-42ef-983b-af63ca588c00","resolution":{"observed_at":"2026-08-07T05:09:18.908261Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14992","last_updated":"2023-10-23T07:24:28Z","snapshot_observed_at":"2026-07-06T15:32:25.931739Z","submitted_at":"2023-05-24T10:28:28Z","title":"Reasoning with Language Model is Planning with World Model","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14992","snapshot_observed_at":"2026-08-07T05:09:18.989445Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:18.989445Z"},"links":{"cited_paper":"/paper/2305.14992","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:c8d256126dcb2bb62851e7066ae2b2f2fa770812df1eaa5f8969d5e171905f41","observation_id":"6b82c278-4165-4991-bafd-c1790c45ff77","resolution":{"observed_at":"2026-08-07T05:09:18.989445Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:26.391818Z","title":null,"venue":null,"work_id":"75d4bcd4-ac69-4f20-8d8e-62667db0d859","year":2025},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:19.196937Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:7c7e59628f70206e270c348001d95532eb4282b0862520cad4d93abaef5ace37","observation_id":"d9b9faf8-5635-468d-ad9f-915166704426","resolution":{"observed_at":"2026-08-07T05:09:26.473220Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:19.387591Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:19.387591Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:2f7d2570a9e60334bd264c96a8465b160f8e4d7bb9997423ce2e9932d2696c12","observation_id":"22e7a611-396f-4de5-9f91-6664fcbf3176","resolution":{"observed_at":"2026-08-07T05:09:19.387591Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.11694","last_updated":"2024-12-31T01:38:12Z","snapshot_observed_at":"2026-07-06T19:52:03.121993Z","submitted_at":"2024-11-18T16:15:17Z","title":"Enhancing LLM Reasoning with Reward-guided Tree Search","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.11694","snapshot_observed_at":"2026-08-07T05:09:19.582648Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:19.582648Z"},"links":{"cited_paper":"/paper/2411.11694","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:86921c210e57b50cd33ea41018a6202933f3708785b310c77f9584d3560640c5","observation_id":"a4cc418c-f9fe-4f52-9f47-9608d2518715","resolution":{"observed_at":"2026-08-07T05:09:19.582648Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:19.753354Z","title":null,"venue":null,"work_id":null,"year":2006},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:19.753354Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:bd5dd68a60f2789736d7789c1af933d06dea9f0eb16f81b64cca1088b51abbc4","observation_id":"cf9a2bda-4e86-44c3-87b2-93b8c94e6932","resolution":{"observed_at":"2026-08-07T05:09:19.753354Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:19.869927Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:19.869927Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:0e696a6c6daa3b8d461652837930b7e31f03c3a621caf9a580ff27309c1d864f","observation_id":"e48b2f53-69e0-4b15-b07c-c0d92a0647c5","resolution":{"observed_at":"2026-08-07T05:09:19.869927Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:26.204510Z","title":null,"venue":null,"work_id":"6d692479-a67d-47f1-af9f-b085dd052fc9","year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.016775Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:0f9bd21cbf40b89fe21ee52b4027753fb6ab8f4f1773a3081780e41cfabcc022","observation_id":"0838fab1-c0a1-4f3b-bca0-ad1570aa6a54","resolution":{"observed_at":"2026-08-07T05:09:26.276145Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.20050","last_updated":"2023-05-31T17:24:00Z","snapshot_observed_at":"2026-08-05T13:11:04.104454Z","submitted_at":"2023-05-31T17:24:00Z","title":"Let's Verify Step by Step","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.20050","snapshot_observed_at":"2026-08-07T05:09:20.155515Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.155515Z"},"links":{"cited_paper":"/paper/2305.20050","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:c5339d117d7c316d6e996b34e916ac7ed37cca65e432d0ac4a8a46556cdf6c29","observation_id":"703522a3-f8db-4ca3-bffe-b10263c8ba16","resolution":{"observed_at":"2026-08-07T05:09:20.155515Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.11236","last_updated":"2024-04-25T03:04:14Z","snapshot_observed_at":"2026-07-06T17:45:55.829892Z","submitted_at":"2024-03-17T14:49:09Z","title":"ChartThinker: A Contextual Chain-of-Thought Approach to Optimized Chart Summarization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.11236","snapshot_observed_at":"2026-08-07T05:09:20.285362Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.285362Z"},"links":{"cited_paper":"/paper/2403.11236","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:fa58ccc29df89b544b36be114a0b1f16db164e5e40196993db1e0dd6f4a8c8b3","observation_id":"53f62501-e248-4ba2-b045-4974a4ac66bb","resolution":{"observed_at":"2026-08-07T05:09:20.285362Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.08291","last_updated":"2023-05-15T01:18:23Z","snapshot_observed_at":"2026-07-06T15:27:02.376191Z","submitted_at":"2023-05-15T01:18:23Z","title":"Large Language Model Guided Tree-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.08291","snapshot_observed_at":"2026-08-07T05:09:20.396124Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.396124Z"},"links":{"cited_paper":"/paper/2305.08291","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:84aae893918c25590ad095c6c57310f419d51227b4cbfab7699da4dd1c1e3fcd","observation_id":"3053ebc7-247d-47bd-a3ba-0dca31b5db42","resolution":{"observed_at":"2026-08-07T05:09:20.396124Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.02255","last_updated":"2024-01-21T03:47:06Z","snapshot_observed_at":"2026-07-06T16:27:15.027202Z","submitted_at":"2023-10-03T17:57:24Z","title":"MathVista: Evaluating Mathematical Reasoning of Foundation Models in Visual Contexts","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.02255","snapshot_observed_at":"2026-08-07T05:09:20.523908Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.523908Z"},"links":{"cited_paper":"/paper/2310.02255","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:4557803eaabf9059000903335f7b8ba4f4717cf9e700e53e20bd12336af931ec","observation_id":"2343b7c5-30cf-4720-be51-62f1247f0507","resolution":{"observed_at":"2026-08-07T05:09:20.523908Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:20.636332Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.636332Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:e27a9b19fd3c7b52ba097edafae6ee6bc1dcefe1eb9882c528607f1c80fc63d8","observation_id":"cf928d51-7b97-4148-b96b-7546f86bebab","resolution":{"observed_at":"2026-08-07T05:09:20.636332Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:26.023492Z","title":null,"venue":null,"work_id":"26ff45d1-60f4-4ccf-8af5-9a3244b660c4","year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.747025Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:404872471a22fdeb949a1568cc6e39afe0e6013d0e005ac7a18aacd22b8dca4b","observation_id":"02824006-03e0-44bd-8de7-a77bf2d4f914","resolution":{"observed_at":"2026-08-07T05:09:26.094013Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:25.801322Z","title":null,"venue":null,"work_id":"d8d66345-251c-4550-a1df-20c87aa20c0e","year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.887191Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:4bba8a5a802de97276d4c14c3c38c737417c010c70576da67d44672d4d631cfd","observation_id":"64f3840c-b3e4-4392-9489-7bf1e8cf09d0","resolution":{"observed_at":"2026-08-07T05:09:25.899175Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:25.608973Z","title":null,"venue":null,"work_id":"32a2e83b-5162-432a-bc85-eafaf3b93858","year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:21.057769Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:95254dcd31c67f34b630b33cc2c87e6e83ec035bbd3973abfb95c17693ce3f83","observation_id":"243c15b1-9f15-4dcf-aac3-c2fadbf53751","resolution":{"observed_at":"2026-08-07T05:09:25.683150Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14804","last_updated":"2024-02-22T18:56:38Z","snapshot_observed_at":"2026-08-06T04:41:26.667478Z","submitted_at":"2024-02-22T18:56:38Z","title":"Measuring Multimodal Mathematical Reasoning with MATH-Vision Dataset","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14804","snapshot_observed_at":"2026-08-07T05:09:21.174185Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:21.174185Z"},"links":{"cited_paper":"/paper/2402.14804","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:c8d0e1ac540c6d1ffe8ce3cb17a9bbdbabf79043f625394e0bcfc2840fb4f7e3","observation_id":"1a6999fa-faf2-47bb-a125-bdb61b8312a0","resolution":{"observed_at":"2026-08-07T05:09:21.174185Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12191","last_updated":"2024-10-03T15:54:49Z","snapshot_observed_at":"2026-08-06T05:35:29.109022Z","submitted_at":"2024-09-18T17:59:32Z","title":"Qwen2-VL: Enhancing Vision-Language Model's Perception of the World at Any Resolution","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12191","snapshot_observed_at":"2026-08-07T05:09:21.380191Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:21.380191Z"},"links":{"cited_paper":"/paper/2409.12191","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:0389104d7b06cda05b3d0a9081ea736268869dd7200bc5531347b9eee3c9a3d5","observation_id":"51686c0b-61a3-46c9-975f-c39cbefbb205","resolution":{"observed_at":"2026-08-07T05:09:21.380191Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2203.11171","last_updated":"2023-03-07T17:57:37Z","snapshot_observed_at":"2026-07-06T12:50:22.773056Z","submitted_at":"2022-03-21T17:48:52Z","title":"Self-Consistency Improves Chain of Thought Reasoning in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2203.11171","snapshot_observed_at":"2026-08-07T05:09:21.545744Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:21.545744Z"},"links":{"cited_paper":"/paper/2203.11171","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:586ea4f841e1c27789c0224c70e1a9a0440dd239b25af8da47d3e08100a5a657","observation_id":"d22d9979-53ec-47a4-af77-8cb356a00a67","resolution":{"observed_at":"2026-08-07T05:09:21.545744Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:21.647142Z","title":"Chi, Sharan Narang, Aakanksha Chowdhery, and Denny Zhou","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:21.647142Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:ada1202ca5551b8b6ed1fbcb5670c4a9cb8787c2af0fb810a9481d2dfc888bbd","observation_id":"59fbe118-6147-4a95-9996-0156d9c7c597","resolution":{"observed_at":"2026-08-07T05:09:21.647142Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.18521","last_updated":"2024-06-26T17:50:11Z","snapshot_observed_at":"2026-08-06T23:04:27.981034Z","submitted_at":"2024-06-26T17:50:11Z","title":"CharXiv: Charting Gaps in Realistic Chart Understanding in Multimodal LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.18521","snapshot_observed_at":"2026-08-07T05:09:21.813886Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:21.813886Z"},"links":{"cited_paper":"/paper/2406.18521","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:646565bc9b74a281f87fd3770ad8dd28a0000faf81cd15bccdcb9c7cd767f469","observation_id":"bc132b93-c994-4b01-8c8a-2b858c3fdb7f","resolution":{"observed_at":"2026-08-07T05:09:21.813886Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:22.026151Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.026151Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:45f70120d2b5909c94f2637a5c0b72a898be2707751bc7da7d5fcacd15b95fda","observation_id":"ecb2c3a1-578c-4746-853a-24f3db9e00e9","resolution":{"observed_at":"2026-08-07T05:09:22.026151Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.07001","last_updated":"2024-11-06T13:56:28Z","snapshot_observed_at":"2026-07-06T18:12:59.925736Z","submitted_at":"2024-05-11T12:33:46Z","title":"ChartInsights: Evaluating Multimodal Large Language Models for Low-Level Chart Question Answering","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.07001","snapshot_observed_at":"2026-08-07T05:09:22.134598Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.134598Z"},"links":{"cited_paper":"/paper/2405.07001","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:109fe061faa3e82d64260a508b558dd6b1fb881b411ab827fa2b860c6c420a0a","observation_id":"a92ce3da-0fea-4661-b610-a2ca319d9984","resolution":{"observed_at":"2026-08-07T05:09:22.134598Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.10332","last_updated":"2025-03-21T12:40:26Z","snapshot_observed_at":"2026-07-06T19:51:01.075884Z","submitted_at":"2024-11-15T16:32:34Z","title":"Number it: Temporal Grounding Videos like Flipping Manga","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.10332","snapshot_observed_at":"2026-08-07T05:09:22.277066Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.277066Z"},"links":{"cited_paper":"/paper/2411.10332","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:2cf1f0f6d901299a2a60de572714cce45bcd5fd946b9c03e530a15638b0386cd","observation_id":"7e1acb8e-75d2-4094-aa64-372b6ae565b4","resolution":{"observed_at":"2026-08-07T05:09:22.277066Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:25.426778Z","title":null,"venue":null,"work_id":"07b2dbcd-14d4-4abd-b06f-1f7cf815ef29","year":2025},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.476292Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:f2b897d750b4f03df6ea34c69fdc7821b87b54ade3e1fd3dc37d4cb356209dad","observation_id":"caded182-4d53-4ea1-bd7c-7174b021f941","resolution":{"observed_at":"2026-08-07T05:09:25.508723Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15915","last_updated":"2024-06-19T03:58:32Z","snapshot_observed_at":"2026-07-06T17:08:00.380938Z","submitted_at":"2023-12-26T07:20:55Z","title":"ChartBench: A Benchmark for Complex Visual Reasoning in Charts","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.15915","snapshot_observed_at":"2026-08-07T05:09:22.590714Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.590714Z"},"links":{"cited_paper":"/paper/2312.15915","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:e65d2839e4d01f577b692a3c743a7355cd2edb076383d8d2fd2139727310e051","observation_id":"0ee7b07b-56ff-4d14-bb91-6822defa8d9e","resolution":{"observed_at":"2026-08-07T05:09:22.590714Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10671","last_updated":"2024-09-10T13:25:53Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-07-15T12:35:42Z","title":"Qwen2 Technical Report","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10671","snapshot_observed_at":"2026-08-07T05:09:22.792599Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.792599Z"},"links":{"cited_paper":"/paper/2407.10671","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:09eb4b2857aa87075e89dde19900233f21e206a9ab51dd3d6278ecfd067eb1ad","observation_id":"9fe46db7-7db3-4442-8aa6-ac1809a3e171","resolution":{"observed_at":"2026-08-07T05:09:22.792599Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:25.251886Z","title":null,"venue":null,"work_id":"e5ec472a-be15-4093-891a-a6c26618d8a5","year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.870253Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:452b6f485cef6f8fa87b24e5f51ba7868e599efbfe9848eb53e9f1332ab42968","observation_id":"1aabe9c5-9116-472c-81e7-422112b144eb","resolution":{"observed_at":"2026-08-07T05:09:25.334215Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:22.962916Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.962916Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:2e6697ac0898fad81df826d00791a886d76e3b8e7b3ce6c67563e9138c07046f","observation_id":"2a01d426-545e-4ba5-b37a-309808586dc8","resolution":{"observed_at":"2026-08-07T05:09:22.962916Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:23.042401Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.042401Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:898f5365eea49507d386fddbbc0bc0e16cc6cf4e69474497fcd766daddc3d116","observation_id":"adadffbe-0eba-41be-8c18-fa335bfa63c7","resolution":{"observed_at":"2026-08-07T05:09:23.042401Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:25.058279Z","title":null,"venue":null,"work_id":"61729eae-ad8a-4cd2-adb5-b93928547257","year":2019},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.127849Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:99bb2f02b68ea6ce2e904186ac194c384492f82553c10edce8b8d3da7e53c532","observation_id":"b712d1f3-ae63-4f03-beab-9d97bf352f50","resolution":{"observed_at":"2026-08-07T05:09:25.128428Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03816","last_updated":"2024-11-18T05:36:16Z","snapshot_observed_at":"2026-08-04T20:04:21.125115Z","submitted_at":"2024-06-06T07:40:00Z","title":"ReST-MCTS*: LLM Self-Training via Process Reward Guided Tree Search","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.03816","snapshot_observed_at":"2026-08-07T05:09:23.240182Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.240182Z"},"links":{"cited_paper":"/paper/2406.03816","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:298083fbfa25a44c65ba1b5a1d8cf20d3b468e629686f9313a2323b0b47e55a1","observation_id":"f4a3a9e7-5c69-4fc8-bffa-56f9115de103","resolution":{"observed_at":"2026-08-07T05:09:23.240182Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.02884","last_updated":"2024-11-21T07:07:59Z","snapshot_observed_at":"2026-08-07T12:03:21.367108Z","submitted_at":"2024-10-03T18:12:29Z","title":"LLaMA-Berry: Pairwise Optimization for O1-like Olympiad-Level Mathematical Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.02884","snapshot_observed_at":"2026-08-07T05:09:23.349125Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.349125Z"},"links":{"cited_paper":"/paper/2410.02884","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:44b908d03846efb83eccfe73ab4e17623b3a192579c0b8020f7a3376df4e0039","observation_id":"f87f5801-3e66-45b7-a293-8186d484b3df","resolution":{"observed_at":"2026-08-07T05:09:23.349125Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:24.850784Z","title":null,"venue":null,"work_id":"a9d87ba7-d228-4605-9680-67ec51edd9eb","year":2025},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.429504Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:cbebec183978b0ecbb72adecb660309307f868ad1d654e42439d7c21d4c3c8a6","observation_id":"4c15f74a-18aa-48b8-a829-3e63b4303cae","resolution":{"observed_at":"2026-08-07T05:09:24.940863Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.03493","last_updated":"2022-10-07T12:28:21Z","snapshot_observed_at":"2026-07-06T14:01:50.333970Z","submitted_at":"2022-10-07T12:28:21Z","title":"Automatic Chain of Thought Prompting in Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.03493","snapshot_observed_at":"2026-08-07T05:09:23.513166Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.513166Z"},"links":{"cited_paper":"/paper/2210.03493","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:8627787bfa0374afc07e89f8f5b8490e577ea0b60c4974f062a819d8fa0ac063","observation_id":"ba1fb63a-f1e6-4641-aa47-f9c3ad038dcd","resolution":{"observed_at":"2026-08-07T05:09:23.513166Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.00923","last_updated":"2024-05-20T06:43:48Z","snapshot_observed_at":"2026-07-06T14:47:34.480641Z","submitted_at":"2023-02-02T07:51:19Z","title":"Multimodal Chain-of-Thought Reasoning in Language Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.00923","snapshot_observed_at":"2026-08-07T05:09:23.620651Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.620651Z"},"links":{"cited_paper":"/paper/2302.00923","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:7c1125ef63a772547644369fe8e2edb2a51702dba9b59cba1d979d350fbd930d","observation_id":"133a0312-fd14-48e7-a0bb-71019c467bc8","resolution":{"observed_at":"2026-08-07T05:09:23.620651Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12742","last_updated":"2024-06-18T16:02:18Z","snapshot_observed_at":"2026-07-06T18:33:07.118735Z","submitted_at":"2024-06-18T16:02:18Z","title":"Benchmarking Multi-Image Understanding in Vision and Language Models: Perception, Knowledge, Reasoning, and Multi-Hop Reasoning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12742","snapshot_observed_at":"2026-08-07T05:09:23.728013Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.728013Z"},"links":{"cited_paper":"/paper/2406.12742","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:81cb5ec84fa4639a934d642b4c31d0e5b8b3ea78635f1730b1081611063e4f6c","observation_id":"00d1e608-90c7-4f33-8808-8973a85ceb53","resolution":{"observed_at":"2026-08-07T05:09:23.728013Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:24.665826Z","title":null,"venue":null,"work_id":"567d36f2-332f-4b77-be9b-b8a55811c128","year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.812478Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:41b7d68e7ca5246d18d2331c7a21cd984cd756cd223b96177de82ae3fb981b00","observation_id":"4c852993-e185-4769-a192-ca990e3e4728","resolution":{"observed_at":"2026-08-07T05:09:24.769823Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:23.911059Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.911059Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:54d79a0510965ccddd5f5e8007b4e77ea36982b187a686b1ebac4bdcb7515ba7","observation_id":"b55daa42-dcb9-4598-b6f7-af149e0545d6","resolution":{"observed_at":"2026-08-07T05:09:23.911059Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:24.020605Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:24.020605Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:7fe980ae9208fb84887f6e617ab57706af3a43ea5674026937ada21fc48b4707","observation_id":"d416c5df-3027-4e69-a81a-896a28d1c2ca","resolution":{"observed_at":"2026-08-07T05:09:24.020605Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-07T12:16:31.533543Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism"},"reference_resolution":{"displayed":45,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":45,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":45},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 45 of 45 outbound references and 1 inbound Pith citation observation for arXiv:2506.08691."}