{"as_of":"2026-08-08T10:16:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:1812542da08bde6fbe0f0c5f01ba0ccb08b32be42c8f13f1a657ef9dc7b41628","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":32,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":32,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":32,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":32,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T15:01:35.933783Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-01T22:36:17.139322Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-07T15:01:35.933783Z","title":"Skywork r1v: pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.16673","last_updated":"2025-05-22T13:39:32Z","snapshot_observed_at":"2026-08-07T14:55:04.641472Z","submitted_at":"2025-05-22T13:39:32Z","title":"R1-ShareVL: Incentivizing Reasoning Capability of Multimodal Large Language Models via Share-GRPO","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T15:01:35.933783Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2505.16673"},"observation_digest":"sha256:8ccb05ae68ce285b87ba4a2a0dbdad8e588b14cb71b0ddda0a2570ee360b1939","observation_id":"02d78a85-0a6d-4a8e-921b-925826914b61","resolution":{"observed_at":"2026-08-07T15:01:35.933783Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-07T14:31:14.986136Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.18536","last_updated":"2025-05-24T06:01:48Z","snapshot_observed_at":"2026-08-07T22:01:21.366221Z","submitted_at":"2025-05-24T06:01:48Z","title":"Reinforcement Fine-Tuning Powers Reasoning Capability of Multimodal Large Language Models","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-07T14:31:14.986136Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2505.18536"},"observation_digest":"sha256:e075ea0bc8c085305864d95bce858e507ea8742b0fffb5a88b637ab7e04f8cf4","observation_id":"be1b0c93-c556-486c-8a10-b029c698a2e9","resolution":{"observed_at":"2026-08-07T14:31:14.986136Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-07T14:01:24.672118Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.20236","last_updated":"2025-05-26T17:16:36Z","snapshot_observed_at":"2026-08-07T13:54:23.640348Z","submitted_at":"2025-05-26T17:16:36Z","title":"Seeing is Believing, but How Much? A Comprehensive Analysis of Verbalized Calibration in Vision-Language Models","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T14:01:24.672118Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2505.20236"},"observation_digest":"sha256:0d18debc4aab11463a7f56f24efbf279a0f68f823900e3dcbf62323667007a53","observation_id":"29f1d055-0c40-4292-bd67-dbf5c2773828","resolution":{"observed_at":"2026-08-07T14:01:24.672118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-07T13:12:55.775370Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.22334","last_updated":"2025-07-23T07:37:08Z","snapshot_observed_at":"2026-08-07T21:24:45.521322Z","submitted_at":"2025-05-28T13:21:38Z","title":"Advancing Multimodal Reasoning via Reinforcement Learning with Cold Start","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T13:12:55.775370Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2505.22334"},"observation_digest":"sha256:b53b05d151f2c72cdc3676e12e6befe1ed8417aedee94a75fd96179b4c55995a","observation_id":"90ef9e79-5272-441e-8a41-c03e13bc9072","resolution":{"observed_at":"2026-08-07T13:12:55.775370Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-07T12:01:04.528992Z","title":"Peng, Chris, X","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.00855","last_updated":"2025-06-01T06:28:36Z","snapshot_observed_at":"2026-08-07T11:54:17.799870Z","submitted_at":"2025-06-01T06:28:36Z","title":"MedBookVQA: A Systematic and Comprehensive Medical Benchmark Derived from Open-Access Book","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T12:01:04.528992Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2506.00855"},"observation_digest":"sha256:f3114b5139f87f526b7d97c3a8773dc5e7822120858ee1ff73eaddd98734fef2","observation_id":"18db10f3-bc3d-486f-862c-0c7fa342febb","resolution":{"observed_at":"2026-08-07T12:01:04.528992Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-07T11:55:13.380294Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought.arXiv preprint arXiv:2504.05599, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.01078","last_updated":"2025-06-01T16:28:26Z","snapshot_observed_at":"2026-08-08T00:05:34.569025Z","submitted_at":"2025-06-01T16:28:26Z","title":"GThinker: Towards General Multimodal Reasoning via Cue-Guided Rethinking","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T11:55:13.380294Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2506.01078"},"observation_digest":"sha256:293ba8e838b81447f88dd19c089c1fbb0623232b3d400663f48a62931d73dddf","observation_id":"d2cd426f-ae60-4c5b-a659-a910706b907c","resolution":{"observed_at":"2026-08-07T11:55:13.380294Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-07T11:26:39.566577Z","title":"Skywork r1v: Pioneering multi- modal reasoning with chain-of-thought","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.02555","last_updated":"2025-06-03T07:44:41Z","snapshot_observed_at":"2026-08-07T11:19:24.597547Z","submitted_at":"2025-06-03T07:44:41Z","title":"SurgVLM: A Large Vision-Language Model and Systematic Evaluation Benchmark for Surgical Intelligence","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T11:26:39.566577Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2506.02555"},"observation_digest":"sha256:fcc9e5a4454e091fc37d0a0ffafe13c831545ceae7191783b6faf18f9be0734c","observation_id":"f64cea09-1498-41e3-a433-a442078a277c","resolution":{"observed_at":"2026-08-07T11:26:39.566577Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-07T10:55:14.779760Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought.arXiv preprint arXiv:2504.05599, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.04034","last_updated":"2025-06-04T14:56:57Z","snapshot_observed_at":"2026-08-07T10:46:30.489099Z","submitted_at":"2025-06-04T14:56:57Z","title":"Rex-Thinker: Grounded Object Referring via Chain-of-Thought Reasoning","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T10:55:14.779760Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2506.04034"},"observation_digest":"sha256:64b526cbcbbb9cfbe62026dcaa6e526f35b48e9fdb0fd06dffc8496f888df83c","observation_id":"bd9dee67-9dcc-46ea-bd33-89aa1ee86ff8","resolution":{"observed_at":"2026-08-07T10:55:14.779760Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":"2504.05599","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-07-01T22:36:17.139322Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":"88d1849d-c6fb-43e1-99db-3615fbcc2f50","year":2025},"citing_paper":{"arxiv_id":"2506.09965","last_updated":"2025-06-19T03:46:55Z","snapshot_observed_at":"2026-07-31T21:40:49.363128Z","submitted_at":"2025-06-11T17:41:50Z","title":"Reinforcing Spatial Reasoning in Vision-Language Models with Interwoven Thinking and Visual Drawing","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-05-17T04:58:10.202784Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2506.09965"},"observation_digest":"sha256:e4703ba58d00fa1d334469eb193eb8367cb6c1e18518cfd4b3a9a5b24462ad5d","observation_id":"9794201d-f274-43cd-84bd-ba8fa3830ebb","resolution":{"observed_at":"2026-05-17T04:58:10.293647Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":"2504.05599","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-07-01T22:36:17.139322Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":"88d1849d-c6fb-43e1-99db-3615fbcc2f50","year":2025},"citing_paper":{"arxiv_id":"2506.20332","last_updated":"2026-04-25T13:54:32Z","snapshot_observed_at":"2026-08-07T18:41:40.994733Z","submitted_at":"2025-06-25T11:34:43Z","title":"Mobile-R1: Towards Interactive Capability for VLM-Based Mobile Agent via Systematic Training","version":4},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-19T08:26:12.441869Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2506.20332"},"observation_digest":"sha256:50590a64e26f0a34308aca04e47e9758263a75236674e5b277ed168ca707c3f1","observation_id":"5aa70ea3-58cb-4792-9299-2b9257ae930f","resolution":{"observed_at":"2026-05-19T08:27:11.267022Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-06T22:29:28.285695Z","title":"Skywork r1v: pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.21655","last_updated":"2025-06-26T17:57:08Z","snapshot_observed_at":"2026-08-08T00:05:38.940053Z","submitted_at":"2025-06-26T17:57:08Z","title":"APO: Enhancing Reasoning Ability of MLLMs via Asymmetric Policy Optimization","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T22:29:28.285695Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2506.21655"},"observation_digest":"sha256:42f0a5184f93d9d91f6389406268039dc37382ff0fb2c6580a6be7e460fad72a","observation_id":"aa884deb-08a2-42ee-8cf7-3c4eb220165b","resolution":{"observed_at":"2026-08-06T22:29:28.285695Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-06T22:10:21.421604Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.22434","last_updated":"2025-06-27T17:59:27Z","snapshot_observed_at":"2026-08-06T22:01:48.405215Z","submitted_at":"2025-06-27T17:59:27Z","title":"MiCo: Multi-image Contrast for Reinforcement Visual Reasoning","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:21.421604Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2506.22434"},"observation_digest":"sha256:d3a09d6e72715b5b5dac4701a0af30885f116d2a8e132208de849ea10604fbe0","observation_id":"41ca8a4b-5b82-4ff3-a3b4-f32a305a2bd7","resolution":{"observed_at":"2026-08-06T22:10:21.421604Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":"2504.05599","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-07-01T22:36:17.139322Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":"88d1849d-c6fb-43e1-99db-3615fbcc2f50","year":2025},"citing_paper":{"arxiv_id":"2507.00748","last_updated":"2026-04-12T11:20:16Z","snapshot_observed_at":"2026-08-08T09:01:59.914584Z","submitted_at":"2025-07-01T13:48:57Z","title":"Improving the Reasoning of Multi-Image Grounding in MLLMs via Reinforcement Learning","version":3},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-19T06:50:02.607136Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2507.00748"},"observation_digest":"sha256:73e271017bf3a881538626c16a1adb865b6a712d1ce83552fbfe16a69ea43fa5","observation_id":"ac111eda-a7b8-4b6a-be5d-43d7950494a5","resolution":{"observed_at":"2026-05-19T06:52:08.143618Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-06T20:43:13.212648Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought.arXiv preprint arXiv:2504.05599,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.02076","last_updated":"2025-07-02T18:27:42Z","snapshot_observed_at":"2026-08-08T01:09:25.758170Z","submitted_at":"2025-07-02T18:27:42Z","title":"Reasoning on a Budget: A Survey of Adaptive and Controllable Test-Time Compute in LLMs","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-06T20:43:13.212648Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2507.02076"},"observation_digest":"sha256:a123e11eef16e26efa63ae6d3b03bd940582acf8dd533731e1071cf09d641067","observation_id":"6af19083-4476-41ec-9c25-8ae61921f423","resolution":{"observed_at":"2026-08-06T20:43:13.212648Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":"2504.05599","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-07-01T22:36:17.139322Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":"88d1849d-c6fb-43e1-99db-3615fbcc2f50","year":2025},"citing_paper":{"arxiv_id":"2507.05920","last_updated":"2026-04-20T01:54:49Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-07-08T12:05:05Z","title":"High-Resolution Visual Reasoning via Multi-Turn Grounding-Based Reinforcement Learning","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-19T06:10:57.219445Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2507.05920"},"observation_digest":"sha256:57e05250e3b8e3d63e9e167543b3fc85b64b102f88899a4ebff7cb46a7f48d05","observation_id":"9bfa58a6-b28e-4251-a279-0ff11e7a0dea","resolution":{"observed_at":"2026-05-19T06:12:07.003212Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-06T19:14:04.302415Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.06167","last_updated":"2025-07-10T15:41:04Z","snapshot_observed_at":"2026-08-07T10:37:03.707334Z","submitted_at":"2025-07-08T16:47:16Z","title":"Skywork-R1V3 Technical Report","version":3},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-06T19:14:04.302415Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2507.06167"},"observation_digest":"sha256:4aec35bb66fb1de031f0c35b2e1a6864188b018957c411c58171cc806692e19b","observation_id":"7ab3dd6f-d702-449b-b23f-dc6d1441d93a","resolution":{"observed_at":"2026-08-06T19:14:04.302415Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-06T15:12:01.937219Z","title":"arXiv preprint arXiv:2504.05599 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.16518","last_updated":"2026-06-24T02:00:50Z","snapshot_observed_at":"2026-08-06T15:05:31.009728Z","submitted_at":"2025-07-22T12:27:08Z","title":"SyncLoop: A Multimodal Dual-Loop Framework for Self-Improving Mathematical Reasoning","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T15:12:01.937219Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2507.16518"},"observation_digest":"sha256:1434be6b0232dd0a29750887f346ba9d2390dc85c309dcb33d71fa4dd004f5ac","observation_id":"3c5ca4cf-5fe4-442f-98a7-a4ec04b56acb","resolution":{"observed_at":"2026-08-06T15:12:01.937219Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-06T11:35:15.419272Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.22607","last_updated":"2025-07-31T09:09:45Z","snapshot_observed_at":"2026-08-06T11:35:06.948357Z","submitted_at":"2025-07-30T12:23:21Z","title":"VL-Cogito: Progressive Curriculum Reinforcement Learning for Advanced Multimodal Reasoning","version":2},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-06T11:35:15.419272Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2507.22607"},"observation_digest":"sha256:f104094f7eabded33ae70e7bc7c852368c1f827a6cd13506fc25120b749c558a","observation_id":"2515639a-a809-49df-bc4c-4378da6d87c5","resolution":{"observed_at":"2026-08-06T11:35:15.419272Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-05T20:28:54.771023Z","title":"Peng et al","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.10955","last_updated":"2025-08-14T07:25:45Z","snapshot_observed_at":"2026-08-07T18:23:08.889281Z","submitted_at":"2025-08-14T07:25:45Z","title":"Empowering Multimodal LLMs with External Tools: A Comprehensive Survey","version":1},"reference_index":131,"source":"arxiv_source","source_observed_at":"2026-08-05T20:28:54.771023Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2508.10955"},"observation_digest":"sha256:fb82579fac6c011c401f665bb5732a50e4576e1e04fd1fac0ae6b2d06f96eb42","observation_id":"404118a3-a047-4128-8830-5c5bea76d8d5","resolution":{"observed_at":"2026-08-05T20:28:54.771023Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-05T10:39:08.181121Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.03871","last_updated":"2025-09-04T04:12:31Z","snapshot_observed_at":"2026-08-08T03:54:26.940042Z","submitted_at":"2025-09-04T04:12:31Z","title":"A Comprehensive Survey on Trustworthiness in Reasoning with Large Language Models","version":1},"reference_index":257,"source":"pdf_text","source_observed_at":"2026-08-05T10:39:08.181121Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2509.03871"},"observation_digest":"sha256:8401f03c41db57f65777729534e3508fa40f2248b0dd23c54c178a6c3c18d1dd","observation_id":"027f1f04-692b-457a-ab22-270ac4a7911e","resolution":{"observed_at":"2026-08-05T10:39:08.181121Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-03T21:09:24.459671Z","title":"Skywork r1v: Pioneering multimodal rea- soning with chain-of-thought.ArXiv, abs/2504.05599, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.16672","last_updated":"2026-06-09T22:19:41Z","snapshot_observed_at":"2026-08-06T04:41:40.984877Z","submitted_at":"2025-11-20T18:59:54Z","title":"EvoLMM: Self-Evolving Large Multimodal Models with Continuous Rewards","version":4},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-03T21:09:24.459671Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2511.16672"},"observation_digest":"sha256:f73b21d1ad6ee22deb00ba66f5691f52125062276dfc713d68b1d275ce93d25d","observation_id":"74a6bab7-dd76-488f-9f44-9342d99902e0","resolution":{"observed_at":"2026-08-03T21:09:24.459671Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-03T05:06:46.007695Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.03300","last_updated":"2026-06-15T19:28:32Z","snapshot_observed_at":"2026-08-07T14:53:21.548464Z","submitted_at":"2026-02-03T09:26:32Z","title":"R1-SyntheticVL: Is Synthetic Data from Generative Models Ready for Multimodal Large Language Model?","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-03T05:06:46.007695Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2602.03300"},"observation_digest":"sha256:0bfb73a534850d7662c2e625c47687129c47d3d971c71a3e151b71e9f76621d4","observation_id":"80734e38-b2c9-4dda-acf6-17b720a4d969","resolution":{"observed_at":"2026-08-03T05:06:46.007695Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-08-03T03:21:13.900352Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.08503","last_updated":"2026-06-04T07:08:18Z","snapshot_observed_at":"2026-08-08T03:26:45.601937Z","submitted_at":"2026-02-09T10:55:13Z","title":"Learning Self-Correction in Vision-Language Models via Rollout Augmentation","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-03T03:21:13.900352Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2602.08503"},"observation_digest":"sha256:7c28a8aea5d57468c6d41cc0c6624d6e751441a5ea71020e62dc2bfd640fd88d","observation_id":"c85aff74-3635-4b8a-bb67-dcca7a2768bb","resolution":{"observed_at":"2026-08-03T03:21:13.900352Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-07-13T21:37:55.887477Z","title":"Skywork r1v: Pioneering multimodal rea- soning with chain-of-thought.ArXiv, abs/2504.05599, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.20190","last_updated":"2026-06-09T22:16:34Z","snapshot_observed_at":"2026-08-04T20:13:17.293244Z","submitted_at":"2026-03-20T17:59:25Z","title":"CoVR-R:Reason-Aware Composed Video Retrieval","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-13T21:37:55.887477Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2603.20190"},"observation_digest":"sha256:ab692b11e51087843aa7033ba8cace78be6d1b3cb3fe5323a0b9a0903b2e5459","observation_id":"659e8145-4fd3-45e0-a805-9c0d857b45e5","resolution":{"observed_at":"2026-07-13T21:37:55.887477Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-07-13T19:34:58.789459Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought.arXiv preprint arXiv:2504.05599, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.23483","last_updated":"2026-07-06T05:32:03Z","snapshot_observed_at":"2026-08-07T04:45:34.480037Z","submitted_at":"2026-03-24T17:45:47Z","title":"SpecEyes: Accelerating Agentic Multimodal LLMs via Speculative Perception and Planning","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-07-13T19:34:58.789459Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2603.23483"},"observation_digest":"sha256:ba75685a033629ca3090b10a5057d498b094803c6f7d757c001ac7109ed27824","observation_id":"68baa483-9c90-4e76-8b4e-2da6fe045052","resolution":{"observed_at":"2026-07-13T19:34:58.789459Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":"2504.05599","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-07-01T22:36:17.139322Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":"88d1849d-c6fb-43e1-99db-3615fbcc2f50","year":2025},"citing_paper":{"arxiv_id":"2604.03318","last_updated":"2026-05-25T15:52:36Z","snapshot_observed_at":"2026-07-13T14:39:49.573224Z","submitted_at":"2026-04-01T15:28:13Z","title":"EgoMind: Activating Spatial Cognition through Linguistic Reasoning in MLLMs","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-13T22:41:09.840792Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2604.03318"},"observation_digest":"sha256:4c37be9a2401d17fa03134a5c08e39e38f2b1a9a0fc0489bc4634f15acd4c36f","observation_id":"52c0c1c7-5226-4120-b4c9-a92a0a738e31","resolution":{"observed_at":"2026-05-13T22:43:22.901724Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-07-13T14:39:55.177552Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.03318","last_updated":"2026-05-25T15:52:36Z","snapshot_observed_at":"2026-07-13T14:39:49.573224Z","submitted_at":"2026-04-01T15:28:13Z","title":"EgoMind: Activating Spatial Cognition through Linguistic Reasoning in MLLMs","version":2},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-07-13T14:39:55.177552Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2604.03318"},"observation_digest":"sha256:3df486cc28c302225c6777c2c3562928ac846563ef558b36f66645782cf89c59","observation_id":"1b1cea7c-eece-4c31-a38b-ad1a103f8ef6","resolution":{"observed_at":"2026-07-13T14:39:55.177552Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-07-12T18:39:55.906516Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2604.21405","last_updated":"2026-07-08T08:18:54Z","snapshot_observed_at":"2026-07-12T18:39:55.261528Z","submitted_at":"2026-04-23T08:20:42Z","title":"Supermassive Black Hole Winds in X-rays: SUBWAYS IV. Tracing Radio Emission and Unveiling the Role of Winds","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-07-12T18:39:55.906516Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2604.21405"},"observation_digest":"sha256:b35a1a5293d952dc128830ad0aad7feedfa7a50ba9c6cb67e7c6c2524f9b151e","observation_id":"e4e9800c-e5b8-4a66-8373-ea859d24cc4e","resolution":{"observed_at":"2026-07-12T18:39:55.906516Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":"2504.05599","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-07-01T22:36:17.139322Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":"88d1849d-c6fb-43e1-99db-3615fbcc2f50","year":2025},"citing_paper":{"arxiv_id":"2604.21409","last_updated":"2026-04-23T08:23:25Z","snapshot_observed_at":"2026-08-03T14:07:48.355881Z","submitted_at":"2026-04-23T08:23:25Z","title":"S1-VL: Scientific Multimodal Reasoning Model with Thinking-with-Images","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-09T22:26:03.080281Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2604.21409"},"observation_digest":"sha256:ff139abf1936c3dc59774d05121949dbd100298258834a4345d9b8d5680d9749","observation_id":"46f7997f-6e51-444c-b435-9973d1f256cc","resolution":{"observed_at":"2026-05-09T22:34:07.596827Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":"2504.05599","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-07-01T22:36:17.139322Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":"88d1849d-c6fb-43e1-99db-3615fbcc2f50","year":2025},"citing_paper":{"arxiv_id":"2605.11723","last_updated":"2026-05-28T12:50:43Z","snapshot_observed_at":"2026-08-01T23:43:09.636688Z","submitted_at":"2026-05-12T08:08:33Z","title":"CaC: Advancing Video Reward Models via Hierarchical Spatiotemporal Concentrating","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-13T06:00:31.582714Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2605.11723"},"observation_digest":"sha256:d7e5b68b5acbf8734b1037623f3b65b58d11edb92bacc164beddaae5663ebfe8","observation_id":"a587b752-de3f-428e-968e-6695b47820dd","resolution":{"observed_at":"2026-05-13T06:02:22.039634Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":"2504.05599","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-07-01T22:36:17.139322Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":"88d1849d-c6fb-43e1-99db-3615fbcc2f50","year":2025},"citing_paper":{"arxiv_id":"2605.11723","last_updated":"2026-05-28T12:50:43Z","snapshot_observed_at":"2026-08-01T23:43:09.636688Z","submitted_at":"2026-05-12T08:08:33Z","title":"CaC: Advancing Video Reward Models via Hierarchical Spatiotemporal Concentrating","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-06-30T22:38:43.102769Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2605.11723"},"observation_digest":"sha256:bb3154a1892061590002f2631fcc17f84cfb0ef6c30bd2728cb9699f65073c61","observation_id":"501ae986-43e8-42bb-922d-4bd8e49b3910","resolution":{"observed_at":"2026-07-01T13:55:45.272929Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought","version":2},"cited_work":{"arxiv_id":"2504.05599","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.05599","snapshot_observed_at":"2026-07-01T22:36:17.139322Z","title":"Skywork r1v: Pioneering multimodal reasoning with chain-of-thought","venue":null,"work_id":"88d1849d-c6fb-43e1-99db-3615fbcc2f50","year":2025},"citing_paper":{"arxiv_id":"2606.02459","last_updated":"2026-06-01T16:30:56Z","snapshot_observed_at":"2026-08-07T18:25:54.197937Z","submitted_at":"2026-06-01T16:30:56Z","title":"Active Exploring like a Pigeon: Reinforcing Spatial Reasoning via Agentic Vision-Language Models","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-06-28T15:16:39.121420Z"},"links":{"cited_paper":"/paper/2504.05599","citing_paper":"/paper/2606.02459"},"observation_digest":"sha256:e069ea3a2396a3dacaa3f752824360663025a45dda9ea6bd2b11bfa0124b833a","observation_id":"2ebd7c6b-8f4e-4dcd-ad81-ae4b2d8cd9b8","resolution":{"observed_at":"2026-07-01T22:36:17.142808Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2504.05599/citation-record","integrity":"/paper/2504.05599/integrity","json":"/paper/2504.05599/citation-record.json","paper":"/paper/2504.05599"},"outbound":[],"paper":{"arxiv_id":"2504.05599","last_updated":"2025-06-09T11:44:18Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-07T16:07:44.263463Z","submitted_at":"2025-04-08T01:19:20Z","title":"Skywork R1V: Pioneering Multimodal Reasoning with Chain-of-Thought"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 32 inbound Pith citation observations for arXiv:2504.05599."}