{"as_of":"2026-08-07T16:25:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:7d20ffae14234cf4d94ffcd93f30d4fe14bb017f21431c402cd0510ac6822b37","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":6,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":6,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":6,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":6,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T21:52:24.557582Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-01T09:45:39.836985Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2401.09712","last_updated":"2024-01-18T04:10:20Z","snapshot_observed_at":"2026-07-06T17:17:11.535092Z","submitted_at":"2024-01-18T04:10:20Z","title":"SkyEyeGPT: Unifying Remote Sensing Vision-Language Tasks via Instruction Tuning with Large Language Model","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.09712","snapshot_observed_at":"2026-08-06T21:52:24.557582Z","title":"Skyeyegpt: Unifying remote sensing vision-language tasks via instruc- tion tuning with large language model","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.23219","last_updated":"2025-06-29T13:04:27Z","snapshot_observed_at":"2026-08-06T21:45:40.384973Z","submitted_at":"2025-06-29T13:04:27Z","title":"UrbanLLaVA: A Multi-modal Large Language Model for Urban Intelligence with Spatial Reasoning and Understanding","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-06T21:52:24.557582Z"},"links":{"cited_paper":"/paper/2401.09712","citing_paper":"/paper/2506.23219"},"observation_digest":"sha256:aa5c9c1275a1e2e057ce98f5bd309af5fd8fa13dbe371db38f2b3ae70eaff0c8","observation_id":"4ea88e8a-8db0-4ce1-8d73-77ef38fac776","resolution":{"observed_at":"2026-08-06T21:52:24.557582Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.09712","last_updated":"2024-01-18T04:10:20Z","snapshot_observed_at":"2026-07-06T17:17:11.535092Z","submitted_at":"2024-01-18T04:10:20Z","title":"SkyEyeGPT: Unifying Remote Sensing Vision-Language Tasks via Instruction Tuning with Large Language Model","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.09712","snapshot_observed_at":"2026-08-06T19:46:12.447693Z","title":"Skyeyegpt: Unifying remote sensing vision-language tasks via instruction tuning with large language model","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.04664","last_updated":"2025-07-07T05:10:15Z","snapshot_observed_at":"2026-08-07T12:14:11.144585Z","submitted_at":"2025-07-07T05:10:15Z","title":"VectorLLM: Human-like Extraction of Structured Building Contours vis Multimodal LLMs","version":1},"reference_index":89,"source":"arxiv_source","source_observed_at":"2026-08-06T19:46:12.447693Z"},"links":{"cited_paper":"/paper/2401.09712","citing_paper":"/paper/2507.04664"},"observation_digest":"sha256:195bde2e2ff924b38ed4b672fc0caa7458ffbf4c491b7b33afc9adff05036f62","observation_id":"22a3060a-2a1b-45da-a5e6-2f7ab87ee0dc","resolution":{"observed_at":"2026-08-06T19:46:12.447693Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.09712","last_updated":"2024-01-18T04:10:20Z","snapshot_observed_at":"2026-07-06T17:17:11.535092Z","submitted_at":"2024-01-18T04:10:20Z","title":"SkyEyeGPT: Unifying Remote Sensing Vision-Language Tasks via Instruction Tuning with Large Language Model","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.09712","snapshot_observed_at":"2026-08-06T19:24:41.352311Z","title":"SkyEyeGPT: Unifying Remote Sensing Vision-Language Tasks via Instruction Tuning with Large Language Model","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.05731","last_updated":"2025-07-08T07:24:34Z","snapshot_observed_at":"2026-08-06T19:16:43.622196Z","submitted_at":"2025-07-08T07:24:34Z","title":"A Satellite-Ground Synergistic Large Vision-Language Model System for Earth Observation","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T19:24:41.352311Z"},"links":{"cited_paper":"/paper/2401.09712","citing_paper":"/paper/2507.05731"},"observation_digest":"sha256:85412e4d03507c9ea9def348c10da95e89a68a7d3634aa12833d4ee17c122901","observation_id":"5d82597f-9204-4330-9b61-f341af8fb186","resolution":{"observed_at":"2026-08-06T19:24:41.352311Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.09712","last_updated":"2024-01-18T04:10:20Z","snapshot_observed_at":"2026-07-06T17:17:11.535092Z","submitted_at":"2024-01-18T04:10:20Z","title":"SkyEyeGPT: Unifying Remote Sensing Vision-Language Tasks via Instruction Tuning with Large Language Model","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.09712","snapshot_observed_at":"2026-08-06T15:07:32.957953Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.16716","last_updated":"2025-07-22T15:54:53Z","snapshot_observed_at":"2026-08-07T01:24:30.524241Z","submitted_at":"2025-07-22T15:54:53Z","title":"Enhancing Remote Sensing Vision-Language Models Through MLLM and LLM-Based High-Quality Image-Text Dataset Generation","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T15:07:32.957953Z"},"links":{"cited_paper":"/paper/2401.09712","citing_paper":"/paper/2507.16716"},"observation_digest":"sha256:cf9f2109534ffbe4889748fa525ee6c7e5c36e17fe86097ddaf4931929e7237e","observation_id":"8c20c456-09d7-4605-89a6-7cf0ff664e8d","resolution":{"observed_at":"2026-08-06T15:07:32.957953Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.09712","last_updated":"2024-01-18T04:10:20Z","snapshot_observed_at":"2026-07-06T17:17:11.535092Z","submitted_at":"2024-01-18T04:10:20Z","title":"SkyEyeGPT: Unifying Remote Sensing Vision-Language Tasks via Instruction Tuning with Large Language Model","version":1},"cited_work":{"arxiv_id":"2401.09712","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2401.09712","snapshot_observed_at":"2026-07-01T09:45:39.836985Z","title":"Skyeyegpt: Unifying remote sensing vision- language tasks via instruction tuning with large language model.ArXiv, abs/2401.09712","venue":null,"work_id":"09a1ae83-890e-4dfe-96b2-307440cef33d","year":2024},"citing_paper":{"arxiv_id":"2605.14475","last_updated":"2026-05-14T07:15:46Z","snapshot_observed_at":"2026-08-02T19:32:52.618513Z","submitted_at":"2026-05-14T07:15:46Z","title":"GeoVista: Visually Grounded Active Perception for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-15T02:39:56.666424Z"},"links":{"cited_paper":"/paper/2401.09712","citing_paper":"/paper/2605.14475"},"observation_digest":"sha256:0f14c6978b081d3a45e1bc5bdeb059ef5aae6f99312c0a7d3d23945b1d9444ea","observation_id":"8656d3e7-bdd5-44bb-9949-15e9c0278521","resolution":{"observed_at":"2026-05-15T02:53:34.044014Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.09712","last_updated":"2024-01-18T04:10:20Z","snapshot_observed_at":"2026-07-06T17:17:11.535092Z","submitted_at":"2024-01-18T04:10:20Z","title":"SkyEyeGPT: Unifying Remote Sensing Vision-Language Tasks via Instruction Tuning with Large Language Model","version":1},"cited_work":{"arxiv_id":"2401.09712","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2401.09712","snapshot_observed_at":"2026-07-01T09:45:39.836985Z","title":"Skyeyegpt: Unifying remote sensing vision- language tasks via instruction tuning with large language model.ArXiv, abs/2401.09712","venue":null,"work_id":"09a1ae83-890e-4dfe-96b2-307440cef33d","year":2024},"citing_paper":{"arxiv_id":"2606.31467","last_updated":"2026-06-30T10:46:23Z","snapshot_observed_at":"2026-07-07T00:05:12.546815Z","submitted_at":"2026-06-30T10:46:23Z","title":"AeroVerse-SatAgent: UAV-Satellite Collaborative Spatial Reasoning Inspired by the Dual Visual Pathway Theory of Cognitive Neuroscience","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-07-01T06:18:35.494240Z"},"links":{"cited_paper":"/paper/2401.09712","citing_paper":"/paper/2606.31467"},"observation_digest":"sha256:3c8eafa979eae47248ae1f3d3e4452d5062a34105e81b0fa635aa8def8c99e95","observation_id":"27edb852-7945-4513-9d5b-8c33a2f0b282","resolution":{"observed_at":"2026-07-01T09:45:39.838498Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2401.09712/citation-record","integrity":"/paper/2401.09712/integrity","json":"/paper/2401.09712/citation-record.json","paper":"/paper/2401.09712"},"outbound":[],"paper":{"arxiv_id":"2401.09712","last_updated":"2024-01-18T04:10:20Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T17:17:11.535092Z","submitted_at":"2024-01-18T04:10:20Z","title":"SkyEyeGPT: Unifying Remote Sensing Vision-Language Tasks via Instruction Tuning with Large Language Model"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 6 inbound Pith citation observations for arXiv:2401.09712."}