{"as_of":"2026-08-14T01:03:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:605c922eb962412e18c32bb82189c587d348f9a1a1269cf6db25cf5f492c305e","coverage":[{"denominator":62,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":62,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-10T15:29:49.681399Z","state":"measured"},{"denominator":63,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":63,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-29T08:13:42.526597Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-06-29T08:23:15.836108Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"cited_work":{"arxiv_id":"2604.11122","doi":null,"metadata_source":"pith","pith_arxiv_id":"2604.11122","snapshot_observed_at":"2026-06-29T08:23:15.836108Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","venue":"cs.CV","work_id":"bf0a6823-f3cd-496d-bac6-4263c6396e16","year":2026},"citing_paper":{"arxiv_id":"2605.30126","last_updated":"2026-05-28T15:57:31Z","snapshot_observed_at":"2026-08-03T03:31:55.994242Z","submitted_at":"2026-05-28T15:57:31Z","title":"PARCEL: Pool-Anchored Resampling with Conditioned Elastic Queries for Efficient Vision-Language Understanding","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-06-29T08:13:42.526597Z"},"links":{"cited_paper":"/paper/2604.11122","citing_paper":"/paper/2605.30126"},"observation_digest":"sha256:41e3414499b8a069d7c55d7df3e7bd0236a9b749bd6fc81c084845203a159ed9","observation_id":"3507f4b2-4b02-42cd-bdef-1db20f1d8cb1","resolution":{"observed_at":"2026-06-29T08:23:15.837500Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2604.11122/citation-record","integrity":"/paper/2604.11122/integrity","json":"/paper/2604.11122/citation-record.json","paper":"/paper/2604.11122"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":"2303.08774","doi":"10.1002/tea.20265","metadata_source":"pith","pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"GPT-4 Technical Report","venue":"cs.CL","work_id":"b928e041-6991-4c08-8c81-0359e4097c7b","year":2023},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:3470ef87d0cce94d0896a1d7c6536a6b204ad5b923eaceee07083f6bdfe91053","observation_id":"682a15fa-604c-4f73-81ec-3fdb3bcff763","resolution":{"observed_at":"2026-05-11T10:26:01.553947Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"NeurIPS35, 23716–23736 (2022) 2, 4","venue":null,"work_id":"e23917a3-b106-4d93-8813-9678dd00a430","year":2022},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:f0df263751f9b885399218b584f10551f4a95b186118badf3ff6c318a8655abd","observation_id":"f7d337c7-beb3-4197-a843-161e1d963f59","resolution":{"observed_at":"2026-05-17T23:15:27.587168Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"2eb6a4bf-5ab6-4743-8a75-5513b9fad31b","year":2023},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:e087edc553defdb503efb63b6170316ad88eeaf380353aa87adb676b183f2f4c","observation_id":"2350c999-240f-47d3-9385-4b99d93acced","resolution":{"observed_at":"2026-05-17T23:15:27.580918Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.12966","last_updated":"2023-10-13T02:41:28Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-08-24T17:59:17Z","title":"Qwen-VL: A Versatile Vision-Language Model for Understanding, Localization, Text Reading, and Beyond","version":3},"cited_work":{"arxiv_id":"2308.12966","doi":"10.48550/arxiv.2308.12966","metadata_source":"pith","pith_arxiv_id":"2308.12966","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Qwen-VL: A Versatile Vision-Language Model for Understanding, Localization, Text Reading, and Beyond","venue":"cs.CV","work_id":"cbc2bb21-b6bb-46c0-80bf-107e195ffe10","year":2023},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2308.12966","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:f1497965edcf8ece9b1fa6ae022a5edf11aa5e28a6b2ae0f180e769034b8cd35","observation_id":"0b0505fb-8b2a-4e20-860a-9cc8a0cbf832","resolution":{"observed_at":"2026-05-11T10:26:01.525600Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2508.15763","doi":"10.48550/arxiv.2508.15763","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Intern-s1: A scientific multimodal foundation model","venue":"arXiv (Cornell University)","work_id":"9aca1953-129b-4f29-9c6f-f3f066dcb5c3","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:595268ff5ffec736fcba6b8a8ca10245824c64787e1ad8744df22995bf2badf7","observation_id":"2e2b0023-fce9-44c7-b175-b10b29e2d4d0","resolution":{"observed_at":"2026-05-11T10:26:01.509011Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.21631","last_updated":"2025-11-27T12:16:54Z","snapshot_observed_at":"2026-08-11T14:42:37.584016Z","submitted_at":"2025-11-26T17:59:08Z","title":"Qwen3-VL Technical Report","version":2},"cited_work":{"arxiv_id":"2511.21631","doi":"10.1016/j.neunet.2025.107777","metadata_source":"pith","pith_arxiv_id":"2511.21631","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Qwen3-VL Technical Report","venue":"cs.CV","work_id":"1fe243aa-e3c0-4da6-b391-4cbcfc88d5c0","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2511.21631","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:00c864a1aa546eeefe83795ad28277adbf3709ea33672db45440c1e103b8ca49","observation_id":"5dbace75-1ed2-49d0-8386-c426739b2c14","resolution":{"observed_at":"2026-05-11T10:26:01.534709Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.13923","last_updated":"2025-02-19T18:00:14Z","snapshot_observed_at":"2026-08-12T17:29:41.806995Z","submitted_at":"2025-02-19T18:00:14Z","title":"Qwen2.5-VL Technical Report","version":1},"cited_work":{"arxiv_id":"2502.13923","doi":"10.48550/arxiv.2502.13923","metadata_source":"pith","pith_arxiv_id":"2502.13923","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Qwen2.5-VL Technical Report","venue":"cs.CV","work_id":"69dffacb-bfe8-442d-be86-48624c60426f","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2502.13923","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:f42b68d05fe7293945763ad54e29454c62a54ab3bb83a19eaebe7b5a5b86b917","observation_id":"9220fe71-9dc3-4d70-9a8c-da3b089bf531","resolution":{"observed_at":"2026-05-11T10:26:01.490076Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-08T16:08:16.864468+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T16:08:16.864468+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.09461","last_updated":"2023-03-01T19:45:11Z","snapshot_observed_at":"2026-08-13T04:00:22.647615Z","submitted_at":"2022-10-17T22:23:40Z","title":"Token Merging: Your ViT But Faster","version":3},"cited_work":{"arxiv_id":"2210.09461","doi":"10.48550/arxiv.2210.09461","metadata_source":"pith","pith_arxiv_id":"2210.09461","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Token Merging: Your ViT But Faster","venue":"cs.CV","work_id":"528509bc-2611-4e7f-a772-ea14d25b6dae","year":2022},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2210.09461","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:564baad0a11848c4063459a7ac4f40a9ea0595e1ead2f07781fb8a05723f6318","observation_id":"5758ac7b-d82d-4573-9bba-614583f0ee6d","resolution":{"observed_at":"2026-05-12T20:52:10.541797Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Machine Learning111(9), 3125–3160 (2022) 2","venue":null,"work_id":"16df5cbc-f0dc-4a51-9a65-cdf2b7f03856","year":2022},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:f02cf804835a7db2772e2cebed20154dfb0065cde7abc426a92215bade52fd5e","observation_id":"d0d01b0f-1c0e-4967-9196-aed86201c6cc","resolution":{"observed_at":"2026-05-17T23:15:27.594129Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"In: ECCV","venue":null,"work_id":"52809788-022f-48fa-bffe-38114ea8c692","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:0d17a26b8eaac3f974199929615107668df7c7eb7fe7a9c9900a74fb6539dbdf","observation_id":"093c6ddb-a74d-4e0d-9d22-10f0a07d8477","resolution":{"observed_at":"2026-05-17T23:15:27.597603Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Science China Information Sciences67(12), 220101 (2024) 2","venue":null,"work_id":"e19b652b-f7d5-449a-bc7f-ab134100dfe4","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:9ad0ccd084a0bab3f4045fe31b90450fa5767ed5e2f3e6781107006c91e0afda","observation_id":"52324a58-4b0d-4655-a4eb-eb21986d803d","resolution":{"observed_at":"2026-05-17T23:15:27.590739Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.06261","last_updated":"2025-12-19T14:25:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-07-07T17:36:04Z","title":"Gemini 2.5: Pushing the Frontier with Advanced Reasoning, Multimodality, Long Context, and Next Generation Agentic Capabilities","version":6},"cited_work":{"arxiv_id":"2507.06261","doi":"10.48550/arxiv.2503.19","metadata_source":"pith","pith_arxiv_id":"2507.06261","snapshot_observed_at":"2026-07-11T03:17:51.364436Z","title":"Gemini 2.5: Pushing the Frontier with Advanced Reasoning, Multimodality, Long Context, and Next Generation Agentic Capabilities","venue":"cs.CL","work_id":"008df105-2fdd-45d8-857a-8e35868aecb6","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2507.06261","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:74f93dfca2c5c6e70d17e8756bed8be0f3dd9faf9eab8022de820e789f083a8a","observation_id":"f8958755-30d4-4583-9daa-e8827f78b2bc","resolution":{"observed_at":"2026-05-11T10:26:01.494111Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"NeurIPS36, 49250–49267 (2023) 2, 4","venue":null,"work_id":"1c0631d0-5f37-41f7-b6a5-a30b9e8b9340","year":2023},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:3c82ac2535761a639a0ebe54a340e7e673928c3f5fcd857068dcc0f377d060fd","observation_id":"7a04eef4-778b-4349-9638-da94cd060bd2","resolution":{"observed_at":"2026-05-17T23:15:27.570587Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09532","last_updated":"2025-02-01T06:13:27Z","snapshot_observed_at":"2026-08-13T20:27:13.721625Z","submitted_at":"2025-01-16T13:34:33Z","title":"AdaFV: Rethinking of Visual-Language alignment for VLM acceleration","version":2},"cited_work":{"arxiv_id":"2501.09532","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.09532","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"arXiv preprint arXiv:2501.09532 , year=","venue":null,"work_id":"10471b34-4093-495d-b865-4a4fc227e047","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2501.09532","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:07763a1b142936c91ffb0f273ef8f0dd578ded3a76b2a0359ef73b5e19878ef4","observation_id":"c5326a5f-2dd8-4602-80aa-e2cda368c97e","resolution":{"observed_at":"2026-05-11T10:26:01.546179Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"ISPRS Journal of Photogrammetry and Remote Sensing224, 272–286 (2025) 2, 5","venue":null,"work_id":"67f399e1-365b-412c-8386-510ee9154dd9","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:4a5f877c8c9bacf8dfbdc20e6d888f52aadb50568309da31a21ca29a6c8621a3","observation_id":"a143544b-f3f4-48bc-900f-add73db724ee","resolution":{"observed_at":"2026-05-17T23:15:27.577689Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"In: CVPR","venue":null,"work_id":"73cea4ed-ab65-4a01-a14f-9cc5cc1e6c5a","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:dcdabfd66a24528fcb0897a625d121b10d023575861f32ae762337200449bf2e","observation_id":"b92593f7-b2c5-4fa8-a67f-26d6be5219ef","resolution":{"observed_at":"2026-05-17T23:15:27.584124Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.03326","last_updated":"2024-10-26T16:35:13Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-08-06T17:59:44Z","title":"LLaVA-OneVision: Easy Visual Task Transfer","version":3},"cited_work":{"arxiv_id":"2408.03326","doi":"10.48550/arxiv.2408.03326","metadata_source":"pith","pith_arxiv_id":"2408.03326","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"LLaVA-OneVision: Easy Visual Task Transfer","venue":"cs.CV","work_id":"f5f2452b-f2a9-49ac-b38d-c76e18cdfe49","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2408.03326","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:b36fde6538573a5b219c03783a1a27e6ca57db2e5e454cc748848090a97951c9","observation_id":"c425d30c-9402-452c-93be-5495452b6de5","resolution":{"observed_at":"2026-05-11T10:26:01.500026Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"In: ICML","venue":null,"work_id":"d91c4ac5-dbca-4ec3-935c-b0fb739fb3f6","year":2023},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:b06f7312f7db8220f4f6fbdadfbc89a568d4d1cd66411973a02da1efd99f1c8b","observation_id":"e44a0dd0-2f55-474c-86a9-a31cb757dbcf","resolution":{"observed_at":"2026-05-17T23:15:27.574205Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.09301","last_updated":"2024-11-14T09:23:40Z","snapshot_observed_at":"2026-08-13T06:30:46.098333Z","submitted_at":"2024-11-14T09:23:40Z","title":"LHRS-Bot-Nova: Improved Multimodal Large Language Model for Remote Sensing Vision-Language Interpretation","version":1},"cited_work":{"arxiv_id":"2411.09301","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2411.09301","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"arXiv preprint arXiv:2411.09301 (2024),https://arxiv.org/abs/ 2411.093015","venue":null,"work_id":"c9c67148-713f-41e6-be43-8b9c1dabc80c","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2411.09301","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:3e01c357e76420b401b0de3ea0b7adff3d89f1030503e87a835a41844076dc3e","observation_id":"9e417430-9f79-4a7d-be2b-d356eaad3b3f","resolution":{"observed_at":"2026-05-11T10:26:01.316488Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"d944ff84-646a-4164-ba1a-0e3bdaffcc46","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:a1a8d2cdc7ec4c5084c6149fddbde6a0cc9674b2e41ba003a5667823d994bbad","observation_id":"5f5e0553-2571-45b0-9085-acd3c6233703","resolution":{"observed_at":"2026-05-17T23:15:27.657711Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"NeurIPS36, 34892– 34916 (2023) 4","venue":null,"work_id":"610cd46c-735e-4270-894b-633a9672878b","year":2023},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:80f886ef4690ef8580aff5ab354646021a1fb0faf5444bcf28574e48417078b0","observation_id":"ce18f5a6-6150-4ae2-8f1c-29a11ac3ce9d","resolution":{"observed_at":"2026-05-17T23:15:27.644850Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2511.12267","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T05:27:39.989410Z","title":"Zoomearth: Active perception for ultra-high-resolution geospatial vision-language tasks","venue":null,"work_id":"ed2c8189-cc00-4577-95e4-d7c11193ad69","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:77c7d94c3c320b67a8490f7494cedcc2ed0a10700b50e88d66843e194078beb3","observation_id":"f7b080fe-e201-4873-81be-560a3d6a5f73","resolution":{"observed_at":"2026-05-11T10:26:01.359193Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.05679","last_updated":"2024-12-10T02:23:30Z","snapshot_observed_at":"2026-08-12T05:19:12.753205Z","submitted_at":"2024-12-07T15:11:21Z","title":"RSUniVLM: A Unified Vision Language Model for Remote Sensing via Granularity-oriented Mixture of Experts","version":2},"cited_work":{"arxiv_id":"2412.05679","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2412.05679","snapshot_observed_at":"2026-07-03T09:47:59.375569Z","title":"Rsunivlm: A unified vision language model for remote sensing via granularity-oriented mixture of experts","venue":null,"work_id":"c3ccb184-35ee-4d20-aad2-28ac05be3be5","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2412.05679","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:c05c9b9c8e0fdef0e0fc0a263e60d512041a80c43e9eeb1e6e4f943484642855","observation_id":"4ed4e817-79c0-40ea-99e1-b327f90eb442","resolution":{"observed_at":"2026-05-11T10:26:01.278915Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"In: ICCV","venue":null,"work_id":"68a16bed-bb1d-4f73-9987-949702e5fee2","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:4c990511ddb65ba3820929b688b027a28ad42403c3ce79b14cb13fa37d278cb6","observation_id":"579b4017-70b8-44c4-adf2-e93e3bfa17ce","resolution":{"observed_at":"2026-05-17T23:15:27.654633Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Nature Reviews Neuroscience13(11), 758–768 (2012) 4","venue":null,"work_id":"8e55bbc3-3f45-4299-bb91-934385da63a9","year":2012},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:a0e37a4ea45eb0a289ce0b42b1f4961c0e717648bea8c7b52a8ac6553201b819","observation_id":"a4c78456-3dd8-4ea1-b760-0fd80a0b6bfe","resolution":{"observed_at":"2026-05-17T23:15:27.624004Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"In: ECCV","venue":null,"work_id":"8a497a19-2192-4c72-b8b4-97e2ee8117b3","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:a74a61651a265722c4d7caa8712fb660e5d9831ae19f24bfdc6c48fa19e7bde8","observation_id":"855dfb92-54e9-4a0e-8027-1920179603bf","resolution":{"observed_at":"2026-05-17T23:15:27.637797Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"00164680-a8ce-4fe2-b40f-cfb248dbf305","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:14e83502d6997aee35416e605361e089b0040c0af11b84389dab56198c0c73cd","observation_id":"7c80d2e9-e022-4006-9b2a-8a78f9dee358","resolution":{"observed_at":"2026-05-17T23:15:27.627721Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.20213","last_updated":"2024-12-19T13:46:40Z","snapshot_observed_at":"2026-08-13T00:42:14.577442Z","submitted_at":"2024-03-29T14:50:43Z","title":"VHM: Versatile and Honest Vision Language Model for Remote Sensing Image Analysis","version":4},"cited_work":{"arxiv_id":"2403.20213","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.20213","snapshot_observed_at":"2026-07-02T13:56:59.147746Z","title":"arXiv preprint arXiv:2403.20213 (2024),https://arxiv.org/abs/ 2403.202135","venue":null,"work_id":"e6962199-88f1-4272-a567-b86365258880","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2403.20213","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:fa05345266643d31565fda2977056226c447c93d7953a45d4f117e84230a22e9","observation_id":"466528d3-ba0d-4c11-837a-69181a9ef3de","resolution":{"observed_at":"2026-05-11T10:26:01.418886Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"In: ICML","venue":null,"work_id":"8d39eec3-4f33-448e-8912-9b3f309819a6","year":2021},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:1381fd57afc3a7441bf21bf1da67a2f33511aac4cc699a3e5787c52ee908f69d","observation_id":"4ac3132f-d77c-4366-b5d5-9fea238bf795","resolution":{"observed_at":"2026-05-17T23:15:27.634410Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"In: ICCV","venue":null,"work_id":"08a09b1b-c0a1-4b49-a2ea-5091a56f46e1","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:2cad1e3540d022ac5608a8a8d1c7de46a3fd9a80b2ea52523edd2f77df51aa1e","observation_id":"2b18b074-4591-446f-8e20-309ef54666b4","resolution":{"observed_at":"2026-05-17T23:15:27.631168Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2506.01667","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T09:47:59.425017Z","title":"Earthmind: Towards multi-granular and multi- sensor earth observation with large multimodal models","venue":null,"work_id":"ac67b65e-e7c0-4fe5-bf00-c97d38125181","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:8c69d8a60ffcbfac735e898f4e883f3374d8099f4494b605b897fe905c87dd35","observation_id":"a90ae10a-e2ff-44c1-b062-9ccfd38ed5d0","resolution":{"observed_at":"2026-05-11T10:26:01.469106Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.06475","last_updated":"2024-02-09T15:31:01Z","snapshot_observed_at":"2026-08-13T04:22:30.692524Z","submitted_at":"2024-02-09T15:31:01Z","title":"Large Language Models for Captioning and Retrieving Remote Sensing Images","version":1},"cited_work":{"arxiv_id":"2402.06475","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.06475","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"arXiv preprint arXiv:2402.06475 (2024), https://arxiv.org/abs/2402.064755","venue":null,"work_id":"297cda91-3d56-4d6f-8005-0bd5ea58fd57","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2402.06475","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:8cd68048215e0fed741f3e88fceaf98d6d3c2248470ba89528bcad79514248cc","observation_id":"76ec8b87-7db7-47eb-8f3b-52265e924568","resolution":{"observed_at":"2026-05-11T10:26:01.323487Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15190","last_updated":"2025-04-07T06:19:02Z","snapshot_observed_at":"2026-08-12T17:01:58.965430Z","submitted_at":"2024-12-19T18:57:13Z","title":"EarthDial: Turning Multi-sensory Earth Observations to Interactive Dialogues","version":2},"cited_work":{"arxiv_id":"2412.15190","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2412.15190","snapshot_observed_at":"2026-07-03T09:47:59.389420Z","title":"arXiv preprint arXiv:2412.15190 (2025),https://arxiv.org/abs/2412.151905","venue":null,"work_id":"7ea1b501-a82c-4b71-961a-f4e10873679e","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2412.15190","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:ee0a3ac5dee13301012a2e63257ac1e9ec036dcff7b4673cb389a0cac473df6a","observation_id":"5179b229-de60-4840-a540-bb0ac5e28f0a","resolution":{"observed_at":"2026-05-11T10:26:01.450581Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2511.21150","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Llava-uhd v3: Progressive visual compression for efficient native-resolution encoding in mllms.ArXiv, abs/2511.21150","venue":null,"work_id":"b3f1e7f9-9c11-4a9f-9b69-7b94ff5a605d","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:d5e4581e3cb6178d36f4d9673c0c210ceedf51cd8f487a97851903b28b1cdf83","observation_id":"e4996d9c-c359-4ea8-8543-6ff5c326c47a","resolution":{"observed_at":"2026-05-11T10:26:01.308433Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.11805","last_updated":"2025-05-09T21:04:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-19T02:39:27Z","title":"Gemini: A Family of Highly Capable Multimodal Models","version":5},"cited_work":{"arxiv_id":"2312.11805","doi":"10.1038/nrn2888","metadata_source":"pith","pith_arxiv_id":"2312.11805","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Gemini: A Family of Highly Capable Multimodal Models","venue":"cs.CL","work_id":"83f7c85b-3f11-450f-ac0c-64d9745220b2","year":2023},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2312.11805","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:c202b014dcee7f1cb8943d75ab03fabf7132d759f8caa50102267f21204b8c2e","observation_id":"f85673ed-15cd-494d-a5a8-9afc8124278f","resolution":{"observed_at":"2026-05-11T10:26:01.388730Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2302.13971","last_updated":"2023-02-27T17:11:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-02-27T17:11:15Z","title":"LLaMA: Open and Efficient Foundation Language Models","version":1},"cited_work":{"arxiv_id":"2302.13971","doi":"10.48550/arxiv.2302.13971","metadata_source":"pith","pith_arxiv_id":"2302.13971","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"LLaMA: Open and Efficient Foundation Language Models","venue":"cs.CL","work_id":"c018fc23-6f3f-4035-9d02-28a2173b2b9d","year":2023},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2302.13971","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:8bad28ecc977d21e76e328192c1d870fd23df62db157daf79fd43ee20dec3e74","observation_id":"574078ba-5053-4614-a603-05b882195095","resolution":{"observed_at":"2026-05-11T10:26:01.337700Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-08T16:08:17.350515+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T16:08:17.350515+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"NeurIPS30(2017) 4","venue":null,"work_id":"a06fd8c2-54fa-426a-ada8-21b3bc766b58","year":2017},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:fc7a230271058a1b3119d166d82136543e18aa9b1317954ed99668c4d252d577","observation_id":"33bde912-6453-4dee-b764-bcba99f4153d","resolution":{"observed_at":"2026-05-17T23:15:27.619836Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2511.20085","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"arXiv preprint arXiv:2511.20085 (2025) 5","venue":null,"work_id":"d3e55704-f213-4fa7-93d2-7de81121ecae","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:ec593ac212a4c18a582e024fdd0734c7b9b6d8fb228033ad20b7e5fae443ddd8","observation_id":"2a72aeb5-c5e2-4db3-a5fa-df23caad0602","resolution":{"observed_at":"2026-05-11T10:26:01.287593Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2505.21375","doi":"10.48550/arxiv.2505.21375","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Geollava-8k: Scaling remote-sensing multimodal large language models to 8k resolution","venue":"ArXiv.org","work_id":"e8b770f4-ab0f-4fa1-9b6e-f9bd0e4e8960","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:03ba2475a84f0d9ac9a80d3b64dced95d35d3bd75ca409991e8039f6ca62eec6","observation_id":"1ac56322-4638-490b-968f-ae45691415a4","resolution":{"observed_at":"2026-05-11T10:26:01.355789Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"ccdb8033-c8f9-4057-b74b-a089a1613444","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:9001c9dec7892f78ed4af887bf328746affd7816956c8f368c037174e0c6f215","observation_id":"d2818960-d3a9-4f77-9ae0-43b86c2c45e7","resolution":{"observed_at":"2026-05-17T23:15:27.616272Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2601.02783","doi":"10.48550/arxiv.2601.02783","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv:2601.02783 [cs] doi:10.48550/arXiv.2601.02783 Mingze Wang, Lili Su, Cilin Yan, Sheng Xu, Pengcheng Yuan, Xiaolong Jiang, and Baochang Zhang","venue":"arXiv (Cornell University)","work_id":"f042ee2f-5dfd-4825-b3cb-e3e49ea260eb","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:83ff55dcc541c0c722a3167129ea680075c97c2f9054e8f830b803fa35ac3372","observation_id":"fb1786ca-0d28-47b6-b70e-b6bb90e12976","resolution":{"observed_at":"2026-05-11T10:26:01.409730Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2024.351083","doi":"10.1109/tgrs.2024.3510833","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"IEEE Transactions on Geoscience and Remote Sensing63, 1–20 (2025).https://doi.org/10.1109/TGRS.2024.3510833 5","venue":"IEEE Transactions on Geoscience and Remote Sensing","work_id":"11a745ea-fdf8-4a3c-8819-3619e8bed60e","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:6832c873d3e664807df5a50a0dbeb6b06ef0e9523af470d72fac91a01c39e190","observation_id":"b67792b7-a70b-4d95-936d-3f363751707c","resolution":{"observed_at":"2026-05-10T15:30:32.177796Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12191","last_updated":"2024-10-03T15:54:49Z","snapshot_observed_at":"2026-08-06T05:35:29.109022Z","submitted_at":"2024-09-18T17:59:32Z","title":"Qwen2-VL: Enhancing Vision-Language Model's Perception of the World at Any Resolution","version":2},"cited_work":{"arxiv_id":"2409.12191","doi":"10.48550/arxiv.2409.12191","metadata_source":"pith","pith_arxiv_id":"2409.12191","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Qwen2-VL: Enhancing Vision-Language Model's Perception of the World at Any Resolution","venue":"cs.CV","work_id":"8abcfe4f-e0fb-44b7-9123-448fac95f90a","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2409.12191","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:19231788cbf275f28f6e434fdb9824fbb25c28e2c9844fb59d7a501433625285","observation_id":"cabd839d-e94c-4be9-9616-70ef5fa93878","resolution":{"observed_at":"2026-05-11T10:26:01.272349Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-11T02:19:33.884263+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T02:19:33.884263+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.10442","last_updated":"2025-04-07T09:09:39Z","snapshot_observed_at":"2026-08-13T18:24:04.904767Z","submitted_at":"2024-11-15T18:59:27Z","title":"Enhancing the Reasoning Ability of Multimodal Large Language Models via Mixed Preference Optimization","version":2},"cited_work":{"arxiv_id":"2411.10442","doi":"10.48550/arxiv.2411.10442","metadata_source":"pith","pith_arxiv_id":"2411.10442","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Enhancing the Reasoning Ability of Multimodal Large Language Models via Mixed Preference Optimization","venue":"cs.CL","work_id":"52f12dd6-1165-4717-9a1b-04ff19d82dfe","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2411.10442","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:1c7dc66e081df6778925206fdc98a5cbe2c1bdc63c6c3ae5394bde0fea322e30","observation_id":"28ac8422-a138-4282-adc0-2a41a2bbb11c","resolution":{"observed_at":"2026-05-16T09:16:17.933281Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"In: Proceedings of the IEEE conference on computer vision and pattern recognition","venue":null,"work_id":"0b52e438-ab4d-4e23-b71a-89032adf9740","year":2018},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:6bd18c54dbe43b06f7ca61d36740755c313e1697df4540f6713666a666af586b","observation_id":"b220adf3-b9f3-4231-984e-0c5f72fb71e4","resolution":{"observed_at":"2026-05-17T23:15:27.612505Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.17247","last_updated":"2025-02-27T11:16:33Z","snapshot_observed_at":"2026-08-12T14:42:48.682739Z","submitted_at":"2024-10-22T17:59:53Z","title":"PyramidDrop: Accelerating Your Large Vision-Language Models via Pyramid Visual Redundancy Reduction","version":2},"cited_work":{"arxiv_id":"2410.17247","doi":"10.48550/arxiv.2410.17247","metadata_source":"pith","pith_arxiv_id":"2410.17247","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"PyramidDrop: Accelerating Your Large Vision-Language Models via Pyramid Visual Redundancy Reduction","venue":"cs.CV","work_id":"ca5cea36-45c6-4d00-8408-14876a5d59c3","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2410.17247","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:0499d0191efd13259dd5a70a4b0f4cd5b3a50f65ad01e6ec873cb5504ac3f1c4","observation_id":"6a36ca47-58e6-4604-9f55-dfdd194f2112","resolution":{"observed_at":"2026-05-15T12:12:14.835646Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"In: CVPR","venue":null,"work_id":"6c2a643c-378b-49b9-90b3-126b61881edd","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:16515d0a693d68bf305ecb068ee305fabb77ff40566fb90a1b585ecb28b7d897","observation_id":"17c2f794-6f35-4d06-b2f0-ac5c369865f3","resolution":{"observed_at":"2026-05-17T23:15:27.608446Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.14178","last_updated":"2024-03-29T08:13:38Z","snapshot_observed_at":"2026-08-13T08:24:19.015641Z","submitted_at":"2023-04-27T13:27:01Z","title":"mPLUG-Owl: Modularization Empowers Large Language Models with Multimodality","version":3},"cited_work":{"arxiv_id":"2304.14178","doi":"10.48550/arxiv.2304.14178","metadata_source":"pith","pith_arxiv_id":"2304.14178","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"mPLUG-Owl: Modularization Empowers Large Language Models with Multimodality","venue":"cs.CL","work_id":"74a7deb6-48be-4132-9d35-882cc5870ebd","year":2023},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2304.14178","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:7655ea2a54d1d6622651c4fce621678e3c0e697f53570f0a4ac86e392be5f908","observation_id":"7a846a2b-d46c-4c2f-b105-de90b785b8a6","resolution":{"observed_at":"2026-05-11T10:26:01.395330Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"In: AAAI","venue":null,"work_id":"b5d61b7c-654d-492b-a794-79a088dbb5ec","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:69298db5f1f0411d46c06b6620a3badd4cf639295d994d492594bde868f7f0c1","observation_id":"adac2e97-61f4-4cc5-a071-815b016206f0","resolution":{"observed_at":"2026-05-17T23:15:27.604601Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"In: CVPR","venue":null,"work_id":"3b74f0f3-4cd7-44b4-ab03-44f858a175a5","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:ca58b5d7ad05081b57a233cf0c0c9849eecf9619f4caed6154cc2bf9a566a17d","observation_id":"be2c3fbd-d08e-40fc-bb27-c0ba22ad2397","resolution":{"observed_at":"2026-05-17T23:15:27.651528Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"In: CVPR","venue":null,"work_id":"f9990727-5544-4567-92e0-2a6e7c77a3b4","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:ad69e23eb78a964d0fea64a425043a3f1c8720b260aa4601102f2eab26dbed25","observation_id":"4ec0b256-be82-43e1-a3ac-abe17a58e66a","resolution":{"observed_at":"2026-05-17T23:15:27.600904Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2601.22674","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T15:59:57.448270Z","title":"Visiontrim: Unified vision token compression for training-free mllm acceleration","venue":null,"work_id":"acac5768-e16e-423c-bc27-e24c2d6d15aa","year":2026},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:4cfee730ce97c858c38b6aa837e215d131a1bccf48a804f2817e7973aaacf19f","observation_id":"9817257f-fe51-4311-beb9-948b26bc20e9","resolution":{"observed_at":"2026-05-11T10:26:01.365392Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.01548","last_updated":"2025-08-03T02:15:43Z","snapshot_observed_at":"2026-08-08T11:58:42.307183Z","submitted_at":"2025-08-03T02:15:43Z","title":"A Glimpse to Compress: Dynamic Visual Token Pruning for Large Vision-Language Models","version":1},"cited_work":{"arxiv_id":"2508.01548","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2508.01548","snapshot_observed_at":"2026-07-04T15:59:57.463228Z","title":"A glimpse to compress: Dynamic visual token prun- ing for large vision-language models","venue":null,"work_id":"c148d5f2-a53a-49d7-b838-a29b79581d3a","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2508.01548","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:fa8afb831ff1ac797a7df0e673f128f5a6bae06804ca7a7415379400bcc638b7","observation_id":"9a84096e-c99e-4b89-8892-ce982fcf3b50","resolution":{"observed_at":"2026-05-11T10:26:01.485770Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"ISPRS Journal of Pho- togrammetry and Remote Sensing221, 64–77 (2025) 5","venue":null,"work_id":"8398bb52-a8fe-41d3-ab13-e2fce736a1d4","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:e727b283831c6231a547c74cd06a15ecd42937878670441ad7186e011356cf63","observation_id":"e4df18e9-3d41-469c-a30f-3ec266b85c22","resolution":{"observed_at":"2026-05-17T23:15:27.648413Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2505.22654","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T14:09:53.807829Z","title":"arXiv preprint arXiv:2505.22654 , year=","venue":null,"work_id":"722ef313-9de6-4921-85a9-66158d76c41b","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:9b81abe0bf41770dd23def1265c5fc15d46fe4956d16aa22c146259a9645d9b6","observation_id":"74170594-6f54-4c95-b8e0-4a4440d9af1a","resolution":{"observed_at":"2026-05-11T10:26:01.368991Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.01818","last_updated":"2025-05-11T17:45:02Z","snapshot_observed_at":"2026-08-13T05:48:56.985392Z","submitted_at":"2024-12-02T18:57:40Z","title":"Beyond Text-Visual Attention: Exploiting Visual Cues for Effective Token Pruning in VLMs","version":2},"cited_work":{"arxiv_id":"2412.01818","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2412.01818","snapshot_observed_at":"2026-07-02T22:17:26.122002Z","title":"[cls] attention is all you need for training- free visual token pruning: Make vlm inference faster","venue":null,"work_id":"0e705883-8702-42cd-b101-1faff544762a","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2412.01818","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:6302ff4dfa167c0fdd664259ed6858aa08ef29a5060060b451dbd9dfee558947","observation_id":"7dfed331-4316-49bf-b658-c03c40981a1b","resolution":{"observed_at":"2026-05-11T10:26:01.376382Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.13596","last_updated":"2024-11-29T10:18:58Z","snapshot_observed_at":"2026-08-12T23:19:22.625658Z","submitted_at":"2024-07-18T15:35:00Z","title":"EarthMarker: A Visual Prompting Multi-modal Large Language Model for Remote Sensing","version":3},"cited_work":{"arxiv_id":"2407.13596","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2407.13596","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"arXiv preprint arXiv:2407.13596 (2024),https://arxiv.org/abs/2407.135965","venue":null,"work_id":"8593a604-58b4-4c1c-bf0e-daba100b366c","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2407.13596","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:ae7c9b2a901b90b912c84a96f9667a1cef5af27854d435688dfe1793b8cc185c","observation_id":"db20ab85-52b2-44d5-b71b-5f4ee6fe9da9","resolution":{"observed_at":"2026-05-11T10:26:01.425718Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"IEEE Transactions on Geoscience and Remote Sensing62, 1–20 (2024) 5","venue":null,"work_id":"4aa34e8a-7c7d-428c-a38e-978480daaa70","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:b52ad57b1d3ff7072d9559efa4826beea2528a4b03322826dd350e6c0a4232f1","observation_id":"3f7a7faf-9c02-44dd-ae82-efd17624040d","resolution":{"observed_at":"2026-05-17T23:15:27.641514Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.04417","last_updated":"2025-06-03T04:12:10Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-10-06T09:18:04Z","title":"SparseVLM: Visual Token Sparsification for Efficient Vision-Language Model Inference","version":4},"cited_work":{"arxiv_id":"2410.04417","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.04417","snapshot_observed_at":"2026-07-09T21:36:34.338168Z","title":"SparseVLM: Visual Token Sparsification for Efficient Vision-Language Model Inference","venue":"cs.CV","work_id":"691499bc-3ae3-49f0-aba6-ab6c7972a7a7","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2410.04417","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:90599458d2891fdb3b4b90c6cc64d5937cc90ffa6b4a0f8f438b591529c5c8e7","observation_id":"a1fb4d6a-f906-41c9-87a0-b760b9169ac0","resolution":{"observed_at":"2026-05-15T14:58:32.742607Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.07688","last_updated":"2025-05-26T09:45:02Z","snapshot_observed_at":"2026-08-12T22:02:14.017869Z","submitted_at":"2024-11-12T10:12:12Z","title":"ImageRAG: Enhancing Ultra High Resolution Remote Sensing Imagery Analysis with ImageRAG","version":4},"cited_work":{"arxiv_id":"2411.07688","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2411.07688","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"arXiv preprint arXiv:2411.07688 (2024) 5 Preprint 19","venue":null,"work_id":"3dbed292-bfd3-47fe-897e-229349b6970f","year":2024},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2411.07688","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:55a74f6abbbbf23a6be1aa94d196405f44e4bbcefc4094930b53aaa92b18ae5f","observation_id":"63f2b567-c4f1-44ba-a0af-5991e7c4ee71","resolution":{"observed_at":"2026-05-11T10:26:01.351825Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.10592","last_updated":"2023-10-02T16:38:35Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-04-20T18:25:35Z","title":"MiniGPT-4: Enhancing Vision-Language Understanding with Advanced Large Language Models","version":2},"cited_work":{"arxiv_id":"2304.10592","doi":"10.18653/v1/2024.findings-emnlp.692","metadata_source":"pith","pith_arxiv_id":"2304.10592","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"MiniGPT-4: Enhancing Vision-Language Understanding with Advanced Large Language Models","venue":"cs.CV","work_id":"a7e3a737-e007-42bc-be89-c4d34c5ee071","year":2023},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2304.10592","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:267cbea049d796643ea1cf269974df41d6a1a1c7398b2b97700dd5c6018108fb","observation_id":"0992bd66-29a8-4424-a98b-afdc3fa54b5b","resolution":{"observed_at":"2026-05-11T10:26:01.384042Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-05-23T17:23:55.433296+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T17:23:55.433296+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.10479","last_updated":"2025-04-19T03:47:21Z","snapshot_observed_at":"2026-08-10T18:37:57.419939Z","submitted_at":"2025-04-14T17:59:25Z","title":"InternVL3: Exploring Advanced Training and Test-Time Recipes for Open-Source Multimodal Models","version":3},"cited_work":{"arxiv_id":"2504.10479","doi":"10.48550/arxiv.2504.10479","metadata_source":"pith","pith_arxiv_id":"2504.10479","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"InternVL3: Exploring Advanced Training and Test-Time Recipes for Open-Source Multimodal Models","venue":"cs.CV","work_id":"fe8637aa-12bc-4434-8d36-9f57b5eebcbe","year":2025},"citing_paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:49.681399Z"},"links":{"cited_paper":"/paper/2504.10479","citing_paper":"/paper/2604.11122"},"observation_digest":"sha256:8ff657e5853884d56842e0f34c36e41e12fabbb0aa53ea692af453ee1793424b","observation_id":"1ef23b35-c07a-4c86-9071-5adf4b21dd47","resolution":{"observed_at":"2026-05-11T10:26:01.478342Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-05-20T07:54:09.017512+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T07:54:09.017512+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2604.11122","last_updated":"2026-04-13T07:36:02Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-10T20:53:25.115704Z","submitted_at":"2026-04-13T07:36:02Z","title":"Semantic-Geometric Dual Compression: Training-Free Visual Token Reduction for Ultra-High-Resolution Remote Sensing Understanding"},"reference_resolution":{"displayed":62,"state_counts":{"malformed_identifier":0,"metadata_mismatch":26,"parse_uncertain":0,"unresolved":4,"verified_exact":10,"verified_fuzzy":22},"total_outbound_references":62},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 14 August 2026, this Paper Citation Record lists 62 of 62 outbound references and 1 inbound Pith citation observation for arXiv:2604.11122."}