{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:JY2S6PZXDN4T2C2KY7LVHUU6AW","short_pith_number":"pith:JY2S6PZX","schema_version":"1.0","canonical_sha256":"4e352f3f371b793d0b4ac7d753d29e058c58a7f5255bbe494d47e2ad5d278c9f","source":{"kind":"arxiv","id":"2411.19325","version":2},"attestation_state":"computed","paper":{"title":"GEOBench-VLM: Benchmarking Vision-Language Models for Geospatial Tasks","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.CV","authors_text":"Alexandre Lacoste, Fahad Shahbaz Khan, Kartik Kuckreja, Muhammad Akhtar Munir, Muhammad Sohail Danish, Paolo Fraccaro, Salman Khan, Syed Roshaan Ali Shah","submitted_at":"2024-11-28T18:59:56Z","abstract_excerpt":"While numerous recent benchmarks focus on evaluating generic Vision-Language Models (VLMs), they do not effectively address the specific challenges of geospatial applications. Generic VLM benchmarks are not designed to handle the complexities of geospatial data, an essential component for applications such as environmental monitoring, urban planning, and disaster management. Key challenges in the geospatial domain include temporal change detection, large-scale object counting, tiny object detection, and understanding relationships between entities in remote sensing imagery. To bridge this gap,"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2411.19325","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CV","submitted_at":"2024-11-28T18:59:56Z","cross_cats_sorted":[],"title_canon_sha256":"4cd2a84d4ab577636cf6fa7f8b168284a6e25b2fb7a3f29dcceeabd96cf93cf8","abstract_canon_sha256":"627e7f0eecf936f6f813f522c8a86c47d4776daf3a0c0bcec84876e75dc7999d"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T10:30:03.245203Z","signature_b64":"f0ROzvQ2KbusEGOcqyOyH9qZpYbq/wAdfnPGmwMGJ4zHl5Xjwv3ITuSkkVHJZL7jxvvczOPPV4JQrKuQcMb/Dw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"4e352f3f371b793d0b4ac7d753d29e058c58a7f5255bbe494d47e2ad5d278c9f","last_reissued_at":"2026-07-05T10:30:03.244595Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T10:30:03.244595Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"GEOBench-VLM: Benchmarking Vision-Language Models for Geospatial Tasks","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.CV","authors_text":"Alexandre Lacoste, Fahad Shahbaz Khan, Kartik Kuckreja, Muhammad Akhtar Munir, Muhammad Sohail Danish, Paolo Fraccaro, Salman Khan, Syed Roshaan Ali Shah","submitted_at":"2024-11-28T18:59:56Z","abstract_excerpt":"While numerous recent benchmarks focus on evaluating generic Vision-Language Models (VLMs), they do not effectively address the specific challenges of geospatial applications. Generic VLM benchmarks are not designed to handle the complexities of geospatial data, an essential component for applications such as environmental monitoring, urban planning, and disaster management. Key challenges in the geospatial domain include temporal change detection, large-scale object counting, tiny object detection, and understanding relationships between entities in remote sensing imagery. To bridge this gap,"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2411.19325","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2411.19325/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2411.19325","created_at":"2026-07-05T10:30:03.244667+00:00"},{"alias_kind":"arxiv_version","alias_value":"2411.19325v2","created_at":"2026-07-05T10:30:03.244667+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2411.19325","created_at":"2026-07-05T10:30:03.244667+00:00"},{"alias_kind":"pith_short_12","alias_value":"JY2S6PZXDN4T","created_at":"2026-07-05T10:30:03.244667+00:00"},{"alias_kind":"pith_short_16","alias_value":"JY2S6PZXDN4T2C2K","created_at":"2026-07-05T10:30:03.244667+00:00"},{"alias_kind":"pith_short_8","alias_value":"JY2S6PZX","created_at":"2026-07-05T10:30:03.244667+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":2,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2606.11740","citing_title":"UniReason-Med: A Shared Grounded Reasoning Interface for 2D-to-3D Transfer in Medical VQA","ref_index":72,"is_internal_anchor":false},{"citing_arxiv_id":"2604.08896","citing_title":"GeoMMBench and GeoMMAgent: Toward Expert-Level Multimodal Intelligence in Geoscience and Remote Sensing","ref_index":10,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/JY2S6PZXDN4T2C2KY7LVHUU6AW","json":"https://pith.science/pith/JY2S6PZXDN4T2C2KY7LVHUU6AW.json","graph_json":"https://pith.science/api/pith-number/JY2S6PZXDN4T2C2KY7LVHUU6AW/graph.json","events_json":"https://pith.science/api/pith-number/JY2S6PZXDN4T2C2KY7LVHUU6AW/events.json","paper":"https://pith.science/paper/JY2S6PZX"},"agent_actions":{"view_html":"https://pith.science/pith/JY2S6PZXDN4T2C2KY7LVHUU6AW","download_json":"https://pith.science/pith/JY2S6PZXDN4T2C2KY7LVHUU6AW.json","view_paper":"https://pith.science/paper/JY2S6PZX","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2411.19325&json=true","fetch_graph":"https://pith.science/api/pith-number/JY2S6PZXDN4T2C2KY7LVHUU6AW/graph.json","fetch_events":"https://pith.science/api/pith-number/JY2S6PZXDN4T2C2KY7LVHUU6AW/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/JY2S6PZXDN4T2C2KY7LVHUU6AW/action/timestamp_anchor","attest_storage":"https://pith.science/pith/JY2S6PZXDN4T2C2KY7LVHUU6AW/action/storage_attestation","attest_author":"https://pith.science/pith/JY2S6PZXDN4T2C2KY7LVHUU6AW/action/author_attestation","sign_citation":"https://pith.science/pith/JY2S6PZXDN4T2C2KY7LVHUU6AW/action/citation_signature","submit_replication":"https://pith.science/pith/JY2S6PZXDN4T2C2KY7LVHUU6AW/action/replication_record"}},"created_at":"2026-07-05T10:30:03.244667+00:00","updated_at":"2026-07-05T10:30:03.244667+00:00"}