{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2022:VHVTRTGED7MOYUVCFHCIKKSYLU","short_pith_number":"pith:VHVTRTGE","schema_version":"1.0","canonical_sha256":"a9eb38ccc41fd8ec52a229c4852a585d0617631a0f377fbf24f75880469f49a9","source":{"kind":"arxiv","id":"2204.13653","version":2},"attestation_state":"computed","paper":{"title":"GRIT: General Robust Image Task Benchmark","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":[],"primary_cat":"cs.CV","authors_text":"Aniruddha Kembhavi, Derek Hoiem, Ryan Marten, Tanmay Gupta","submitted_at":"2022-04-28T17:13:23Z","abstract_excerpt":"Computer vision models excel at making predictions when the test distribution closely resembles the training distribution. Such models have yet to match the ability of biological vision to learn from multiple sources and generalize to new data sources and tasks. To facilitate the development and evaluation of more general vision systems, we introduce the General Robust Image Task (GRIT) benchmark. GRIT evaluates the performance, robustness, and calibration of a vision system across a variety of image prediction tasks, concepts, and data sources. The seven tasks in GRIT are selected to cover a "},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2204.13653","kind":"arxiv","version":2},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CV","submitted_at":"2022-04-28T17:13:23Z","cross_cats_sorted":[],"title_canon_sha256":"148cdd2d4f486a3b05016cb1b81d104c1086a28d65d562baf28260fec9345064","abstract_canon_sha256":"3fb64549e9fa77c1723ef6f0e6671f1985e17d0584ac8190636f5ac2d8d9966f"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T04:19:59.794216Z","signature_b64":"5cm0jB6NO3C9w488EvaGP/zdZfAUU1xnFzJ8wAxaHva5I2WH8wqlsAUyIAVVDFCKJEbtTfKqKJadN9H+zgjKAg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"a9eb38ccc41fd8ec52a229c4852a585d0617631a0f377fbf24f75880469f49a9","last_reissued_at":"2026-07-05T04:19:59.793744Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T04:19:59.793744Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"GRIT: General Robust Image Task Benchmark","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":[],"primary_cat":"cs.CV","authors_text":"Aniruddha Kembhavi, Derek Hoiem, Ryan Marten, Tanmay Gupta","submitted_at":"2022-04-28T17:13:23Z","abstract_excerpt":"Computer vision models excel at making predictions when the test distribution closely resembles the training distribution. Such models have yet to match the ability of biological vision to learn from multiple sources and generalize to new data sources and tasks. To facilitate the development and evaluation of more general vision systems, we introduce the General Robust Image Task (GRIT) benchmark. GRIT evaluates the performance, robustness, and calibration of a vision system across a variety of image prediction tasks, concepts, and data sources. The seven tasks in GRIT are selected to cover a "},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2204.13653","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2204.13653/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2204.13653","created_at":"2026-07-05T04:19:59.793803+00:00"},{"alias_kind":"arxiv_version","alias_value":"2204.13653v2","created_at":"2026-07-05T04:19:59.793803+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2204.13653","created_at":"2026-07-05T04:19:59.793803+00:00"},{"alias_kind":"pith_short_12","alias_value":"VHVTRTGED7MO","created_at":"2026-07-05T04:19:59.793803+00:00"},{"alias_kind":"pith_short_16","alias_value":"VHVTRTGED7MOYUVC","created_at":"2026-07-05T04:19:59.793803+00:00"},{"alias_kind":"pith_short_8","alias_value":"VHVTRTGE","created_at":"2026-07-05T04:19:59.793803+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":2,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2606.17030","citing_title":"Qwen-RobotWorld Technical Report: Unifying Embodied World Modeling through Language-Conditioned Video Generation","ref_index":152,"is_internal_anchor":false},{"citing_arxiv_id":"2607.00338","citing_title":"DroneFINE: Domain-Aware Parameter-Efficient Fine-Tuning of Vision-Language Detectors for Drone Images","ref_index":10,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/VHVTRTGED7MOYUVCFHCIKKSYLU","json":"https://pith.science/pith/VHVTRTGED7MOYUVCFHCIKKSYLU.json","graph_json":"https://pith.science/api/pith-number/VHVTRTGED7MOYUVCFHCIKKSYLU/graph.json","events_json":"https://pith.science/api/pith-number/VHVTRTGED7MOYUVCFHCIKKSYLU/events.json","paper":"https://pith.science/paper/VHVTRTGE"},"agent_actions":{"view_html":"https://pith.science/pith/VHVTRTGED7MOYUVCFHCIKKSYLU","download_json":"https://pith.science/pith/VHVTRTGED7MOYUVCFHCIKKSYLU.json","view_paper":"https://pith.science/paper/VHVTRTGE","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2204.13653&json=true","fetch_graph":"https://pith.science/api/pith-number/VHVTRTGED7MOYUVCFHCIKKSYLU/graph.json","fetch_events":"https://pith.science/api/pith-number/VHVTRTGED7MOYUVCFHCIKKSYLU/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/VHVTRTGED7MOYUVCFHCIKKSYLU/action/timestamp_anchor","attest_storage":"https://pith.science/pith/VHVTRTGED7MOYUVCFHCIKKSYLU/action/storage_attestation","attest_author":"https://pith.science/pith/VHVTRTGED7MOYUVCFHCIKKSYLU/action/author_attestation","sign_citation":"https://pith.science/pith/VHVTRTGED7MOYUVCFHCIKKSYLU/action/citation_signature","submit_replication":"https://pith.science/pith/VHVTRTGED7MOYUVCFHCIKKSYLU/action/replication_record"}},"created_at":"2026-07-05T04:19:59.793803+00:00","updated_at":"2026-07-05T04:19:59.793803+00:00"}