{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:4YES4GWNZVXYZD67O6WLI43HXI","short_pith_number":"pith:4YES4GWN","schema_version":"1.0","canonical_sha256":"e6092e1acdcd6f8c8fdf77acb47367ba280a25a9b3bc4a6029721ddd730d18c0","source":{"kind":"arxiv","id":"2408.05831","version":1},"attestation_state":"computed","paper":{"title":"Robust Domain Generalization for Multi-modal Object Recognition","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.CV","authors_text":"Chufeng Jiang, Haoyu Yang, Junhong Lin, Keqin Li, Rong Wei, Yang Luo, Yuxin Qiao","submitted_at":"2024-08-11T17:13:21Z","abstract_excerpt":"In multi-label classification, machine learning encounters the challenge of domain generalization when handling tasks with distributions differing from the training data. Existing approaches primarily focus on vision object recognition and neglect the integration of natural language. Recent advancements in vision-language pre-training leverage supervision from extensive visual-language pairs, enabling learning across diverse domains and enhancing recognition in multi-modal scenarios. However, these approaches face limitations in loss function utilization, generality across backbones, and class"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2408.05831","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CV","submitted_at":"2024-08-11T17:13:21Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"26abd0255da5b592ff024af4fcf8cbeebc704274ae68c7c766c79a9a7449c97a","abstract_canon_sha256":"a3e993b26dbbd715f8ec4cb503bbaf79b70cd57bfcd08ec7ab4d90f925348847"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T08:54:25.809613Z","signature_b64":"0LflHzESHpIuMxDtbg4679t0/1LKVxP2UWtwbBFE/Upvi/608B+Yz3Riy6mDJwJPcfNUicoFPrcraMDbY+Z6AQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"e6092e1acdcd6f8c8fdf77acb47367ba280a25a9b3bc4a6029721ddd730d18c0","last_reissued_at":"2026-07-05T08:54:25.809220Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T08:54:25.809220Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Robust Domain Generalization for Multi-modal Object Recognition","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.CV","authors_text":"Chufeng Jiang, Haoyu Yang, Junhong Lin, Keqin Li, Rong Wei, Yang Luo, Yuxin Qiao","submitted_at":"2024-08-11T17:13:21Z","abstract_excerpt":"In multi-label classification, machine learning encounters the challenge of domain generalization when handling tasks with distributions differing from the training data. Existing approaches primarily focus on vision object recognition and neglect the integration of natural language. Recent advancements in vision-language pre-training leverage supervision from extensive visual-language pairs, enabling learning across diverse domains and enhancing recognition in multi-modal scenarios. However, these approaches face limitations in loss function utilization, generality across backbones, and class"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2408.05831","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2408.05831/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2408.05831","created_at":"2026-07-05T08:54:25.809279+00:00"},{"alias_kind":"arxiv_version","alias_value":"2408.05831v1","created_at":"2026-07-05T08:54:25.809279+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2408.05831","created_at":"2026-07-05T08:54:25.809279+00:00"},{"alias_kind":"pith_short_12","alias_value":"4YES4GWNZVXY","created_at":"2026-07-05T08:54:25.809279+00:00"},{"alias_kind":"pith_short_16","alias_value":"4YES4GWNZVXYZD67","created_at":"2026-07-05T08:54:25.809279+00:00"},{"alias_kind":"pith_short_8","alias_value":"4YES4GWN","created_at":"2026-07-05T08:54:25.809279+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":1,"sample":[{"citing_arxiv_id":"2412.16935","citing_title":"Detecting and Classifying Defective Products in Images Using YOLO","ref_index":30,"is_internal_anchor":true}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/4YES4GWNZVXYZD67O6WLI43HXI","json":"https://pith.science/pith/4YES4GWNZVXYZD67O6WLI43HXI.json","graph_json":"https://pith.science/api/pith-number/4YES4GWNZVXYZD67O6WLI43HXI/graph.json","events_json":"https://pith.science/api/pith-number/4YES4GWNZVXYZD67O6WLI43HXI/events.json","paper":"https://pith.science/paper/4YES4GWN"},"agent_actions":{"view_html":"https://pith.science/pith/4YES4GWNZVXYZD67O6WLI43HXI","download_json":"https://pith.science/pith/4YES4GWNZVXYZD67O6WLI43HXI.json","view_paper":"https://pith.science/paper/4YES4GWN","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2408.05831&json=true","fetch_graph":"https://pith.science/api/pith-number/4YES4GWNZVXYZD67O6WLI43HXI/graph.json","fetch_events":"https://pith.science/api/pith-number/4YES4GWNZVXYZD67O6WLI43HXI/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/4YES4GWNZVXYZD67O6WLI43HXI/action/timestamp_anchor","attest_storage":"https://pith.science/pith/4YES4GWNZVXYZD67O6WLI43HXI/action/storage_attestation","attest_author":"https://pith.science/pith/4YES4GWNZVXYZD67O6WLI43HXI/action/author_attestation","sign_citation":"https://pith.science/pith/4YES4GWNZVXYZD67O6WLI43HXI/action/citation_signature","submit_replication":"https://pith.science/pith/4YES4GWNZVXYZD67O6WLI43HXI/action/replication_record"}},"created_at":"2026-07-05T08:54:25.809279+00:00","updated_at":"2026-07-05T08:54:25.809279+00:00"}