{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2021:3JW4H5HGUBAWK64X6SJFTD5RA5","short_pith_number":"pith:3JW4H5HG","schema_version":"1.0","canonical_sha256":"da6dc3f4e6a041657b97f492598fb107742dcc27ceee3e3003a9ebd43a7cc744","source":{"kind":"arxiv","id":"2107.13782","version":3},"attestation_state":"computed","paper":{"title":"Multimodal Co-learning: Challenges, Applications with Datasets, Recent Advances and Future Directions","license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.LG","authors_text":"Anil Rahate, Ketan Kotecha, Rahee Walambe, Sheela Ramanna","submitted_at":"2021-07-29T07:25:21Z","abstract_excerpt":"Multimodal deep learning systems which employ multiple modalities like text, image, audio, video, etc., are showing better performance in comparison with individual modalities (i.e., unimodal) systems. Multimodal machine learning involves multiple aspects: representation, translation, alignment, fusion, and co-learning. In the current state of multimodal machine learning, the assumptions are that all modalities are present, aligned, and noiseless during training and testing time. However, in real-world tasks, typically, it is observed that one or more modalities are missing, noisy, lacking ann"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2107.13782","kind":"arxiv","version":3},"metadata":{"license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","primary_cat":"cs.LG","submitted_at":"2021-07-29T07:25:21Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"a9f91e5ddb2b4fd027603d57b4b46f046a08bd9f55bd9374395a7337674d1ed7","abstract_canon_sha256":"8b426014269660365018f63d26149ca80e2d3a0c2d9567fcbc7493f043367e51"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T03:48:43.820032Z","signature_b64":"unXezFlcoZNnrlOKjA4QUexdjMuj59z/D95xIysEN5RjePX+qIrmzl0gHN46YJ1iqeR4O4EyRAOaStEWhEMNDA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"da6dc3f4e6a041657b97f492598fb107742dcc27ceee3e3003a9ebd43a7cc744","last_reissued_at":"2026-07-05T03:48:43.819556Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T03:48:43.819556Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Multimodal Co-learning: Challenges, Applications with Datasets, Recent Advances and Future Directions","license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.LG","authors_text":"Anil Rahate, Ketan Kotecha, Rahee Walambe, Sheela Ramanna","submitted_at":"2021-07-29T07:25:21Z","abstract_excerpt":"Multimodal deep learning systems which employ multiple modalities like text, image, audio, video, etc., are showing better performance in comparison with individual modalities (i.e., unimodal) systems. Multimodal machine learning involves multiple aspects: representation, translation, alignment, fusion, and co-learning. In the current state of multimodal machine learning, the assumptions are that all modalities are present, aligned, and noiseless during training and testing time. However, in real-world tasks, typically, it is observed that one or more modalities are missing, noisy, lacking ann"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2107.13782","kind":"arxiv","version":3},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2107.13782/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2107.13782","created_at":"2026-07-05T03:48:43.819613+00:00"},{"alias_kind":"arxiv_version","alias_value":"2107.13782v3","created_at":"2026-07-05T03:48:43.819613+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2107.13782","created_at":"2026-07-05T03:48:43.819613+00:00"},{"alias_kind":"pith_short_12","alias_value":"3JW4H5HGUBAW","created_at":"2026-07-05T03:48:43.819613+00:00"},{"alias_kind":"pith_short_16","alias_value":"3JW4H5HGUBAWK64X","created_at":"2026-07-05T03:48:43.819613+00:00"},{"alias_kind":"pith_short_8","alias_value":"3JW4H5HG","created_at":"2026-07-05T03:48:43.819613+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":0,"internal_anchor_count":0,"sample":[]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/3JW4H5HGUBAWK64X6SJFTD5RA5","json":"https://pith.science/pith/3JW4H5HGUBAWK64X6SJFTD5RA5.json","graph_json":"https://pith.science/api/pith-number/3JW4H5HGUBAWK64X6SJFTD5RA5/graph.json","events_json":"https://pith.science/api/pith-number/3JW4H5HGUBAWK64X6SJFTD5RA5/events.json","paper":"https://pith.science/paper/3JW4H5HG"},"agent_actions":{"view_html":"https://pith.science/pith/3JW4H5HGUBAWK64X6SJFTD5RA5","download_json":"https://pith.science/pith/3JW4H5HGUBAWK64X6SJFTD5RA5.json","view_paper":"https://pith.science/paper/3JW4H5HG","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2107.13782&json=true","fetch_graph":"https://pith.science/api/pith-number/3JW4H5HGUBAWK64X6SJFTD5RA5/graph.json","fetch_events":"https://pith.science/api/pith-number/3JW4H5HGUBAWK64X6SJFTD5RA5/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/3JW4H5HGUBAWK64X6SJFTD5RA5/action/timestamp_anchor","attest_storage":"https://pith.science/pith/3JW4H5HGUBAWK64X6SJFTD5RA5/action/storage_attestation","attest_author":"https://pith.science/pith/3JW4H5HGUBAWK64X6SJFTD5RA5/action/author_attestation","sign_citation":"https://pith.science/pith/3JW4H5HGUBAWK64X6SJFTD5RA5/action/citation_signature","submit_replication":"https://pith.science/pith/3JW4H5HGUBAWK64X6SJFTD5RA5/action/replication_record"}},"created_at":"2026-07-05T03:48:43.819613+00:00","updated_at":"2026-07-05T03:48:43.819613+00:00"}