{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2023:UJTW7STL67TZL7L5L7NQIDILAH","short_pith_number":"pith:UJTW7STL","schema_version":"1.0","canonical_sha256":"a2676fca6bf7e795fd7d5fdb040d0b01dc0243f0bbd0892d02366c5f40979d4a","source":{"kind":"arxiv","id":"2309.15494","version":1},"attestation_state":"computed","paper":{"title":"VideoAdviser: Video Knowledge Distillation for Multimodal Transfer Learning","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.CL"],"primary_cat":"cs.CV","authors_text":"Donghuo Zeng, Satoshi Kurihara, Shinya Wada, Yanan Wang","submitted_at":"2023-09-27T08:44:04Z","abstract_excerpt":"Multimodal transfer learning aims to transform pretrained representations of diverse modalities into a common domain space for effective multimodal fusion. However, conventional systems are typically built on the assumption that all modalities exist, and the lack of modalities always leads to poor inference performance. Furthermore, extracting pretrained embeddings for all modalities is computationally inefficient for inference. In this work, to achieve high efficiency-performance multimodal transfer learning, we propose VideoAdviser, a video knowledge distillation method to transfer multimoda"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2309.15494","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CV","submitted_at":"2023-09-27T08:44:04Z","cross_cats_sorted":["cs.CL"],"title_canon_sha256":"240a1ebfc454724b1f69b826f62da93184acf837ca795f552ae8ebb9df1ae719","abstract_canon_sha256":"81ee9b1979dff21ee48405530b0e05ea78b4371176dcfb8862a2d53c63110ed2"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T06:54:46.966809Z","signature_b64":"Ez8hkn6TRTB3OC6/owB1RnWGzJ2PXZwIHISkA1lECwO1iClLJ4WS6CPp3UUhcAgCOSjqKgjIXlyI7Top0k8kBA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"a2676fca6bf7e795fd7d5fdb040d0b01dc0243f0bbd0892d02366c5f40979d4a","last_reissued_at":"2026-07-05T06:54:46.966373Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T06:54:46.966373Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"VideoAdviser: Video Knowledge Distillation for Multimodal Transfer Learning","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.CL"],"primary_cat":"cs.CV","authors_text":"Donghuo Zeng, Satoshi Kurihara, Shinya Wada, Yanan Wang","submitted_at":"2023-09-27T08:44:04Z","abstract_excerpt":"Multimodal transfer learning aims to transform pretrained representations of diverse modalities into a common domain space for effective multimodal fusion. However, conventional systems are typically built on the assumption that all modalities exist, and the lack of modalities always leads to poor inference performance. Furthermore, extracting pretrained embeddings for all modalities is computationally inefficient for inference. In this work, to achieve high efficiency-performance multimodal transfer learning, we propose VideoAdviser, a video knowledge distillation method to transfer multimoda"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2309.15494","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2309.15494/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2309.15494","created_at":"2026-07-05T06:54:46.966435+00:00"},{"alias_kind":"arxiv_version","alias_value":"2309.15494v1","created_at":"2026-07-05T06:54:46.966435+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2309.15494","created_at":"2026-07-05T06:54:46.966435+00:00"},{"alias_kind":"pith_short_12","alias_value":"UJTW7STL67TZ","created_at":"2026-07-05T06:54:46.966435+00:00"},{"alias_kind":"pith_short_16","alias_value":"UJTW7STL67TZL7L5","created_at":"2026-07-05T06:54:46.966435+00:00"},{"alias_kind":"pith_short_8","alias_value":"UJTW7STL","created_at":"2026-07-05T06:54:46.966435+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":0,"internal_anchor_count":0,"sample":[]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/UJTW7STL67TZL7L5L7NQIDILAH","json":"https://pith.science/pith/UJTW7STL67TZL7L5L7NQIDILAH.json","graph_json":"https://pith.science/api/pith-number/UJTW7STL67TZL7L5L7NQIDILAH/graph.json","events_json":"https://pith.science/api/pith-number/UJTW7STL67TZL7L5L7NQIDILAH/events.json","paper":"https://pith.science/paper/UJTW7STL"},"agent_actions":{"view_html":"https://pith.science/pith/UJTW7STL67TZL7L5L7NQIDILAH","download_json":"https://pith.science/pith/UJTW7STL67TZL7L5L7NQIDILAH.json","view_paper":"https://pith.science/paper/UJTW7STL","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2309.15494&json=true","fetch_graph":"https://pith.science/api/pith-number/UJTW7STL67TZL7L5L7NQIDILAH/graph.json","fetch_events":"https://pith.science/api/pith-number/UJTW7STL67TZL7L5L7NQIDILAH/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/UJTW7STL67TZL7L5L7NQIDILAH/action/timestamp_anchor","attest_storage":"https://pith.science/pith/UJTW7STL67TZL7L5L7NQIDILAH/action/storage_attestation","attest_author":"https://pith.science/pith/UJTW7STL67TZL7L5L7NQIDILAH/action/author_attestation","sign_citation":"https://pith.science/pith/UJTW7STL67TZL7L5L7NQIDILAH/action/citation_signature","submit_replication":"https://pith.science/pith/UJTW7STL67TZL7L5L7NQIDILAH/action/replication_record"}},"created_at":"2026-07-05T06:54:46.966435+00:00","updated_at":"2026-07-05T06:54:46.966435+00:00"}