{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:5ACV665A2U64MVJPIFTKHCQU2N","short_pith_number":"pith:5ACV665A","schema_version":"1.0","canonical_sha256":"e8055f7ba0d53dc6552f4166a38a14d36aaf12e1429a2a43bc3f2a0b1724762c","source":{"kind":"arxiv","id":"2502.02885","version":3},"attestation_state":"computed","paper":{"title":"Expertized Caption Auto-Enhancement for Video-Text Retrieval","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI","cs.LG"],"primary_cat":"cs.CV","authors_text":"Baoyao Yang, Junxiang Chen, Wanyun Li, Wenbin Yao, Yang Zhou","submitted_at":"2025-02-05T04:51:46Z","abstract_excerpt":"Video-text retrieval has been stuck in the information mismatch caused by personalized and inadequate textual descriptions of videos. The substantial information gap between the two modalities hinders an effective cross-modal representation alignment, resulting in ambiguous retrieval results. Although text rewriting methods have been proposed to broaden text expressions, the modality gap remains significant, as the text representation space is hardly expanded with insufficient semantic enrichment.Instead, this paper turns to enhancing visual presentation, bridging video expression closer to te"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2502.02885","kind":"arxiv","version":3},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CV","submitted_at":"2025-02-05T04:51:46Z","cross_cats_sorted":["cs.AI","cs.LG"],"title_canon_sha256":"32b860f0c521d2699767a7c65c313e8371dbb49580016539428ed007bbd87c7a","abstract_canon_sha256":"19c6fac4d04f1e4bdf0382ecde382f72f69cad84726aca8328933a5ac1bbaf36"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T10:45:59.909074Z","signature_b64":"XyLpfcf11104kcaFGP6rJb77kqrXIq4EOKe0hfT88a3y1ilZI+sLY7ZD8MdAxqbhxoj8IfSXAzfcEyHd/PfPCw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"e8055f7ba0d53dc6552f4166a38a14d36aaf12e1429a2a43bc3f2a0b1724762c","last_reissued_at":"2026-07-05T10:45:59.908578Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T10:45:59.908578Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Expertized Caption Auto-Enhancement for Video-Text Retrieval","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI","cs.LG"],"primary_cat":"cs.CV","authors_text":"Baoyao Yang, Junxiang Chen, Wanyun Li, Wenbin Yao, Yang Zhou","submitted_at":"2025-02-05T04:51:46Z","abstract_excerpt":"Video-text retrieval has been stuck in the information mismatch caused by personalized and inadequate textual descriptions of videos. The substantial information gap between the two modalities hinders an effective cross-modal representation alignment, resulting in ambiguous retrieval results. Although text rewriting methods have been proposed to broaden text expressions, the modality gap remains significant, as the text representation space is hardly expanded with insufficient semantic enrichment.Instead, this paper turns to enhancing visual presentation, bridging video expression closer to te"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2502.02885","kind":"arxiv","version":3},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2502.02885/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2502.02885","created_at":"2026-07-05T10:45:59.908636+00:00"},{"alias_kind":"arxiv_version","alias_value":"2502.02885v3","created_at":"2026-07-05T10:45:59.908636+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2502.02885","created_at":"2026-07-05T10:45:59.908636+00:00"},{"alias_kind":"pith_short_12","alias_value":"5ACV665A2U64","created_at":"2026-07-05T10:45:59.908636+00:00"},{"alias_kind":"pith_short_16","alias_value":"5ACV665A2U64MVJP","created_at":"2026-07-05T10:45:59.908636+00:00"},{"alias_kind":"pith_short_8","alias_value":"5ACV665A","created_at":"2026-07-05T10:45:59.908636+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2604.11174","citing_title":"EmbodiedGovBench: A Benchmark for Governance, Recovery, and Upgrade Safety in Embodied Agent Systems","ref_index":44,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/5ACV665A2U64MVJPIFTKHCQU2N","json":"https://pith.science/pith/5ACV665A2U64MVJPIFTKHCQU2N.json","graph_json":"https://pith.science/api/pith-number/5ACV665A2U64MVJPIFTKHCQU2N/graph.json","events_json":"https://pith.science/api/pith-number/5ACV665A2U64MVJPIFTKHCQU2N/events.json","paper":"https://pith.science/paper/5ACV665A"},"agent_actions":{"view_html":"https://pith.science/pith/5ACV665A2U64MVJPIFTKHCQU2N","download_json":"https://pith.science/pith/5ACV665A2U64MVJPIFTKHCQU2N.json","view_paper":"https://pith.science/paper/5ACV665A","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2502.02885&json=true","fetch_graph":"https://pith.science/api/pith-number/5ACV665A2U64MVJPIFTKHCQU2N/graph.json","fetch_events":"https://pith.science/api/pith-number/5ACV665A2U64MVJPIFTKHCQU2N/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/5ACV665A2U64MVJPIFTKHCQU2N/action/timestamp_anchor","attest_storage":"https://pith.science/pith/5ACV665A2U64MVJPIFTKHCQU2N/action/storage_attestation","attest_author":"https://pith.science/pith/5ACV665A2U64MVJPIFTKHCQU2N/action/author_attestation","sign_citation":"https://pith.science/pith/5ACV665A2U64MVJPIFTKHCQU2N/action/citation_signature","submit_replication":"https://pith.science/pith/5ACV665A2U64MVJPIFTKHCQU2N/action/replication_record"}},"created_at":"2026-07-05T10:45:59.908636+00:00","updated_at":"2026-07-05T10:45:59.908636+00:00"}