{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:HCFDK2F7QLR6CS7IZWPGQZ5YEY","short_pith_number":"pith:HCFDK2F7","schema_version":"1.0","canonical_sha256":"388a3568bf82e3e14be8cd9e6867b8261678f1a24b10ae818da03dbfcaa932bf","source":{"kind":"arxiv","id":"2410.03459","version":1},"attestation_state":"computed","paper":{"title":"Generative Semantic Communication for Text-to-Speech Synthesis","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.IT","cs.LG","eess.AS","math.IT"],"primary_cat":"cs.SD","authors_text":"Fangxin Wang, Gui Gui, Jiahao Zheng, Jie Xu, Jinke Ren, Peng Xu, Shuguang Cui, Zhihao Yuan","submitted_at":"2024-10-04T14:18:31Z","abstract_excerpt":"Semantic communication is a promising technology to improve communication efficiency by transmitting only the semantic information of the source data. However, traditional semantic communication methods primarily focus on data reconstruction tasks, which may not be efficient for emerging generative tasks such as text-to-speech (TTS) synthesis. To address this limitation, this paper develops a novel generative semantic communication framework for TTS synthesis, leveraging generative artificial intelligence technologies. Firstly, we utilize a pre-trained large speech model called WavLM and the r"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2410.03459","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.SD","submitted_at":"2024-10-04T14:18:31Z","cross_cats_sorted":["cs.IT","cs.LG","eess.AS","math.IT"],"title_canon_sha256":"ec031b3c0a6bac803e31cca4aa40d440096fe0cc77e4d65198b9698ce4c4ff58","abstract_canon_sha256":"a006a7539bb8edefa1c8f832eb2e5fa26ffd03161c1049fff5429858be8c4f27"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T09:16:01.408216Z","signature_b64":"yn0roA/EGrUEHXyRBnKZpxaEbq5C9ciuY06e8Sp+hcTb4BBRsUQ24MIhoUj5LKd647GLCAqa89daAxfSK/V7Bw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"388a3568bf82e3e14be8cd9e6867b8261678f1a24b10ae818da03dbfcaa932bf","last_reissued_at":"2026-07-05T09:16:01.407732Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T09:16:01.407732Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Generative Semantic Communication for Text-to-Speech Synthesis","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.IT","cs.LG","eess.AS","math.IT"],"primary_cat":"cs.SD","authors_text":"Fangxin Wang, Gui Gui, Jiahao Zheng, Jie Xu, Jinke Ren, Peng Xu, Shuguang Cui, Zhihao Yuan","submitted_at":"2024-10-04T14:18:31Z","abstract_excerpt":"Semantic communication is a promising technology to improve communication efficiency by transmitting only the semantic information of the source data. However, traditional semantic communication methods primarily focus on data reconstruction tasks, which may not be efficient for emerging generative tasks such as text-to-speech (TTS) synthesis. To address this limitation, this paper develops a novel generative semantic communication framework for TTS synthesis, leveraging generative artificial intelligence technologies. Firstly, we utilize a pre-trained large speech model called WavLM and the r"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2410.03459","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2410.03459/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2410.03459","created_at":"2026-07-05T09:16:01.407789+00:00"},{"alias_kind":"arxiv_version","alias_value":"2410.03459v1","created_at":"2026-07-05T09:16:01.407789+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2410.03459","created_at":"2026-07-05T09:16:01.407789+00:00"},{"alias_kind":"pith_short_12","alias_value":"HCFDK2F7QLR6","created_at":"2026-07-05T09:16:01.407789+00:00"},{"alias_kind":"pith_short_16","alias_value":"HCFDK2F7QLR6CS7I","created_at":"2026-07-05T09:16:01.407789+00:00"},{"alias_kind":"pith_short_8","alias_value":"HCFDK2F7","created_at":"2026-07-05T09:16:01.407789+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":2,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2605.18457","citing_title":"Sense Smarter, Think Better: A Survey on Edge Perception for Next-Generation Networks","ref_index":124,"is_internal_anchor":false},{"citing_arxiv_id":"2605.18457","citing_title":"Sense Smarter, Think Better: A Survey on Edge Perception for Next-Generation Networks","ref_index":124,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/HCFDK2F7QLR6CS7IZWPGQZ5YEY","json":"https://pith.science/pith/HCFDK2F7QLR6CS7IZWPGQZ5YEY.json","graph_json":"https://pith.science/api/pith-number/HCFDK2F7QLR6CS7IZWPGQZ5YEY/graph.json","events_json":"https://pith.science/api/pith-number/HCFDK2F7QLR6CS7IZWPGQZ5YEY/events.json","paper":"https://pith.science/paper/HCFDK2F7"},"agent_actions":{"view_html":"https://pith.science/pith/HCFDK2F7QLR6CS7IZWPGQZ5YEY","download_json":"https://pith.science/pith/HCFDK2F7QLR6CS7IZWPGQZ5YEY.json","view_paper":"https://pith.science/paper/HCFDK2F7","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2410.03459&json=true","fetch_graph":"https://pith.science/api/pith-number/HCFDK2F7QLR6CS7IZWPGQZ5YEY/graph.json","fetch_events":"https://pith.science/api/pith-number/HCFDK2F7QLR6CS7IZWPGQZ5YEY/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/HCFDK2F7QLR6CS7IZWPGQZ5YEY/action/timestamp_anchor","attest_storage":"https://pith.science/pith/HCFDK2F7QLR6CS7IZWPGQZ5YEY/action/storage_attestation","attest_author":"https://pith.science/pith/HCFDK2F7QLR6CS7IZWPGQZ5YEY/action/author_attestation","sign_citation":"https://pith.science/pith/HCFDK2F7QLR6CS7IZWPGQZ5YEY/action/citation_signature","submit_replication":"https://pith.science/pith/HCFDK2F7QLR6CS7IZWPGQZ5YEY/action/replication_record"}},"created_at":"2026-07-05T09:16:01.407789+00:00","updated_at":"2026-07-05T09:16:01.407789+00:00"}