{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2021:E7QH3XO4PAUOBQW5UZKPJ666YM","short_pith_number":"pith:E7QH3XO4","schema_version":"1.0","canonical_sha256":"27e07ddddc7828e0c2dda654f4fbdec31c96d79bd509241a029d89fb9373d632","source":{"kind":"arxiv","id":"2110.00570","version":1},"attestation_state":"computed","paper":{"title":"Leveraging Low-Distortion Target Estimates for Improved Speech Enhancement","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["eess.AS"],"primary_cat":"cs.SD","authors_text":"Gordon Wichern, Jonathan Le Roux, Zhong-Qiu Wang","submitted_at":"2021-10-01T17:53:40Z","abstract_excerpt":"A promising approach for multi-microphone speech separation involves two deep neural networks (DNN), where the predicted target speech from the first DNN is used to compute signal statistics for time-invariant minimum variance distortionless response (MVDR) beamforming, and the MVDR result is then used as extra features for the second DNN to predict target speech. Previous studies suggested that the MVDR result can provide complementary information for the second DNN to better predict target speech. However, on fixed-geometry arrays, both DNNs can take in, for example, the real and imaginary ("},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2110.00570","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.SD","submitted_at":"2021-10-01T17:53:40Z","cross_cats_sorted":["eess.AS"],"title_canon_sha256":"c78521433df740fdbfc5e8fc6446117a50ff4f7a460cd963c59efab15df6dd20","abstract_canon_sha256":"15627e4762eca28c98f0714292bd932b560c427f97fc6693f081c337fe858ffd"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T03:19:30.978857Z","signature_b64":"3Ia6nFZQQTk8Z+aZlnnO4wKGbaK3dD0eYxGYQTdt0aioJbYKsq6jM6t0hMYU/zwyRgfrHARo4KKA3YdkPTlpDw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"27e07ddddc7828e0c2dda654f4fbdec31c96d79bd509241a029d89fb9373d632","last_reissued_at":"2026-07-05T03:19:30.978398Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T03:19:30.978398Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Leveraging Low-Distortion Target Estimates for Improved Speech Enhancement","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["eess.AS"],"primary_cat":"cs.SD","authors_text":"Gordon Wichern, Jonathan Le Roux, Zhong-Qiu Wang","submitted_at":"2021-10-01T17:53:40Z","abstract_excerpt":"A promising approach for multi-microphone speech separation involves two deep neural networks (DNN), where the predicted target speech from the first DNN is used to compute signal statistics for time-invariant minimum variance distortionless response (MVDR) beamforming, and the MVDR result is then used as extra features for the second DNN to predict target speech. Previous studies suggested that the MVDR result can provide complementary information for the second DNN to better predict target speech. However, on fixed-geometry arrays, both DNNs can take in, for example, the real and imaginary ("},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2110.00570","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2110.00570/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2110.00570","created_at":"2026-07-05T03:19:30.978459+00:00"},{"alias_kind":"arxiv_version","alias_value":"2110.00570v1","created_at":"2026-07-05T03:19:30.978459+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2110.00570","created_at":"2026-07-05T03:19:30.978459+00:00"},{"alias_kind":"pith_short_12","alias_value":"E7QH3XO4PAUO","created_at":"2026-07-05T03:19:30.978459+00:00"},{"alias_kind":"pith_short_16","alias_value":"E7QH3XO4PAUOBQW5","created_at":"2026-07-05T03:19:30.978459+00:00"},{"alias_kind":"pith_short_8","alias_value":"E7QH3XO4","created_at":"2026-07-05T03:19:30.978459+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":1,"sample":[{"citing_arxiv_id":"2507.02192","citing_title":"An Investigation on Combining Geometry and Consistency Constraints into Phase Estimation for Speech Enhancement","ref_index":32,"is_internal_anchor":true}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/E7QH3XO4PAUOBQW5UZKPJ666YM","json":"https://pith.science/pith/E7QH3XO4PAUOBQW5UZKPJ666YM.json","graph_json":"https://pith.science/api/pith-number/E7QH3XO4PAUOBQW5UZKPJ666YM/graph.json","events_json":"https://pith.science/api/pith-number/E7QH3XO4PAUOBQW5UZKPJ666YM/events.json","paper":"https://pith.science/paper/E7QH3XO4"},"agent_actions":{"view_html":"https://pith.science/pith/E7QH3XO4PAUOBQW5UZKPJ666YM","download_json":"https://pith.science/pith/E7QH3XO4PAUOBQW5UZKPJ666YM.json","view_paper":"https://pith.science/paper/E7QH3XO4","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2110.00570&json=true","fetch_graph":"https://pith.science/api/pith-number/E7QH3XO4PAUOBQW5UZKPJ666YM/graph.json","fetch_events":"https://pith.science/api/pith-number/E7QH3XO4PAUOBQW5UZKPJ666YM/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/E7QH3XO4PAUOBQW5UZKPJ666YM/action/timestamp_anchor","attest_storage":"https://pith.science/pith/E7QH3XO4PAUOBQW5UZKPJ666YM/action/storage_attestation","attest_author":"https://pith.science/pith/E7QH3XO4PAUOBQW5UZKPJ666YM/action/author_attestation","sign_citation":"https://pith.science/pith/E7QH3XO4PAUOBQW5UZKPJ666YM/action/citation_signature","submit_replication":"https://pith.science/pith/E7QH3XO4PAUOBQW5UZKPJ666YM/action/replication_record"}},"created_at":"2026-07-05T03:19:30.978459+00:00","updated_at":"2026-07-05T03:19:30.978459+00:00"}