{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2021:QYVSGIHSU5W2ZJKNF6P5GNY7GK","short_pith_number":"pith:QYVSGIHS","schema_version":"1.0","canonical_sha256":"862b2320f2a76daca54d2f9fd3371f3281c9d0ae4e6655a4c3c173718dbf6cef","source":{"kind":"arxiv","id":"2104.08806","version":2},"attestation_state":"computed","paper":{"title":"Best Practices for Noise-Based Augmentation to Improve the Performance of Deployable Speech-Based Emotion Recognition Systems","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.LG","eess.AS"],"primary_cat":"cs.SD","authors_text":"Emily Mower Provost, Mimansa Jaiswal","submitted_at":"2021-04-18T10:33:38Z","abstract_excerpt":"Speech emotion recognition is an important component of any human centered system. But speech characteristics produced and perceived by a person can be influenced by a multitude of reasons, both desirable such as emotion, and undesirable such as noise. To train robust emotion recognition models, we need a large, yet realistic data distribution, but emotion datasets are often small and hence are augmented with noise. Often noise augmentation makes one important assumption, that the prediction label should remain the same in presence or absence of noise, which is true for automatic speech recogn"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2104.08806","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.SD","submitted_at":"2021-04-18T10:33:38Z","cross_cats_sorted":["cs.LG","eess.AS"],"title_canon_sha256":"01506558a47808beaf9f8cb152a240786fdf44f1eef75daec7800c87abada21a","abstract_canon_sha256":"52f8f6b815ab104eee12f4b2e9502cc9d8e461cbc4a09c09db9b9d458a880e9c"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T06:46:42.363555Z","signature_b64":"UD2mJEUr/9Ltlz5sq9aYm7SLqHfTfWOJGyLFQKICGC+NqHL5oumXBE/5ySayxvM0ePDIWhh+Y5/XQ1fFO3yyDA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"862b2320f2a76daca54d2f9fd3371f3281c9d0ae4e6655a4c3c173718dbf6cef","last_reissued_at":"2026-07-05T06:46:42.363013Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T06:46:42.363013Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Best Practices for Noise-Based Augmentation to Improve the Performance of Deployable Speech-Based Emotion Recognition Systems","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.LG","eess.AS"],"primary_cat":"cs.SD","authors_text":"Emily Mower Provost, Mimansa Jaiswal","submitted_at":"2021-04-18T10:33:38Z","abstract_excerpt":"Speech emotion recognition is an important component of any human centered system. But speech characteristics produced and perceived by a person can be influenced by a multitude of reasons, both desirable such as emotion, and undesirable such as noise. To train robust emotion recognition models, we need a large, yet realistic data distribution, but emotion datasets are often small and hence are augmented with noise. Often noise augmentation makes one important assumption, that the prediction label should remain the same in presence or absence of noise, which is true for automatic speech recogn"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2104.08806","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2104.08806/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2104.08806","created_at":"2026-07-05T06:46:42.363071+00:00"},{"alias_kind":"arxiv_version","alias_value":"2104.08806v2","created_at":"2026-07-05T06:46:42.363071+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2104.08806","created_at":"2026-07-05T06:46:42.363071+00:00"},{"alias_kind":"pith_short_12","alias_value":"QYVSGIHSU5W2","created_at":"2026-07-05T06:46:42.363071+00:00"},{"alias_kind":"pith_short_16","alias_value":"QYVSGIHSU5W2ZJKN","created_at":"2026-07-05T06:46:42.363071+00:00"},{"alias_kind":"pith_short_8","alias_value":"QYVSGIHS","created_at":"2026-07-05T06:46:42.363071+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2509.08470","citing_title":"Joint Learning using Mixture-of-Expert-Based Representation for Speech Enhancement and Robust Emotion Recognition","ref_index":15,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/QYVSGIHSU5W2ZJKNF6P5GNY7GK","json":"https://pith.science/pith/QYVSGIHSU5W2ZJKNF6P5GNY7GK.json","graph_json":"https://pith.science/api/pith-number/QYVSGIHSU5W2ZJKNF6P5GNY7GK/graph.json","events_json":"https://pith.science/api/pith-number/QYVSGIHSU5W2ZJKNF6P5GNY7GK/events.json","paper":"https://pith.science/paper/QYVSGIHS"},"agent_actions":{"view_html":"https://pith.science/pith/QYVSGIHSU5W2ZJKNF6P5GNY7GK","download_json":"https://pith.science/pith/QYVSGIHSU5W2ZJKNF6P5GNY7GK.json","view_paper":"https://pith.science/paper/QYVSGIHS","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2104.08806&json=true","fetch_graph":"https://pith.science/api/pith-number/QYVSGIHSU5W2ZJKNF6P5GNY7GK/graph.json","fetch_events":"https://pith.science/api/pith-number/QYVSGIHSU5W2ZJKNF6P5GNY7GK/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/QYVSGIHSU5W2ZJKNF6P5GNY7GK/action/timestamp_anchor","attest_storage":"https://pith.science/pith/QYVSGIHSU5W2ZJKNF6P5GNY7GK/action/storage_attestation","attest_author":"https://pith.science/pith/QYVSGIHSU5W2ZJKNF6P5GNY7GK/action/author_attestation","sign_citation":"https://pith.science/pith/QYVSGIHSU5W2ZJKNF6P5GNY7GK/action/citation_signature","submit_replication":"https://pith.science/pith/QYVSGIHSU5W2ZJKNF6P5GNY7GK/action/replication_record"}},"created_at":"2026-07-05T06:46:42.363071+00:00","updated_at":"2026-07-05T06:46:42.363071+00:00"}