{"as_of":"2026-08-08T09:15:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:8a152e0a3897897a9f242aca29eae06032508c98fa626360c074e9c8da96d452","coverage":[{"denominator":49,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":49,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T12:31:39.694445Z","state":"measured"},{"denominator":50,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":50,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T12:31:34.636801Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-07T12:31:40.495855Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"cited_work":{"arxiv_id":"2505.24248","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.24248","snapshot_observed_at":"2026-08-07T12:31:40.495855Z","title":"Probing the Robustness Properties of Neural Speech Codecs","venue":"eess.AS","work_id":"070783ec-069f-4d85-b547-7c3b0733b1ef","year":2025},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:34.636801Z"},"links":{"cited_paper":"/paper/2505.24248","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:259f08311b6230d0a37c1ed2fd3317cff49aa06d725be64d1288f2ea80428520","observation_id":"2ebc279a-ec1f-4857-932e-8d04bad44066","resolution":{"observed_at":"2026-08-07T12:31:40.571669Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2505.24248/citation-record","integrity":"/paper/2505.24248/integrity","json":"/paper/2505.24248/citation-record.json","paper":"/paper/2505.24248"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"cited_work":{"arxiv_id":"2505.24248","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.24248","snapshot_observed_at":"2026-08-07T12:31:40.495855Z","title":"Probing the Robustness Properties of Neural Speech Codecs","venue":"eess.AS","work_id":"070783ec-069f-4d85-b547-7c3b0733b1ef","year":2025},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:34.636801Z"},"links":{"cited_paper":"/paper/2505.24248","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:259f08311b6230d0a37c1ed2fd3317cff49aa06d725be64d1288f2ea80428520","observation_id":"2ebc279a-ec1f-4857-932e-8d04bad44066","resolution":{"observed_at":"2026-08-07T12:31:40.571669Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:43.902576Z","title":null,"venue":null,"work_id":"342433b7-a841-419c-bc96-b3add6d0eab3","year":null},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:34.710512Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:bc3ecfd580d86042cbd9cfc0dbfa85b71a7077bd90df9e156cb9a385990d0c96","observation_id":"cf5ecc7b-3650-479a-8435-ed52719b7d27","resolution":{"observed_at":"2026-08-07T12:31:43.969553Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:43.765147Z","title":"3.1), the evaluated neu- ral speech codecs(sec","venue":null,"work_id":"a58ef885-7355-427a-9bf9-c1b0eb79a117","year":null},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:34.799065Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:0b86ae137eb1ec692e74bace32e6112b78b27c49601a3c68069a61a41c44d14f","observation_id":"adbb5a53-2a4a-40a6-9b21-bf7b64e0013e","resolution":{"observed_at":"2026-08-07T12:31:43.825291Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:43.620156Z","title":"However, their deployment in real-world environ- ments—where background noise is common—raises concerns about their robustness","venue":null,"work_id":"e96888e3-e559-4664-beed-b997131739fd","year":null},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:34.923634Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:5b5b1c0e701f3ed730c673d3aae281e65a9660c8c422fc6cbf163564c70df589","observation_id":"9610600b-ec35-4fc9-b09a-9f3d9d11c2fe","resolution":{"observed_at":"2026-08-07T12:31:43.684231Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:43.434884Z","title":null,"venue":null,"work_id":"df0778c4-f671-4c9b-9655-51d1c7038d22","year":null},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:35.030713Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:fcb360230071ee577771f6aa2c144d6378800287e35b302d94c542729e770871","observation_id":"ce441ca0-ade2-449e-9417-7e5a56e0ec23","resolution":{"observed_at":"2026-08-07T12:31:43.541799Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:43.319568Z","title":"While compression inevitably introduces distortions, a robust neural codec should maintain spectral integrity","venue":null,"work_id":"e8cede21-f4bc-4f38-b8a4-2c4d0104931d","year":null},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:35.139267Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:e85982597a93a01612286da6c5a48fc59c68c7c717286d22282344502d746043","observation_id":"7cc5dbfb-3f52-4485-b51a-2eb5eb3faed5","resolution":{"observed_at":"2026-08-07T12:31:43.384431Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:43.168988Z","title":"Our findings reveal sig- nificant variations in their noise robustness, influenced by fac- tors such as training data diversity, operating bitrate, and quanti- zation strategies","venue":null,"work_id":"0d0bbb8d-f11d-412a-b7f8-4abaa9414f1b","year":null},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:35.301578Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:6b19bfda7175a3b3936c9cd12f143915862b0010d616ade90596c7d938908fb7","observation_id":"6259b171-a1bc-49e6-94ef-8d7a95ec6da7","resolution":{"observed_at":"2026-08-07T12:31:43.235442Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:43.010617Z","title":"7 khz audio coding within 64 kbit/s,","venue":null,"work_id":"fb73060f-bafc-4bcc-b81e-5764671ca85b","year":1988},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:35.406798Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:6e8e2d31ab3142c7cadfcf04902dc4cfe47f58106334361f1c77275f2ece348e","observation_id":"56a1f766-f27d-4329-8e20-d3fa722efac8","resolution":{"observed_at":"2026-08-07T12:31:43.084592Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:42.837574Z","title":"Mpeg-4 low delay audio coding based on the aac codec,","venue":null,"work_id":"9a35a76e-d993-4f08-8e86-c4fccbd7dad6","year":1999},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:35.570135Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:67b261c95aea61912e1d838bb262a2e9e7fe4df549104f7796e42b1bbfe8ae32","observation_id":"22e69f6d-c1f6-4ebc-9f07-586cd636e59c","resolution":{"observed_at":"2026-08-07T12:31:42.911813Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1602.04845","last_updated":"2016-02-15T21:30:54Z","snapshot_observed_at":"2026-07-06T04:46:15.993249Z","submitted_at":"2016-02-15T21:30:54Z","title":"High-Quality, Low-Delay Music Coding in the Opus Codec","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1602.04845","snapshot_observed_at":"2026-08-07T12:31:35.681076Z","title":"High- quality, low-delay music coding in the opus codec,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:35.681076Z"},"links":{"cited_paper":"/paper/1602.04845","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:5f6c610a25b3a7640ea0b331e0a71c79e28c433b809b6f108caee2c6f4025327","observation_id":"17f135f1-b807-45a2-81d2-a368ca5eecab","resolution":{"observed_at":"2026-08-07T12:31:35.681076Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:42.687463Z","title":"Review of methods for coding of speech sig- nals,","venue":null,"work_id":"d0270488-60d8-4f89-9e1a-198628f93f6e","year":2023},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:35.789096Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:15e75cc76fe266be85778f8d09c48515c793a88ba3f1a6a5990b5e272202cdaa","observation_id":"c8b41d1a-e60d-46a8-b560-d7417d22742a","resolution":{"observed_at":"2026-08-07T12:31:42.767134Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:42.532510Z","title":"Soundstream: An end-to-end neural au- dio codec,","venue":null,"work_id":"0a797a5a-2522-46e4-a3d0-349ee714584d","year":2021},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:35.905788Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:cdd608605cfa8933a19ebe11a99eaa3e532e2da06357a7d164c1f2bda4da5d1b","observation_id":"efceb1ba-df7d-4a86-8699-1fa0a8bc8372","resolution":{"observed_at":"2026-08-07T12:31:42.594803Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.13438","last_updated":"2022-10-24T17:52:02Z","snapshot_observed_at":"2026-08-03T16:47:47.192907Z","submitted_at":"2022-10-24T17:52:02Z","title":"High Fidelity Neural Audio Compression","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.13438","snapshot_observed_at":"2026-08-07T12:31:35.995896Z","title":"High fidelity neural audio compression,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:35.995896Z"},"links":{"cited_paper":"/paper/2210.13438","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:6a4eed63bc4521c5ab357b08e3e7c017922e5cc6b847440655d9822f951add02","observation_id":"a946eb23-00c5-48a5-a75a-bf3493c4e97b","resolution":{"observed_at":"2026-08-07T12:31:35.995896Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:36.074333Z","title":"High-fidelity audio compression with improved rvqgan,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:36.074333Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:618119cfe8ec9a33f17828849b329a98d8e387b568e6eb48440cb3597a818491","observation_id":"56a079a7-2e81-4d80-8528-3e66bfc0b8f6","resolution":{"observed_at":"2026-08-07T12:31:36.074333Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:42.347757Z","title":"Speechtokenizer: Unified speech tokenizer for speech language models,","venue":null,"work_id":"959e3d4c-1ccd-486b-a847-4c1ecf74dfa2","year":2023},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:36.200999Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:38e98343cbae3f309e9520b6364fda5f53192894ea539e82566241c7786f7ff6","observation_id":"82dc9378-00f3-4afb-afb6-452e4ae0dd63","resolution":{"observed_at":"2026-08-07T12:31:42.419818Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.02765","last_updated":"2023-05-07T09:22:04Z","snapshot_observed_at":"2026-08-06T10:45:04.364170Z","submitted_at":"2023-05-04T12:11:13Z","title":"HiFi-Codec: Group-residual Vector quantization for High Fidelity Audio Codec","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.02765","snapshot_observed_at":"2026-08-07T12:31:36.315885Z","title":"Hifi- codec: Group-residual vector quantization for high fidelity audio codec,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:36.315885Z"},"links":{"cited_paper":"/paper/2305.02765","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:820898d3ebd1ee3e2835aa732455858572de32e40a88c683c409e429f43abd46","observation_id":"314333f8-cf24-4d99-ab6f-6c3aa6839524","resolution":{"observed_at":"2026-08-07T12:31:36.315885Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:42.167276Z","title":"Funcodec: A funda- mental, reproducible and integrable open-source toolkit for neural speech codec,","venue":null,"work_id":"7e8e40b2-c3de-4040-9ecd-9a488608ad48","year":2024},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:36.396118Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:04a216b4781f6135d98f4087e5e2e652eebefa5efd02021e398436a5b1a2ad64","observation_id":"0ee52ea8-31f0-402c-b6bd-72652b9ac12a","resolution":{"observed_at":"2026-08-07T12:31:42.236405Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10533","last_updated":"2024-09-24T01:14:07Z","snapshot_observed_at":"2026-07-06T17:31:04.675896Z","submitted_at":"2024-02-16T09:38:16Z","title":"APCodec: A Neural Audio Codec with Parallel Amplitude and Phase Spectrum Encoding and Decoding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.10533","snapshot_observed_at":"2026-08-07T12:31:36.482077Z","title":"Apcodec: A neural audio codec with parallel amplitude and phase spectrum encoding and decoding,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:36.482077Z"},"links":{"cited_paper":"/paper/2402.10533","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:b9037359a9a983c607048338ebb436291cce59ff49ff08dc6aef84bd4499ec4a","observation_id":"4b657165-07a3-433e-869b-79f756f43896","resolution":{"observed_at":"2026-08-07T12:31:36.482077Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.19842","last_updated":"2024-11-29T16:58:02Z","snapshot_observed_at":"2026-07-06T19:59:01.756640Z","submitted_at":"2024-11-29T16:58:02Z","title":"Scaling Transformers for Low-Bitrate High-Quality Speech Coding","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.19842","snapshot_observed_at":"2026-08-07T12:31:36.571687Z","title":"Scaling transformers for low-bitrate high-quality speech coding,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:36.571687Z"},"links":{"cited_paper":"/paper/2411.19842","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:bae255ce741746df685b55a1447f053f7fa88452c6cda255b99334b84b5863db","observation_id":"8d38116b-38ef-441e-a007-d4492b2d1929","resolution":{"observed_at":"2026-08-07T12:31:36.571687Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.00233","last_updated":"2024-11-28T12:31:04Z","snapshot_observed_at":"2026-07-06T18:08:00.410061Z","submitted_at":"2024-04-30T22:51:36Z","title":"SemantiCodec: An Ultra Low Bitrate Semantic Audio Codec for General Sound","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.00233","snapshot_observed_at":"2026-08-07T12:31:36.650499Z","title":"Semanticodec: An ultra low bitrate semantic audio codec for general sound,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:36.650499Z"},"links":{"cited_paper":"/paper/2405.00233","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:d9384cf0455ad3531afcb6ad648da34a88fd12f929cc75baff8f94b77d2e1083","observation_id":"ce31894c-bcdf-4b26-af0c-1d88a1a15bff","resolution":{"observed_at":"2026-08-07T12:31:36.650499Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:41.956713Z","title":"Autoregres- sive image generation using residual quantization,","venue":null,"work_id":"b077cc5f-493d-4dbf-a917-0ff24d573c29","year":2022},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:36.737380Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:2058e4e1c40c0069eabe99c40caa2b41b3db218481e30308aaa7bacffcc5e292","observation_id":"0b8eb200-e0b4-492c-835d-a9b5e751a96d","resolution":{"observed_at":"2026-08-07T12:31:42.049443Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:36.832036Z","title":"Recent advances in discrete speech tokens: A review,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:36.832036Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:d7f5987b1b0d4a0d96afffa4b55134d7ee7e2db906f737dfe622f30fea12f24d","observation_id":"fa763c0d-f017-4990-be7b-ab30ded36ec7","resolution":{"observed_at":"2026-08-07T12:31:36.832036Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.12792","last_updated":"2023-09-22T01:44:46Z","snapshot_observed_at":"2026-08-04T07:24:03.964437Z","submitted_at":"2023-08-24T13:47:16Z","title":"Sparks of Large Audio Models: A Survey and Outlook","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.12792","snapshot_observed_at":"2026-08-07T12:31:36.944186Z","title":"Sparks of large audio models: A survey and out- look,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:36.944186Z"},"links":{"cited_paper":"/paper/2308.12792","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:1c18679a9ac59f385ee0f0a1d79d68d36b5290ce033767fe9ea5f2761991cdf7","observation_id":"25ba0e21-4489-4a94-aef1-40529f76f3ca","resolution":{"observed_at":"2026-08-07T12:31:36.944186Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:37.026469Z","title":"A survey on speech large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:37.026469Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:6d6a7a3129f503ca450dee24ab8a5a64be1873458f43a462927f7877eee67b5e","observation_id":"5d71bcdf-fade-4634-ad4b-7a54382f75eb","resolution":{"observed_at":"2026-08-07T12:31:37.026469Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.02111","last_updated":"2023-01-05T15:37:15Z","snapshot_observed_at":"2026-08-07T10:11:17.796562Z","submitted_at":"2023-01-05T15:37:15Z","title":"Neural Codec Language Models are Zero-Shot Text to Speech Synthesizers","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.02111","snapshot_observed_at":"2026-08-07T12:31:37.116039Z","title":"Neural codec language models are zero-shot text to speech synthesizers,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:37.116039Z"},"links":{"cited_paper":"/paper/2301.02111","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:d5e853f36f22ce061464bb7975dffe4dd52f6d2f69505d8015b1ff4fd3f45158","observation_id":"0c8ee0e7-771f-415a-a6e0-7b0342511fc1","resolution":{"observed_at":"2026-08-07T12:31:37.116039Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.16107","last_updated":"2023-05-25T14:39:47Z","snapshot_observed_at":"2026-08-06T20:26:21.695165Z","submitted_at":"2023-05-25T14:39:47Z","title":"VioLA: Unified Codec Language Models for Speech Recognition, Synthesis, and Translation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.16107","snapshot_observed_at":"2026-08-07T12:31:37.177571Z","title":"Viola: Unified codec language models for speech recognition, synthesis, and translation,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:37.177571Z"},"links":{"cited_paper":"/paper/2305.16107","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:241b7cf1d75fa70158cc94d7f4a07c7e302ba7d64b988f9fd0ad99968a4966d4","observation_id":"d051f26a-f4d5-4dc1-a1a0-71add2d27203","resolution":{"observed_at":"2026-08-07T12:31:37.177571Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:41.775109Z","title":"Speechx: Neural codec language model as a versatile speech transformer,","venue":null,"work_id":"3092cbcd-85c3-4b3a-b20d-9f094f1a5750","year":2024},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:37.265455Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:7c629c57245b8b732f25d548e81a237d7a9235e1d704d04cfc5e286441071920","observation_id":"a685d257-0a88-439e-8e9e-d74966e9582b","resolution":{"observed_at":"2026-08-07T12:31:41.842452Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.16973","last_updated":"2024-06-14T00:29:46Z","snapshot_observed_at":"2026-07-06T17:50:18.017647Z","submitted_at":"2024-03-25T17:38:32Z","title":"VoiceCraft: Zero-Shot Speech Editing and Text-to-Speech in the Wild","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.16973","snapshot_observed_at":"2026-08-07T12:31:37.381685Z","title":"V oicecraft: Zero-shot speech editing and text-to-speech in the wild,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:37.381685Z"},"links":{"cited_paper":"/paper/2403.16973","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:75d0b3e4e122ed772a0d9fc264aa52dff11e441bb52fb40022ed2be23ebb5e74","observation_id":"364e8b81-d0c9-4f55-a7cb-2cdb826ba4b0","resolution":{"observed_at":"2026-08-07T12:31:37.381685Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.00037","last_updated":"2024-10-02T09:11:45Z","snapshot_observed_at":"2026-07-30T10:21:14.474746Z","submitted_at":"2024-09-17T17:55:39Z","title":"Moshi: a speech-text foundation model for real-time dialogue","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.00037","snapshot_observed_at":"2026-08-07T12:31:37.477873Z","title":"Moshi: a speech-text foundation model for real-time dialogue,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:37.477873Z"},"links":{"cited_paper":"/paper/2410.00037","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:be2f7d4559d5614accc0aff2a866dd8d36de8875513bc310626e781d355bf7e8","observation_id":"167bcdb1-45ee-4618-9db4-07a1d78c3c6b","resolution":{"observed_at":"2026-08-07T12:31:37.477873Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12226","last_updated":"2025-09-08T07:04:17Z","snapshot_observed_at":"2026-07-06T17:32:15.447061Z","submitted_at":"2024-02-19T15:33:10Z","title":"AnyGPT: Unified Multimodal LLM with Discrete Sequence Modeling","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12226","snapshot_observed_at":"2026-08-07T12:31:37.589253Z","title":"Anygpt: Unified multimodal llm with discrete se- quence modeling,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:37.589253Z"},"links":{"cited_paper":"/paper/2402.12226","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:55a8ffda202597eb91b83965154d37ec04307d9f4157bb80d77d9113d01863e0","observation_id":"84fd0d6f-d564-4b80-a7a3-2e87d8697218","resolution":{"observed_at":"2026-08-07T12:31:37.589253Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:41.530493Z","title":"Discrete audio representation as an alter- native to mel-spectrograms for speaker and speech recognition,","venue":null,"work_id":"c7a73679-7e96-45ab-9d14-e143db82447c","year":2024},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:37.713642Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:874454e3f4d9d22177ac7bdba38f1b5f5ce0fec220806387bfd52149b8d0e19b","observation_id":"70a27bac-e8a1-4cc6-a4cd-728d54ff315a","resolution":{"observed_at":"2026-08-07T12:31:41.636060Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.03495","last_updated":"2024-07-03T20:51:41Z","snapshot_observed_at":"2026-07-06T18:41:13.986887Z","submitted_at":"2024-07-03T20:51:41Z","title":"Codec-ASR: Training Performant Automatic Speech Recognition Systems with Discrete Speech Representations","version":1},"cited_work":{"arxiv_id":"2407.03495","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.03495","snapshot_observed_at":"2026-08-07T12:31:39.989072Z","title":"Codec-ASR: Training Performant Automatic Speech Recognition Systems with Discrete Speech Representations","venue":"eess.AS","work_id":"f66b32a5-8a49-42fa-b0e9-4601d2b3c756","year":2024},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:37.863241Z"},"links":{"cited_paper":"/paper/2407.03495","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:20a9936afd207f9f2277bb7a2cc91be46ad18a49da79120240a0a8813ebfa172","observation_id":"d5e6c052-e71b-4012-b8d9-7dcacf1bdc9b","resolution":{"observed_at":"2026-08-07T12:31:40.077465Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.13071","last_updated":"2024-09-18T12:02:47Z","snapshot_observed_at":"2026-08-06T09:11:48.466453Z","submitted_at":"2024-02-20T15:13:38Z","title":"Codec-SUPERB: An In-Depth Analysis of Sound Codec Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.13071","snapshot_observed_at":"2026-08-07T12:31:37.992844Z","title":"Codec-superb: An in-depth analysis of sound codec models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:37.992844Z"},"links":{"cited_paper":"/paper/2402.13071","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:e5997df28a084ce23ee2c9c870af86e96311886fbde8c29218f576c5a4f7c946","observation_id":"cfb51a57-77e6-43be-aa61-f45e5f97fb94","resolution":{"observed_at":"2026-08-07T12:31:37.992844Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.14294","last_updated":"2026-04-21T15:27:17Z","snapshot_observed_at":"2026-07-30T17:14:02.349402Z","submitted_at":"2024-06-20T13:23:27Z","title":"DASB - Discrete Audio and Speech Benchmark","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.14294","snapshot_observed_at":"2026-08-07T12:31:38.099625Z","title":"Dasb–discrete audio and speech benchmark,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:38.099625Z"},"links":{"cited_paper":"/paper/2406.14294","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:cf20a11a035b79b574679ba62a6b01ab223aaaace4c9157cadb3aa7b461d9f8e","observation_id":"3caaa792-47b2-4697-a4eb-8c70a9964f82","resolution":{"observed_at":"2026-08-07T12:31:38.099625Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:41.325636Z","title":"Espnet-codec: Comprehensive training and evalua- tion of neural codecs for audio, music, and speech,","venue":null,"work_id":"ba00c75d-6a1b-434e-a957-dbfb4c0e63ad","year":null},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:38.195522Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:77bd0322675ab7fd481f77662dcc48dbeae6460df7639288614d17ecb91ab6e2","observation_id":"6a885727-6800-44d8-a91a-a2f69b8776a5","resolution":{"observed_at":"2026-08-07T12:31:41.414713Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:40.659430Z","title":"Perceptual evaluation of speech quality (pesq)- a new method for speech quality assessment of telephone net- works and codecs,","venue":null,"work_id":"7be44072-6e5d-434f-985f-2c088c67ba12","year":2001},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:39.258093Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:5369505a340873e210f44cde9a6e5d3b1007627d6ff1848ba6ed3e2e1cdc5635","observation_id":"902dfcdf-92fa-4457-88db-e3df0dac5a96","resolution":{"observed_at":"2026-08-07T12:31:40.728082Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:38.506178Z","title":"Lib- rispeech: an asr corpus based on public domain audio books,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:38.506178Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:fce6f93a900b2db05c23a3da1d328657d1f2f9841d81ac22b257b451990185e7","observation_id":"b31b17fd-b76d-4bee-96a0-98d6c2c6ea91","resolution":{"observed_at":"2026-08-07T12:31:38.506178Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:38.632101Z","title":"The ryerson audio-visual database of emotional speech and song (ravdess): A dynamic, multimodal set of facial and vocal expressions in north american english,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:38.632101Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:829743df3a07298e5968eeefc1c893aa2f2d3402e04db199c1448cd2c756add5","observation_id":"57163fc6-9cc0-4b06-b671-3f0603db4395","resolution":{"observed_at":"2026-08-07T12:31:38.632101Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2009.02095","last_updated":"2020-10-01T17:26:35Z","snapshot_observed_at":"2026-07-06T09:52:52.148226Z","submitted_at":"2020-09-04T10:22:43Z","title":"SEANet: A Multi-modal Speech Enhancement Network","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.02095","snapshot_observed_at":"2026-08-07T12:31:38.772273Z","title":"Seanet: A multi-modal speech enhancement network,","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:38.772273Z"},"links":{"cited_paper":"/paper/2009.02095","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:e779b97e22362b78606be85791766c1dd5acf838e0d645101c73481db8190ffc","observation_id":"0b34c496-5bfa-4938-9f58-2861210d8665","resolution":{"observed_at":"2026-08-07T12:31:38.772273Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:38.859389Z","title":"Long short-term memory,","venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:38.859389Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:44eeb382511922ba722b87ba9b09c8a8816c3aa1371b528faf3114feee03f25f","observation_id":"e9f67569-9931-49b3-8f4a-54573173683c","resolution":{"observed_at":"2026-08-07T12:31:38.859389Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:41.005430Z","title":"Wavlm: Large-scale self- supervised pre-training for full stack speech processing,","venue":null,"work_id":"a14dd269-cfb2-45aa-86a2-adc463964c17","year":2022},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:38.941302Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:becee8369af53a4ac350b656652e7d4c64692590b334f93c847e0277f2ce5f25","observation_id":"d3de5ff9-f977-4155-be71-2606cffc9cb9","resolution":{"observed_at":"2026-08-07T12:31:41.096553Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1904.02882","last_updated":"2019-04-05T06:05:00Z","snapshot_observed_at":"2026-08-07T13:32:24.356336Z","submitted_at":"2019-04-05T06:05:00Z","title":"LibriTTS: A Corpus Derived from LibriSpeech for Text-to-Speech","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1904.02882","snapshot_observed_at":"2026-08-07T12:31:39.059469Z","title":"Libritts: A corpus derived from librispeech for text-to-speech,","venue":null,"work_id":null,"year":1904},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:39.059469Z"},"links":{"cited_paper":"/paper/1904.02882","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:21a6322a90d9a7c8b1225f885fcc6e802d227afc400fe1edb5989deca9e80032","observation_id":"bfbe0925-bfd6-4e58-9b21-2a813bd359e8","resolution":{"observed_at":"2026-08-07T12:31:39.059469Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:40.857520Z","title":"Mel-cepstral distance measure for objective speech quality assessment,","venue":null,"work_id":"a7bbcdc7-a1b3-4f2a-bc68-6ff31889e2e3","year":1993},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:39.150975Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:5429461a1048af804b84a733f2fd929d45c807dace4435cc82ad3b704ba89dd9","observation_id":"150170b4-2ff3-4407-9ea0-6af3af8ac80c","resolution":{"observed_at":"2026-08-07T12:31:40.927826Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:39.342757Z","title":"Robust speech recognition via large-scale weak supervision,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:39.342757Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:9a42c451159d5c8e2972bc375c5428b354c4a5783eea159f3b7f912e8530eb25","observation_id":"a57e1d6f-431a-42fc-96e4-2eb755eb01f1","resolution":{"observed_at":"2026-08-07T12:31:39.342757Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2005.07143","last_updated":"2020-08-10T13:50:24Z","snapshot_observed_at":"2026-08-07T15:21:18.262728Z","submitted_at":"2020-05-14T17:02:15Z","title":"ECAPA-TDNN: Emphasized Channel Attention, Propagation and Aggregation in TDNN Based Speaker Verification","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2005.07143","snapshot_observed_at":"2026-08-07T12:31:39.430961Z","title":"Ecapa- tdnn: Emphasized channel attention, propagation and ag- gregation in tdnn based speaker verification,","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:39.430961Z"},"links":{"cited_paper":"/paper/2005.07143","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:86dd7c11cdda464d11cfd30972467067a38a0d451921895e7fcd6ec9f68b216f","observation_id":"9f765aba-dcf3-4379-955d-c7976a02072f","resolution":{"observed_at":"2026-08-07T12:31:39.430961Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15185","last_updated":"2023-12-23T07:46:55Z","snapshot_observed_at":"2026-08-08T00:25:44.100410Z","submitted_at":"2023-12-23T07:46:55Z","title":"emotion2vec: Self-Supervised Pre-Training for Speech Emotion Representation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.15185","snapshot_observed_at":"2026-08-07T12:31:39.523254Z","title":"emotion2vec: Self-supervised pre-training for speech emotion representation,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:39.523254Z"},"links":{"cited_paper":"/paper/2312.15185","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:c964e1b1fa2d59d5a4625ab211443e3e2571eb28883b0d709a870d3ca4c6a7a1","observation_id":"1283c4a1-509d-4616-a396-e1175b8febbd","resolution":{"observed_at":"2026-08-07T12:31:39.523254Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1907.01160","last_updated":"2019-07-02T04:27:55Z","snapshot_observed_at":"2026-07-06T08:04:17.909965Z","submitted_at":"2019-07-02T04:27:55Z","title":"WHAM!: Extending Speech Separation to Noisy Environments","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1907.01160","snapshot_observed_at":"2026-08-07T12:31:39.627183Z","title":"Wham!: Extend- ing speech separation to noisy environments,","venue":null,"work_id":null,"year":1907},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:39.627183Z"},"links":{"cited_paper":"/paper/1907.01160","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:e2df2ca0c0ea5c0e691c394e910000b0360c14a6e6edd2bc56c3f1b8145c990c","observation_id":"2415ad9f-7ab7-421b-9435-e5469f932cc4","resolution":{"observed_at":"2026-08-07T12:31:39.627183Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2005.13981","last_updated":"2020-10-18T04:36:21Z","snapshot_observed_at":"2026-07-31T23:56:07.410341Z","submitted_at":"2020-05-16T23:48:37Z","title":"The INTERSPEECH 2020 Deep Noise Suppression Challenge: Datasets, Subjective Testing Framework, and Challenge Results","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2005.13981","snapshot_observed_at":"2026-08-07T12:31:39.694445Z","title":"The interspeech 2020 deep noise suppression challenge: Datasets, subjective testing framework, and challenge results,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:39.694445Z"},"links":{"cited_paper":"/paper/2005.13981","citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:f0604dbaadce2630d255b32e513608c823e07ed3b5b05b041efd011a992be545","observation_id":"6bb3dce6-761b-408c-a565-cd752132f230","resolution":{"observed_at":"2026-08-07T12:31:39.694445Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:31:41.200623Z","title":null,"venue":null,"work_id":"24dd2fb6-6b76-45fe-9539-166fa5a716f7","year":2024},"citing_paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-07T12:31:38.381810Z"},"links":{"citing_paper":"/paper/2505.24248"},"observation_digest":"sha256:f54cd97b4e44d14da56b822d2ef848afb4e7a8d6c97109ae8f305572c1578491","observation_id":"7ad1a252-d32c-41fb-8ec1-2a9342d8e997","resolution":{"observed_at":"2026-08-07T12:31:41.243873Z","resolver_source":"raw_fallback","status":"parse_uncertain"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2505.24248","last_updated":"2025-05-30T06:12:52Z","latest_version":1,"primary_category":"eess.AS","snapshot_observed_at":"2026-08-08T04:48:33.525066Z","submitted_at":"2025-05-30T06:12:52Z","title":"Probing the Robustness Properties of Neural Speech Codecs"},"reference_resolution":{"displayed":49,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":1,"unresolved":29,"verified_exact":2,"verified_fuzzy":17},"total_outbound_references":49},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 49 of 49 outbound references and 1 inbound Pith citation observation for arXiv:2505.24248."}