{"as_of":"2026-08-08T14:54:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:545c77d7478b5c335fae418c06ef744a7ed186f6986d590a6add8f0c4135f27a","coverage":[{"denominator":55,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":55,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T13:43:11.818073Z","state":"measured"},{"denominator":56,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":56,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T13:43:07.380177Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-07T13:43:11.971534Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"cited_work":{"arxiv_id":"2505.21237","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.21237","snapshot_observed_at":"2026-08-07T13:43:11.971534Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","venue":"cs.SD","work_id":"7b522658-188f-4ea0-808c-e133f6cc439c","year":2025},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:07.380177Z"},"links":{"cited_paper":"/paper/2505.21237","citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:444beed871d5b51db7e8ec07b62c9199e5d84c1c315bdf100339c9f9653fd3d8","observation_id":"a01e73e8-971b-4289-ba77-49917627baa8","resolution":{"observed_at":"2026-08-07T13:43:12.052007Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2505.21237/citation-record","integrity":"/paper/2505.21237/integrity","json":"/paper/2505.21237/citation-record.json","paper":"/paper/2505.21237"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"cited_work":{"arxiv_id":"2505.21237","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.21237","snapshot_observed_at":"2026-08-07T13:43:11.971534Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","venue":"cs.SD","work_id":"7b522658-188f-4ea0-808c-e133f6cc439c","year":2025},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:07.380177Z"},"links":{"cited_paper":"/paper/2505.21237","citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:444beed871d5b51db7e8ec07b62c9199e5d84c1c315bdf100339c9f9653fd3d8","observation_id":"a01e73e8-971b-4289-ba77-49917627baa8","resolution":{"observed_at":"2026-08-07T13:43:12.052007Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:21.865421Z","title":null,"venue":null,"work_id":"fa9809c6-3c03-4caf-a112-32faf1727af9","year":null},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:07.422727Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:403783f812b7b9fe55b0cfd061c15a994973b5840d69b32b1000c8cf04fdb4dd","observation_id":"e8216090-5335-4d83-8319-5e46171442fb","resolution":{"observed_at":"2026-08-07T13:43:21.931960Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:21.694637Z","title":"layers” and “blocks","venue":null,"work_id":"4a0563bc-24a4-44c7-8f61-28e65da6e0dd","year":null},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:07.488025Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:3b61ea8c45979c94f116d96d74b83d802d027eb06d08bcd01cfc086df6cd37ba","observation_id":"04ca0274-f973-4fbb-a149-601eb63d885e","resolution":{"observed_at":"2026-08-07T13:43:21.777172Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:21.517434Z","title":"Experimental Setup Models and Data.For supervised learning, we take Conformer AED system configured with the ESPnet [45] recipe2","venue":null,"work_id":"bc47b1f3-47db-4099-8ef3-37480eb891ad","year":null},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:07.552904Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:154c6f49a806216db09b7c09257489fa72e29137d3dd49cfc2d5162c5e43e27d","observation_id":"bc15c7df-9463-4229-8440-51046a599862","resolution":{"observed_at":"2026-08-07T13:43:21.600217Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:21.186494Z","title":null,"venue":null,"work_id":"25a5ce89-97a2-4a59-936c-50596ed88307","year":null},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:07.636975Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:6f9d553d57871e6501e201afe2b7e1dab3b65dd3fa5d8211129c3ae86302aa32","observation_id":"4fa7592e-8a24-4a51-80aa-8cf15a16c188","resolution":{"observed_at":"2026-08-07T13:43:21.271346Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:21.085552Z","title":"14200220, 14200021, 14200324, Innovation Technology Fund grant No","venue":null,"work_id":"24f4d035-5430-47b9-8ea9-b2f13caf5776","year":null},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:07.689427Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:77a04222b8f231be15b6a8cdc58abb61bfd3d728da162d23d32e9f66390b3fd1","observation_id":"6715130b-38a3-41be-9bf1-f435b6e1cae8","resolution":{"observed_at":"2026-08-07T13:43:21.135289Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:19.885661Z","title":"2-bit conformer quantization for automatic speech recognition,","venue":null,"work_id":"6ccdd84f-16e3-4451-9eeb-01c975ea0c19","year":2023},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:08.196038Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:8960c00f29d2f8098e3445c4df6d1273bcf92356e97685f66235fdb3fa384e2d","observation_id":"d897bbf7-c59a-4fc6-8e58-48a795dff800","resolution":{"observed_at":"2026-08-07T13:43:19.972606Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:21.350522Z","title":"2 and 3), and the 8-layer system that trained at the unfolded depth of 12 that can not be folded back (i.e., Sys","venue":null,"work_id":"834f9248-090a-46d7-8882-1f2e0104e028","year":null},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:07.594177Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:24d8eadab4f15bb31dbe724a694645a97ef77d0260c45abf4a334a49c3e19de9","observation_id":"19ee9b6d-04ab-4134-bf12-eb8638c5e798","resolution":{"observed_at":"2026-08-07T13:43:21.405275Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:20.940309Z","title":"wav2vec 2.0: A framework for self- supervised learning of speech representations,","venue":null,"work_id":"cda4e15a-b791-4027-84f6-dc2e9fb08a54","year":2020},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:07.767612Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:24960d9065644758052ca7ce9920534c4c73f80c70986cd650d2f1e00431fb36","observation_id":"c1b064d0-5332-4ef2-96ee-b3c3304d02c9","resolution":{"observed_at":"2026-08-07T13:43:21.005214Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:20.811173Z","title":"HuBERT: Self-supervised speech representation learning by masked prediction of hidden units,","venue":null,"work_id":"c1cb7502-a61a-471c-aaee-6809de06cb9c","year":2021},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:07.836164Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:2aded99b913eace1302b04861df4d902d7ff6d930b4b3487828c5cacd8a4b9c7","observation_id":"dacc2fed-c0d5-44d4-a08c-c941a5027e9b","resolution":{"observed_at":"2026-08-07T13:43:20.876500Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:20.591211Z","title":"Data2vec: A general framework for self-supervised learning in speech, vision and language,","venue":null,"work_id":"6af59a49-f5c0-406e-aad2-624d3646f012","year":2022},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:07.906516Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:c1038c85d1106294d2a9b4e48a0171285bdbe7b9a9bc21800548619cab83ac78","observation_id":"9d77b614-38e7-470b-8a2a-0af15030a474","resolution":{"observed_at":"2026-08-07T13:43:20.709167Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:20.436405Z","title":"WavLM: Large-scale self-supervised pre-training for full stack speech processing,","venue":null,"work_id":"23739b5f-0c7c-4127-b983-87e6c1755797","year":2022},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:07.984689Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:4fa60341f6d1f9781c557b350e9f531197c80b9f532318be902c8ef08634a0cd","observation_id":"e2bb097f-33a2-464a-8de2-ba22feb3b4fe","resolution":{"observed_at":"2026-08-07T13:43:20.524588Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:20.262377Z","title":"4-bit conformer with native quan- tization aware training for speech recognition,","venue":null,"work_id":"9ba35b07-699e-4532-9209-94f8c4d04a46","year":2022},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:08.061213Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:fa4c291a8921b08e04056ae6a49a84fa2c8db4464be183b01f80e0e8b90255e4","observation_id":"9dbdde20-340c-47ae-826e-4fbbf3ee65ff","resolution":{"observed_at":"2026-08-07T13:43:20.337788Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:20.068173Z","title":"Integer-only zero-shot quantization for efficient speech recognition,","venue":null,"work_id":"0b7e8722-139a-42a8-88e9-287d547b7694","year":2022},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:08.129326Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:fe9ad54fead2c02058c81b44e2d6715ba5915e2b1aad24dbb8df83772bf87420","observation_id":"c0d14317-719a-43ac-85a9-6432039bfaea","resolution":{"observed_at":"2026-08-07T13:43:20.145116Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:19.732601Z","title":"Effective and efficient mixed precision quan- tization of speech foundation models,","venue":null,"work_id":"a81e5bbf-4d04-4234-9910-e96d30eb4e37","year":2025},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:08.268580Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:b1c5d9ee63241da18cdce0b49c8b86526ec97b37bd436c5570ebb5f69a29e546","observation_id":"ef536f06-bcd0-4c70-adf9-26ed315da8b1","resolution":{"observed_at":"2026-08-07T13:43:19.797576Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:19.569761Z","title":"Multi-stage progressive com- pression of conformer transducer for on-device speech recogni- tion,","venue":null,"work_id":"3b0f9a21-0fcc-4336-92b4-92742c3d6ab2","year":2022},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:08.342033Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:4477bff70c3219a37ae6c76f1df2aae85a2be343161bc00e6cbc6892858f1df6","observation_id":"6d5209bc-e6ae-430f-9414-4baef42eccec","resolution":{"observed_at":"2026-08-07T13:43:19.652190Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:19.423918Z","title":"DistilHuBERT: Speech represen- tation learning by layer-wise distillation of hidden-unit bert,","venue":null,"work_id":"ce13d19a-a433-41d7-a4a7-6b33be005749","year":2022},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:08.446336Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:dc225210502a664b404f8c4565d20aa186204a037b2901dee8b141f47b85d04c","observation_id":"0f7550c7-1ad4-4e2f-b85c-44054cafcafb","resolution":{"observed_at":"2026-08-07T13:43:19.490554Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:19.249560Z","title":"FitHuBERT: Going thinner and deeper for knowledge distillation of speech self-supervised learning,","venue":null,"work_id":"ce5db3fc-8b28-40f3-91de-3fa8e2cf4544","year":2022},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:08.517020Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:f9751eb600e616b5f65d96d2b16fd7b5bfbd64ce628db3e59f052778f26589a3","observation_id":"af473499-aff0-4bf7-b9db-09305326defa","resolution":{"observed_at":"2026-08-07T13:43:19.343594Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:19.076350Z","title":"Deep versus wide: An analysis of student architectures for task-agnostic knowledge distillation of self-supervised speech models,","venue":null,"work_id":"f6fbfd3f-732b-4abc-817d-a1063beb8ce4","year":2022},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:08.586275Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:ca9980440f83a61cdf4bd88b49c39b97867a7ea272b7af707850d70ae54d076e","observation_id":"c11cc13c-c2b7-4506-a5c2-1c4bab1fc1c8","resolution":{"observed_at":"2026-08-07T13:43:19.144847Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.09278","last_updated":"2023-03-16T12:59:17Z","snapshot_observed_at":"2026-08-04T00:20:40.261026Z","submitted_at":"2023-03-16T12:59:17Z","title":"DistillW2V2: A Small and Streaming Wav2vec 2.0 Based ASR Model","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.09278","snapshot_observed_at":"2026-08-07T13:43:08.666942Z","title":"DistillW2V2: A small and streaming wav2vec 2.0 based asr model,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:08.666942Z"},"links":{"cited_paper":"/paper/2303.09278","citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:b38b34a398184ab5f7dd7246c53b9154f0a30027337ff594e0f5958d895c3462","observation_id":"1700499c-873e-4142-85f3-a4d25961bcbb","resolution":{"observed_at":"2026-08-07T13:43:08.666942Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:18.921714Z","title":"Distilling HuBERT with LSTMs via decoupled knowledge distillation,","venue":null,"work_id":"f7a099c5-49fc-45da-8acb-b2497ba37e49","year":2024},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:08.679927Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:085dbee4f9234a5ae6a7e82b55548c5c2e252e8b54948a61f5d01b97e9cd17f5","observation_id":"45225458-cd02-4995-a8a1-f81ed1e40c72","resolution":{"observed_at":"2026-08-07T13:43:19.014413Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:18.734498Z","title":"Conformer-based on-device streaming speech recognition with KD compression and two-pass architec- ture,","venue":null,"work_id":"8d07ca26-625e-4459-98d2-68574d67e74e","year":2023},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:08.800652Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:7a38312c2f4a0d23968125f356bd762633dadd074bb115beb692d5792fd1bca8","observation_id":"6bcd0f5c-197d-43de-82e1-e02a86acd0ee","resolution":{"observed_at":"2026-08-07T13:43:18.822013Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:18.508460Z","title":"Dynamic sparsity neural networks for au- tomatic speech recognition,","venue":null,"work_id":"eec878d4-ac04-45de-9da2-9d9343871a4c","year":2021},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:08.899527Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:c86b8a96f96f83cbf4174bf7089e68955104538e6e5708a48457a7da07981d8b","observation_id":"0f532c98-5325-407a-9be6-0c6cb3f3b3f6","resolution":{"observed_at":"2026-08-07T13:43:18.607389Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:18.313208Z","title":"PARP: Prune, adjust and re-prune for self-supervised speech recognition,","venue":null,"work_id":"d376d96b-71eb-4d07-9457-8def19bfc36c","year":2021},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:08.975950Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:12f9e8518a4429a039758664f2929b7ebe095622cb32c53ba22c4e1af5342656","observation_id":"655fb50f-5a43-4f44-bdf8-657e410785f2","resolution":{"observed_at":"2026-08-07T13:43:18.405328Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:18.142449Z","title":"Layer pruning on demand with intermediate ctc,","venue":null,"work_id":"ad4b25e1-4e8a-46df-8f57-b0e73430b853","year":2021},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:09.048207Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:f1b64388b3cb5444268da3e95d10df9f00f933ae7dfe7a9037fbcb58b943f38a","observation_id":"fce3a570-f306-4c4f-9606-30d1498b6c23","resolution":{"observed_at":"2026-08-07T13:43:18.214466Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:17.975482Z","title":"DPHuBERT: Joint distillation and prun- ing of self-supervised speech models,","venue":null,"work_id":"d0516d42-197c-4036-9fb3-0c41750e84e5","year":2023},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:09.201200Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:bef038b3c670831b539579ec9e77a597f8d742c384a92b7f375ac9f15fe61c24","observation_id":"ed453b9c-462f-4237-a907-0bb010e32512","resolution":{"observed_at":"2026-08-07T13:43:18.054476Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:17.806497Z","title":"Accurate and structured pruning for efficient automatic speech recognition,","venue":null,"work_id":"2fae2201-096a-4c07-944f-090412b3428d","year":2023},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:09.313430Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:68f003835256411cf828ae4b2e92463d2c81896ec0d99fcbe045729ce4ced80d","observation_id":"6a80f6ad-d054-4718-8aae-0a949b599daf","resolution":{"observed_at":"2026-08-07T13:43:17.906178Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:17.634610Z","title":"PADA: Pruning assisted domain adaptation for self-supervised speech representations,","venue":null,"work_id":"06b2380a-e614-4bb9-9355-ac56b4a53b8a","year":2023},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:09.417020Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:2bce76b9fcec06cbab0339fff23574249c4fc4efab0de6893db771b2e80c07b1","observation_id":"fefd0017-24ea-4caa-9826-b891d1be1345","resolution":{"observed_at":"2026-08-07T13:43:17.708458Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:17.400136Z","title":"Structured pruning of self-supervised pre-trained models for speech recognition and understanding,","venue":null,"work_id":"3b1d3acd-aecb-4b6c-b6ce-c98bcbda388b","year":2023},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:09.541148Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:dfca9856b0f53c8c98c48b380364ccd0a98f5278f5bdbdca3b3c188a70e09d23","observation_id":"875650d8-cb07-4942-bafa-8dc7b4e8ea48","resolution":{"observed_at":"2026-08-07T13:43:17.545328Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:17.214783Z","title":"Task-agnostic structured pruning of speech representation models,","venue":null,"work_id":"0f5349a6-dcbe-490d-9ddf-2173caa8248a","year":2023},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:09.649449Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:8c60a1f5ba08ad4a802b3b5a8cea2ff307a4448c3e82678d81e68afe347397ee","observation_id":"7d9d327d-2ca1-4463-a4db-665476dc0bd1","resolution":{"observed_at":"2026-08-07T13:43:17.313917Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:17.003831Z","title":"SparseW A V: Fast and accurate one-shot un- structured pruning for large speech foundation models,","venue":null,"work_id":"f45ac5d1-25da-42e9-b170-daf9c354cdd0","year":2024},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:09.772196Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:c000597d7e0021defa72dc4971432f8880731e1640c8ebc5dcff31fdd1cd76c8","observation_id":"3bda6774-74e5-4f47-9c98-fc6f5831f8ef","resolution":{"observed_at":"2026-08-07T13:43:17.098892Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:16.813227Z","title":"Semi-orthogonal low-rank matrix fac- torization for deep neural networks,","venue":null,"work_id":"aeb730bd-75a4-4402-8c07-91434960f1a2","year":2018},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:09.831632Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:e4f616e319afa380fc9ff42aa44c8bcf9e3b4475f211e2e35984974ac4cd2c86","observation_id":"69733129-8516-4289-b2df-f623fdb8acf4","resolution":{"observed_at":"2026-08-07T13:43:16.921178Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:16.629885Z","title":"Neural architecture search for LF-MMI trained time delay neural networks,","venue":null,"work_id":"c54ad666-ca4b-4125-86f3-a7b45ed32374","year":2022},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:09.875891Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:d69b8bf488e37691e3fe199e0efb77db259dfd7a6d50ff138ce0dacfd30d0770","observation_id":"207ca7ae-df4c-46a8-8026-2a00b2c5f358","resolution":{"observed_at":"2026-08-07T13:43:16.710212Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:16.471725Z","title":"Efficient conformer-based speech recognition with linear attention,","venue":null,"work_id":"2eef5bb9-bffd-4df0-8e4d-a9baa5b81761","year":2021},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:09.975754Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:1df97b8a43d8a43aa7f4be90a30b30555a6f9d3a98be2a6b6da8288e69818b80","observation_id":"00db2f45-c5ed-4e83-8f2c-2f6f878179fc","resolution":{"observed_at":"2026-08-07T13:43:16.548847Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:16.313406Z","title":"Lossless 4-bit quantization of architecture compressed conformer asr systems on the 300-hr switchboard cor- pus,","venue":null,"work_id":"78d12376-1735-4746-9763-873c6aad2ac6","year":2023},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:10.086077Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:3695dde4b05a54920bc904f89a35aa14abfb6820eacdd6f08e1387ffba3cbd18","observation_id":"5e399492-2abe-417b-8ca3-f45ac5948bb3","resolution":{"observed_at":"2026-08-07T13:43:16.383257Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:16.098073Z","title":"ALBERT: A lite bert for self-supervised learning of language representations,","venue":null,"work_id":"74dc2eb5-4d9d-4a43-b8ee-59aa9c3cf9e0","year":2020},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:10.151303Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:9e1a9619c623943ee91c38df999120ebaae6909d2a99fa83ceb51a40ebfd516b","observation_id":"ea5c415f-0991-49bd-9a26-99b0f5f1ecec","resolution":{"observed_at":"2026-08-07T13:43:16.198027Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:15.920123Z","title":"Extremely low footprint end-to-end asr system for smart device,","venue":null,"work_id":"32fb9ef8-48dd-4a80-b451-c7786babe43a","year":2021},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:10.232689Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:6061453759a9b2324ea82974d7cbbcfb0ee5eba7571cb03518fe1d74c4e7c065","observation_id":"212d1c2b-b4b2-430c-a7d7-2deeeffdf259","resolution":{"observed_at":"2026-08-07T13:43:16.023507Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:15.710414Z","title":"Weight-sharing supernet for searching specialized acoustic event classification networks across device constraints,","venue":null,"work_id":"f32bfa93-a793-485f-aaa8-6e54b4fe119d","year":2023},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:10.329505Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:13b31f6f2bdbc4fa382d9705f6e8795e1620d3212ce67a892d7a50524a7ac31d","observation_id":"9b3f8fd7-4e57-4fe6-9507-ca76f74dc3de","resolution":{"observed_at":"2026-08-07T13:43:15.822574Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:15.482224Z","title":"Sharing low rank conformer weights for tiny always-on ambient speech recognition models,","venue":null,"work_id":"db52d258-f326-41bd-b40e-2b6b4620e3e9","year":2023},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:10.420796Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:fc4044d0e0b7cff1682f79f2e502a53f32a579e20b784d84ce8fb9c245833639","observation_id":"baf31a38-c01c-4a99-a3c7-dbd7dfc6602d","resolution":{"observed_at":"2026-08-07T13:43:15.593433Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:15.293590Z","title":"ResidualTransformer: Residual low-rank learning with weight-sharing for transformer layers,","venue":null,"work_id":"21cbde62-32e6-4900-8f30-f3543365535e","year":2024},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:10.526774Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:a652a6029dc50a5ddb0cc1816cf107a3e3dce41729914b8eb5a5f1707e5d8878","observation_id":"7b5fc68d-63df-4c61-8371-05dbbfc6e238","resolution":{"observed_at":"2026-08-07T13:43:15.400104Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:15.079338Z","title":"Once for all: Train one network and spe- cialize it for efficient deployment,","venue":null,"work_id":"2ea94ac1-3705-4ba3-a3e3-e0f4453a3bac","year":2020},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:10.627127Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:973d3b7ab2946a49b683620a8dc4e9031147f3c1a69f9b93d176b68a9e5cd5b9","observation_id":"952352c2-cc15-40f7-ae5d-09782935b63f","resolution":{"observed_at":"2026-08-07T13:43:15.150055Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:14.868179Z","title":"Universally slimmable networks and im- proved training techniques,","venue":null,"work_id":"6f22b436-494a-4cae-941d-29b004ba16c1","year":2019},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:10.698683Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:5c2f65058141af487e3f10a4d7749dbba60f710c9106043405ea163cee5ed380","observation_id":"892d5b50-6082-499b-a623-cc2a0ee53a2f","resolution":{"observed_at":"2026-08-07T13:43:14.968833Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:14.713702Z","title":"Self-distillation: Towards efficient and compact neural networks,","venue":null,"work_id":"711de415-1b9c-46f0-8b59-7165ed64da35","year":2021},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:10.814701Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:c2fb3c7e8d7893947ae0f496c480f6e0df613d23c0617a1af4a8bfedab394265","observation_id":"bf643f84-bf20-46cb-894b-b1d88708889b","resolution":{"observed_at":"2026-08-07T13:43:14.781678Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:14.531268Z","title":"Collaborative training of acoustic en- coders for speech recognition,","venue":null,"work_id":"6e7590b0-3253-439a-b359-380ab31d1e41","year":2021},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:10.922117Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:6730a979861cb4afaab021b43811e83dfe3ac87a5b8cbe2f7e4f0e80258acb93","observation_id":"248ef7cf-cf00-4c26-b8df-68d2e4778aed","resolution":{"observed_at":"2026-08-07T13:43:14.626532Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:14.345413Z","title":"LightHuBERT: Lightweight and con- figurable speech representation learning with once-for-all hidden- unit bert,","venue":null,"work_id":"9fc75fe9-eac9-423d-aadc-eef201e192d6","year":2022},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:11.016488Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:825149a08e2afa3f005267d96d49a654c4ee26d0a6846d20653404b12060534b","observation_id":"d2ebdf17-9c94-43ee-9996-8a67960f0a2a","resolution":{"observed_at":"2026-08-07T13:43:14.447473Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:14.107365Z","title":"One-pass multiple conformer and foundation speech systems compression and quantization using an all-in-one neural model,","venue":null,"work_id":"3b46b1fa-8de9-4087-b42e-87e296b1c08a","year":2024},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:11.108368Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:08131d769a7c33fff6efc73a0977c6680bf807cbea4dd36e00dd996499a504a3","observation_id":"32cb1af8-aab7-48b7-b9e6-86df7896fba3","resolution":{"observed_at":"2026-08-07T13:43:14.257772Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:13.854147Z","title":"Small-footprint slimmable networks for keyword spotting,","venue":null,"work_id":"ee749c5a-a32d-4a81-b2e0-1c9f17dc0eca","year":2023},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:11.212859Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:037adb12bc185d23df74b87fe08e40dda6afb17a3fc74baa8d7b0616f9f86226","observation_id":"74e735e2-aa56-4fc9-869e-2a4625078488","resolution":{"observed_at":"2026-08-07T13:43:14.000088Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:13.549685Z","title":"Conformer: Convolution-augmented transformer for speech recognition,","venue":null,"work_id":"cbb2ced2-06cf-4bec-89e1-dedb439d7062","year":2020},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:11.285371Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:f572f0d9989936e476a548f6c4d71f9b25605be0adfe6a0259fd206b7f73d53d","observation_id":"7e2a922d-55dc-44b6-8fab-daa672eba023","resolution":{"observed_at":"2026-08-07T13:43:13.683737Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:13.202655Z","title":"Hybrid CTC/attention architecture for end-to-end speech recognition,","venue":null,"work_id":"b93e3426-a880-4f25-b8ab-ce084df4ad81","year":2017},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:11.356354Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:4c29d8ddf72d321b08bf501cd6b6b2c6623f085fd84476545b132e573e043be9","observation_id":"157f47f6-e3bd-42f3-8115-acf55a4e457c","resolution":{"observed_at":"2026-08-07T13:43:13.323439Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:12.892713Z","title":"Improving transformer-based speech recognition systems with compressed structure and speech at- tributes augmentation","venue":null,"work_id":"618f8664-a7bc-4b31-8c2e-f931cde86a3c","year":2019},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:11.439126Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:c002c3fdd42a4a8e7a800ccab380fb76b625994d8e4cfdead42669dafd8eae48","observation_id":"fc00da22-5a94-4337-b17d-9ad461cfba72","resolution":{"observed_at":"2026-08-07T13:43:13.034065Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:12.559242Z","title":"Non-autoregressive asr with self-conditioned folded encoders,","venue":null,"work_id":"e0965014-4dbe-4b12-b5e0-b8c52f7211f5","year":2022},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:11.537910Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:85651ec527b05257c52da07c3a568f0986271da1a4a6d3e256d4c99148dd0cf2","observation_id":"ea74fb35-6bd4-4177-ba4b-c64c5093a18e","resolution":{"observed_at":"2026-08-07T13:43:12.720401Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1804.00015","last_updated":"2018-03-30T18:09:39Z","snapshot_observed_at":"2026-07-06T06:31:07.605442Z","submitted_at":"2018-03-30T18:09:39Z","title":"ESPnet: End-to-End Speech Processing Toolkit","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1804.00015","snapshot_observed_at":"2026-08-07T13:43:11.622366Z","title":"Espnet: End-to-end speech process- ing toolkit,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:11.622366Z"},"links":{"cited_paper":"/paper/1804.00015","citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:2c88fc48d8e6200e006b36fa14f8b3db14ca4458ded055ce9b1b51622cf86141","observation_id":"ca7c6e29-cd31-43fa-a06d-6804a150a26c","resolution":{"observed_at":"2026-08-07T13:43:11.622366Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:12.352473Z","title":"SWITCHBOARD: Tele- phone speech corpus for research and development,","venue":null,"work_id":"617517b9-c30e-4ef3-9956-7592bd45d59e","year":1992},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:11.695465Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:84235902fc70a6dc2b2ab02a4d4037b7218fe95edee0df93fdbca20271bbf558","observation_id":"20cbcf8d-b378-45a9-a1c1-9809e1611514","resolution":{"observed_at":"2026-08-07T13:43:12.444631Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:12.190061Z","title":"LibriSpeech: an asr corpus based on public domain audio books,","venue":null,"work_id":"4fa84ad3-3a75-4afb-87d2-e01a36d02c3f","year":2015},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:11.746269Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:1289d6874d4e580351cdb7f23d2132a47a0c10ff0694b94bdce521a422f68529","observation_id":"32db411a-7660-45ac-bd1d-59cd30072332","resolution":{"observed_at":"2026-08-07T13:43:12.251585Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:43:11.818073Z","title":"Some statistical issues in the comparison of speech recognition algorithms,","venue":null,"work_id":null,"year":1989},"citing_paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-07T13:43:11.818073Z"},"links":{"citing_paper":"/paper/2505.21237"},"observation_digest":"sha256:48b8782e87ec17280010900ed16435168d7febda37f7192599d12bcd96d5bfc6","observation_id":"5697be6e-e07b-438a-be74-12a188469a03","resolution":{"observed_at":"2026-08-07T13:43:11.818073Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2505.21237","last_updated":"2025-05-27T14:19:21Z","latest_version":1,"primary_category":"cs.SD","snapshot_observed_at":"2026-08-07T13:30:33.274867Z","submitted_at":"2025-05-27T14:19:21Z","title":"Unfolding A Few Structures for The Many: Memory-Efficient Compression of Conformer and Speech Foundation Models"},"reference_resolution":{"displayed":55,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":5,"verified_exact":1,"verified_fuzzy":48},"total_outbound_references":55},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 55 of 55 outbound references and 1 inbound Pith citation observation for arXiv:2505.21237."}