{"as_of":"2026-08-06T17:05:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:89924b64018416cd285893fc83f76dedb869c673a613f94182ffa95ec67f5dd3","coverage":[{"denominator":44,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":44,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-30T07:26:38.380118Z","state":"measured"},{"denominator":46,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":46,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-06T06:34:29.942622+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-01T16:37:24.239501Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-06-30T07:34:21.772463Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"cited_work":{"arxiv_id":"2606.29534","doi":null,"metadata_source":"pith","pith_arxiv_id":"2606.29534","snapshot_observed_at":"2026-06-30T07:34:21.772463Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","venue":"cs.CL","work_id":"cda87a86-7482-4875-a65f-900700e50aa9","year":2026},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"cited_paper":"/paper/2606.29534","citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:14b555c488eaa0b1dc4450f45ee928d0f254a49cf3764914f4186a304ebb16cb","observation_id":"4e1c45ee-beb8-4513-9b94-6224a6ad7915","resolution":{"observed_at":"2026-06-30T07:34:21.773819Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2606.29534","snapshot_observed_at":"2026-08-01T16:37:24.239501Z","title":"Preference-asr: A preference-aware test set for benchmarking asr in the era of speech llms.arXiv preprint arXiv:2606.29534, 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.26410","last_updated":"2026-07-29T02:42:56Z","snapshot_observed_at":"2026-08-03T19:09:06.764457Z","submitted_at":"2026-07-29T02:42:56Z","title":"Voice Memory for Agentic Speech Recognition","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-01T16:37:24.239501Z"},"links":{"cited_paper":"/paper/2606.29534","citing_paper":"/paper/2607.26410"},"observation_digest":"sha256:9fab040000e07d7ca762e576bdaa712c4f86478b77850f1329f19c067893e7ea","observation_id":"e52b8d79-183e-4cf1-894a-2fcd6af005f9","resolution":{"observed_at":"2026-08-01T16:37:24.239501Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2606.29534/citation-record","integrity":"/paper/2606.29534/integrity","json":"/paper/2606.29534/citation-record.json","paper":"/paper/2606.29534"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"cited_work":{"arxiv_id":"2606.29534","doi":null,"metadata_source":"pith","pith_arxiv_id":"2606.29534","snapshot_observed_at":"2026-06-30T07:34:21.772463Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","venue":"cs.CL","work_id":"cda87a86-7482-4875-a65f-900700e50aa9","year":2026},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"cited_paper":"/paper/2606.29534","citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:14b555c488eaa0b1dc4450f45ee928d0f254a49cf3764914f4186a304ebb16cb","observation_id":"4e1c45ee-beb8-4513-9b94-6224a6ad7915","resolution":{"observed_at":"2026-06-30T07:34:21.773819Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:ef83304a6bbe82d072c0305fb85f8d9bac05e16045906d70636394e546ff3791","observation_id":"f7965e44-9410-49c7-ad1a-d0d74184c56f","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:2d37246f26f16cba72973e3740476fe41efe6b16407e2af69f543264c4dc0f3f","observation_id":"adf18de2-0f98-4b3f-969c-53195acfd7cb","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:5d7d774a63b8c62e1ec95c7c4a3d70ff41b8094d0ef91ae58a093f94fbbfbea6","observation_id":"950b88e3-6f68-43d6-a25b-7278493b6090","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:317b6e2ec496f13cc9d3ab807fe39736b3c720da1d21146287ed5c2f314c54af","observation_id":"a3aeea7d-27cd-4d42-ab3a-cc4ab58f8cf1","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"22”) or spoken words (TN, e.g., “twenty-two","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:a7bd02417f3d2607417ae8e91ef8e6d5595c4127abe74ec97ff646c7c7cd26af","observation_id":"39896001-e407-4d00-b522-8e13104cbd2e","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"22nd” following an ITN instruction, a standard normalizer converts it to “twenty second,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:08307616eeba9bd6ae6f8e6cdcfd0f548775eed8500b661f2982fb834935e3e0","observation_id":"c07e22fe-74d6-4002-a754-c1e98a6569a5","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"Standard","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:3099629e6420b42cc7e9c3a021ab05b2e3e1440db6bdcda6bb2b4f13d118c355","observation_id":"76e307f9-446f-4065-8b15-471d91a92dd9","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:42ce26d56722530627b590c4778f7e11eb5c1a24fdc5815879419bf8fec711fc","observation_id":"ed2e110b-dbf5-4c58-a20d-89f32bbeee53","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"All content was reviewed and validated by the authors","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:8967b0f626ee58597ff8a8fc35cd8e698e65b6768e5d50dafe9ad92fb97327ba","observation_id":"4712ad91-6399-486b-88f0-a4dd83ec1f06","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"SALMONN: Towards generic hearing abilities for large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:24e99a1b0bd87002b5c54ee9d6195e9ad6fc78299ba70d83728e627ef2174507","observation_id":"1f7f1991-29d5-4e81-8569-8ec9eee7042b","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.07919","last_updated":"2023-12-21T10:20:42Z","snapshot_observed_at":"2026-07-06T16:47:12.709738Z","submitted_at":"2023-11-14T05:34:50Z","title":"Qwen-Audio: Advancing Universal Audio Understanding via Unified Large-Scale Audio-Language Models","version":2},"cited_work":{"arxiv_id":"2311.07919","doi":"10.48550/arxiv.2311.07919","metadata_source":"pith","pith_arxiv_id":"2311.07919","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Qwen-Audio: Advancing Universal Audio Understanding via Unified Large-Scale Audio-Language Models","venue":"eess.AS","work_id":"d3f033ac-bfa8-4143-9d0d-51f3f5bd3f0e","year":2023},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"cited_paper":"/paper/2311.07919","citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:753dbbb4f44f9ff95f92de441f4341d88ff70327a083f0df709e55b0226d0228","observation_id":"0891cb16-4399-4f0e-b866-501bbb99ede7","resolution":{"observed_at":"2026-06-30T07:34:21.799793Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.20215","last_updated":"2025-03-26T04:17:55Z","snapshot_observed_at":"2026-08-06T08:46:20.194739Z","submitted_at":"2025-03-26T04:17:55Z","title":"Qwen2.5-Omni Technical Report","version":1},"cited_work":{"arxiv_id":"2503.20215","doi":"10.48550/arxiv.2503.20215","metadata_source":"pith","pith_arxiv_id":"2503.20215","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Qwen2.5-Omni Technical Report","venue":"cs.CL","work_id":"438f105c-fa9b-44aa-ad52-43acb8045cda","year":2025},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"cited_paper":"/paper/2503.20215","citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:d9cd82ef61dfcbb1244f548e97d2269875729a8af45cc7586ad07d4272b5fdde","observation_id":"72db6cf8-a841-4ab7-8e03-ce8448ebaf02","resolution":{"observed_at":"2026-06-30T07:34:21.790114Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-07-10T18:18:44.887499+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-10T18:18:44.887499+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"Lib- rispeech: An ASR corpus based on public domain audio books,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:48a205f6bedbb7512b824364dc4c774e5851ff5f78504207db5d80621d246812","observation_id":"d078c01e-ca32-45ef-b5a8-ffdba9edee69","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"SPGISpeech: 5,000 hours of transcribed financial audio for fully formatted end-to-end speech recognition,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:6cd6b2064e7940209ce21758132304b93a91bfa565ff895d2d9e3ccedb7dec8f","observation_id":"e72b7596-415c-4644-b4a7-e63a3c0a868a","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"Earnings-22: A practical benchmark for accents in the wild,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:12c3630c86d89d396fed41842013e639fcd3167e3c34bb81f4541634048c582d","observation_id":"c3d06017-5b5e-4330-b650-6f5f611c1426","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"GigaSpeech: An evolving, multi-domain ASR corpus with 10,000 hours of transcribed au- dio,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:316efbb02565b9ccb10d2fd8c1f4937c1f7a350348cafb887976f202cc5598c5","observation_id":"fba494af-0b5d-4eb0-a2a1-155ef55d10a3","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"The ami meeting corpus,","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:99a62e694259b28837bc819944f5f9a9c8262c18947025debc4b9436d33be716","observation_id":"3ab4f8fc-ede6-499e-b474-c009a1df4376","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"V oxPopuli: A large-scale multilingual speech corpus for representation learning, semi-supervised learning and interpreta- tion,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:4f7a2f1d05a0dd7f3566b5782616c6f3a863049b524851d63ff9ee22b6f1a534","observation_id":"43e704e1-e2e7-45e8-9dae-8c280bdc10f6","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"Common V oice: A massively-multilingual speech corpus,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:f8687c6ef55bc010c68b24c293655dff50a8a9d7a25c0066d0881ba02e9b4855","observation_id":"4db163c6-883b-4d16-94c9-5bafb678780d","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.05382","last_updated":"2021-07-07T12:52:49Z","snapshot_observed_at":"2026-08-03T13:17:40.818010Z","submitted_at":"2021-07-07T12:52:49Z","title":"End-to-End Rich Transcription-Style Automatic Speech Recognition with Semi-Supervised Learning","version":1},"cited_work":{"arxiv_id":"2107.05382","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.05382","snapshot_observed_at":"2026-06-30T07:34:21.788136Z","title":"End-to-end rich transcription-style automatic speech recognition with semi- supervised learning,","venue":null,"work_id":"f7acea6b-63e9-4678-9962-c851f77fa8e4","year":2021},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"cited_paper":"/paper/2107.05382","citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:233794285b7059e9505c16448f1fa8497a96d4e209dbb7c2fc31d84d5b527a1a","observation_id":"f7ab28a0-4a3e-4606-b5a6-bb0e85d638ed","resolution":{"observed_at":"2026-06-30T07:34:21.790001Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.18961","last_updated":"2024-08-15T21:32:57Z","snapshot_observed_at":"2026-08-02T17:13:54.584404Z","submitted_at":"2024-07-18T00:58:41Z","title":"MMAU: A Holistic Benchmark of Agent Capabilities Across Diverse Domains","version":3},"cited_work":{"arxiv_id":"2407.18961","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2407.18961","snapshot_observed_at":"2026-06-30T07:34:21.791157Z","title":"Mmau: A holistic benchmark of agent capabilities across diverse domains","venue":null,"work_id":"2018cdbe-c61a-4a44-9743-1748ea0c5e09","year":2024},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"cited_paper":"/paper/2407.18961","citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:73314a6eb847e8c9a0c641452aee7c3178828e92ef490ea5da70e43de69bdaa3","observation_id":"7caa8452-4898-4ae9-990a-869cd9a2b726","resolution":{"observed_at":"2026-06-30T07:34:21.793214Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.19037","last_updated":"2025-05-25T08:37:55Z","snapshot_observed_at":"2026-08-01T16:50:49.776636Z","submitted_at":"2025-05-25T08:37:55Z","title":"Speech-IFEval: Evaluating Instruction-Following and Quantifying Catastrophic Forgetting in Speech-Aware Language Models","version":1},"cited_work":{"arxiv_id":"2505.19037","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.19037","snapshot_observed_at":"2026-07-07T14:53:55.827006Z","title":"Speech-ifeval: Evaluating instruction-following and quantifying catastrophic forgetting in speech- aware language models","venue":"eess.AS","work_id":"486d8c9b-c1eb-4f21-8da6-20a0d8e2e34b","year":2025},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"cited_paper":"/paper/2505.19037","citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:197a2e9b5810ecc7c1e141602f34b6a7df7bdf9b87423b24f5f875ccaf26efed","observation_id":"4b5dee97-94de-4cce-8e2c-254715369ddf","resolution":{"observed_at":"2026-06-30T07:34:21.796509Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.09843","last_updated":"2023-09-18T14:59:10Z","snapshot_observed_at":"2026-08-05T10:54:37.106698Z","submitted_at":"2023-09-18T14:59:10Z","title":"Instruction-Following Speech Recognition","version":1},"cited_work":{"arxiv_id":"2309.09843","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2309.09843","snapshot_observed_at":"2026-06-30T07:34:21.801114Z","title":"Instruction-following speech recognition","venue":null,"work_id":"303c6c91-c7b7-4a50-b3c8-8b74031e2970","year":2023},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"cited_paper":"/paper/2309.09843","citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:4bb8303cba94701e6d0fb579e1f59d08ce92261bebc0b6caf46bd3c068c68955","observation_id":"3e9aedf8-aa48-42e6-94b8-84fb77e2848a","resolution":{"observed_at":"2026-06-30T07:34:21.802883Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"Open ASR leaderboard,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:10d296bf3ee6c9ee52141373d1d99ed83718f1a8b3fca00fa0f2e4fdff887f17","observation_id":"78197144-cf55-44c1-95c4-75cbb3926c46","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1611.00068","last_updated":"2017-01-24T19:51:12Z","snapshot_observed_at":"2026-07-06T05:16:53.476907Z","submitted_at":"2016-10-31T22:42:02Z","title":"RNN Approaches to Text Normalization: A Challenge","version":2},"cited_work":{"arxiv_id":"1611.00068","doi":null,"metadata_source":"pith","pith_arxiv_id":"1611.00068","snapshot_observed_at":"2026-07-04T16:59:58.647155Z","title":"RNN Approaches to Text Normalization: A Challenge","venue":"cs.CL","work_id":"498d8a5a-a311-4da7-9298-b4675401bf2a","year":2016},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"cited_paper":"/paper/1611.00068","citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:4e92aba52d7302e90096f3e968bbb02e675d273e0473cfad96ced7969c97166f","observation_id":"a2ef0593-a380-4d4f-8e62-6c0c92e35815","resolution":{"observed_at":"2026-06-30T07:34:21.805699Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"Deep context: End-to-end contex- tual speech recognition,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:8e6da6aba9f7f48d08647dc4dd24636055a8ed46f1439012c8ecdbf175db681b","observation_id":"08cce4f0-19e8-4719-b6b3-1a7d37a4a802","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"Spontaneous speech: How people really talk and why engineers should care,","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:44a29d6308bb39595b8fc88cbcaab57cde198879701b6cdc28ab7f30c562b2ff","observation_id":"bf751df5-580b-4a09-92a4-df3d1fa61f22","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"Bidirectional recurrent neural network with attention mechanism for punctuation restoration,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:bf7ff56ef11c7dc273501e2dbabbb39169244daef8328ef0c213a85090eaa3bb","observation_id":"7de4f260-105d-4f08-89ef-373d151a4191","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"Longer is (not necessarily) stronger: Punc- tuated long-sequence training for enhanced speech recognition and translation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:699bf1620daab7cd964bf0b6b6b7f8a4944c4aefcb0760664686ce04dcd7ebcd","observation_id":"da891a08-1b62-48fa-bd79-4f7acefea2c4","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":"2505.09388","doi":"10.1016/j.aiopen.2022.12","metadata_source":"pith","pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Qwen3 Technical Report","venue":"cs.CL","work_id":"25a4e30c-1232-48e7-9925-02fa12ba7c9e","year":2025},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:b058b8b9df7086a522484b0b419b51ffcaae4a9b4d120ecbe1f56e14af3d25f2","observation_id":"4aa9218a-5c3d-40ea-a763-e7e53761ea12","resolution":{"observed_at":"2026-06-30T07:34:21.798432Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"open asr leaderboard/normalizer,","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:28f07b1870e734a169b7f81675560c62598ba6d17db736746cae2b86aa625816","observation_id":"505b05c3-a23a-4bda-8df0-3c75223abfa8","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"Robust speech recognition via large-scale weak su- pervision,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:1f6c5d8ec61334084eb66d127dc187447b47cf7f320e2de5ff85b639de0a658e","observation_id":"8b507b33-a595-464f-8249-7bbf769c5f16","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"What is lost in normalization? Exploring pitfalls in multilingual ASR model evaluations,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:551b4baa02d3eee02fb1098c7e54551b93424688200adc8fa9d1890999e94ad0","observation_id":"cae87b92-83cc-4a4d-bdbb-1d24c1d07701","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2509.14128","doi":"10.48550/arxiv.2509.14128","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2509.14128 , year =","venue":"ArXiv.org","work_id":"7cfdf406-6e6c-4f9d-8072-9e9796b4ced8","year":2025},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:a49b466d314d0b32156c1647a60dcbbfcbf9f56ea049b2c26e3e7cd9e69e79db","observation_id":"cc1ff3ed-950d-4dfc-bddc-cf8196f1f909","resolution":{"observed_at":"2026-06-30T07:34:21.782612Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"Fast conformer with linearly scalable attention for efficient speech recognition,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:45cca5385b4c5ec87834217b2dbaf47af5a7ae5b3fd0702c8fc1e2222d257696","observation_id":"b71f859e-edd8-45cd-b987-468deb4c40d0","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"Efficient sequence transduction by jointly predicting tokens and durations,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:c591955896eef1a591b04b9ef05121102908d3b6403529d0ddbe9eb3f2670001","observation_id":"4c2c7284-fdd4-421f-a083-c59b72951ba4","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"Qwen2. 5 technical report,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:d3a568bd80408336c5e47cc96e3ccd0279d5be79c8f1affa51a20c8bb09b3495","observation_id":"bb454d88-8374-46c4-b610-317859993007","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"Lora: Low-rank adaptation of large language models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:745e7af7e6f1290e39291e8e8db23f87d7aeced7480deb5de09a940ccaae1a36","observation_id":"8cc82e96-1ff4-4949-a65d-98a87d4c6cb4","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.01743","last_updated":"2025-03-07T09:05:58Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-03-03T17:05:52Z","title":"Phi-4-Mini Technical Report: Compact yet Powerful Multimodal Language Models via Mixture-of-LoRAs","version":2},"cited_work":{"arxiv_id":"2503.01743","doi":"10.18653/v1/2023.wmt-1.23","metadata_source":"pith","pith_arxiv_id":"2503.01743","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Phi-4-Mini Technical Report: Compact yet Powerful Multimodal Language Models via Mixture-of-LoRAs","venue":"cs.CL","work_id":"83956045-536a-41ff-af02-b80e2a614eab","year":2025},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"cited_paper":"/paper/2503.01743","citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:7421921cb26fa7393a42ed42806213db8dc344cab7a00624ebc9c4aa95823aec","observation_id":"b03c37ba-07ad-4462-a21e-b1dfda67028e","resolution":{"observed_at":"2026-06-30T07:34:21.779789Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.17765","last_updated":"2025-09-22T13:26:24Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-09-22T13:26:24Z","title":"Qwen3-Omni Technical Report","version":1},"cited_work":{"arxiv_id":"2509.17765","doi":"10.48550/arxiv.2509.17765","metadata_source":"pith","pith_arxiv_id":"2509.17765","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Qwen3-Omni Technical Report","venue":"cs.CL","work_id":"ae43e594-8bab-4471-b6af-92a300f6a048","year":2025},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"cited_paper":"/paper/2509.17765","citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:c9fa13622505b0fc6a815dc44048044bfe52b61778a76d85d5db55ab07677f2b","observation_id":"3fe6b5a9-40fc-4198-9b60-0daa379a17e9","resolution":{"observed_at":"2026-06-30T07:34:21.808452Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"NVIDIA NeMo Canary-Qwen-2.5B,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:67b74299baef023410d18b6b6b3cd8da692ac15350b0966af8af9cd28512baae","observation_id":"27bdba1f-014e-4fdb-9280-fba7b745c280","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":"Um, and profit aim is fifty million Euros, which is uh","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:3b2f445a29937bb03d360bc3dff9febbe41af51a97f7f17f373aba155730b503","observation_id":"7c1fb2fb-025c-4e7f-ba48-7963825cf84a","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T07:26:38.380118Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-06-30T07:26:38.380118Z"},"links":{"citing_paper":"/paper/2606.29534"},"observation_digest":"sha256:62550a5d6e45229bcf738c2bb3c21f13d7b54b9b7aa49b54bdc2569c0227531a","observation_id":"b21faa45-485a-4b5b-8122-55a784e69741","resolution":{"observed_at":"2026-06-30T07:26:38.380118Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2606.29534","last_updated":"2026-06-28T17:57:41Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-07-07T00:03:31.715960Z","submitted_at":"2026-06-28T17:57:41Z","title":"Preference-ASR: A Preference-Aware Test Set for Benchmarking ASR in the Era of Speech LLMs"},"reference_resolution":{"displayed":44,"state_counts":{"malformed_identifier":2,"metadata_mismatch":2,"parse_uncertain":0,"unresolved":30,"verified_exact":10,"verified_fuzzy":0},"total_outbound_references":44},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"thesis":"As of 6 August 2026, this Paper Citation Record lists 44 of 44 outbound references and 2 inbound Pith citation observations for arXiv:2606.29534."}