{"as_of":"2026-08-06T12:20:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:f4c25cf0cdcdc974173189423afef987b62f40d6360191891cef27fe26c114a3","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":32,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":32,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-06T06:34:29.942622+00:00","state":"measured"},{"denominator":32,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":32,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T04:17:35.289242Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-10T20:37:34.440026Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-08-05T21:45:22.388008Z","title":"An embarrassingly simple approach for llm with strong asr capacity,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.08095","last_updated":"2025-08-11T15:33:44Z","snapshot_observed_at":"2026-08-05T21:45:21.468156Z","submitted_at":"2025-08-11T15:33:44Z","title":"Dual Information Speech Language Models for Emotional Conversations","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-05T21:45:22.388008Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2508.08095"},"observation_digest":"sha256:b3bb615601605addd0246f144ba36b8f9d1b0c4c0edb2791eeb79641e24955c6","observation_id":"d2ab4cb9-f143-45b2-9f74-ba441f2be27a","resolution":{"observed_at":"2026-08-05T21:45:22.388008Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-08-06T04:17:35.289242Z","title":"Ma and et al., ``An embarrassingly simple approach for llm with strong asr capacity,'' arXiv preprint arXiv:2402.08846, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.14048","last_updated":"2025-08-05T17:53:11Z","snapshot_observed_at":"2026-08-06T04:17:23.253592Z","submitted_at":"2025-08-05T17:53:11Z","title":"RAG-Boost: Retrieval-Augmented Generation Enhanced LLM-based Speech Recognition","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-06T04:17:35.289242Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2508.14048"},"observation_digest":"sha256:d997147a034de5f9e549e7b3d1a97ca7f5f97d3dbb803408e6fd902ac25c387e","observation_id":"8db43db0-dba1-4fca-8ce9-bd411e41521b","resolution":{"observed_at":"2026-08-06T04:17:35.289242Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-08-05T20:01:03.613617Z","title":"An embarrassingly simple approach for llm with strong asr capacity,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.14916","last_updated":"2025-08-15T10:39:05Z","snapshot_observed_at":"2026-08-05T20:01:03.228253Z","submitted_at":"2025-08-15T10:39:05Z","title":"Transsion Multilingual Speech Recognition System for MLC-SLM 2025 Challenge","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-05T20:01:03.613617Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2508.14916"},"observation_digest":"sha256:17fa70eab611813a14e9179c41a34c76d221b5a6896cd4d07da08fa5af7c2a5f","observation_id":"b8005c93-cc00-44fa-bb0f-4ec3663f192a","resolution":{"observed_at":"2026-08-05T20:01:03.613617Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-08-05T16:46:49.196106Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.17863","last_updated":"2025-08-25T10:16:07Z","snapshot_observed_at":"2026-08-06T09:02:37.213294Z","submitted_at":"2025-08-25T10:16:07Z","title":"Speech Discrete Tokens or Continuous Features? A Comparative Analysis for Spoken Language Understanding in SpeechLLMs","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-05T16:46:49.196106Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2508.17863"},"observation_digest":"sha256:2d28c993fcfdc6213b661ac2ad22f905bb20c340db0a0d637cbc760e93499bb7","observation_id":"ca7456a1-cfec-42dc-8ed9-fd7c83588c02","resolution":{"observed_at":"2026-08-05T16:46:49.196106Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-08-05T15:27:29.280697Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.19856","last_updated":"2025-08-27T13:16:31Z","snapshot_observed_at":"2026-08-05T15:27:28.932842Z","submitted_at":"2025-08-27T13:16:31Z","title":"TokenVerse++: Towards Flexible Multitask Learning with Dynamic Task Activation","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-05T15:27:29.280697Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2508.19856"},"observation_digest":"sha256:0a7b3bd17e0ca158f5d6fdb48b32c0bf0d50aa0bfccabdee5f999091d6ad8067","observation_id":"38f2f1b3-3766-4a55-a261-a2edd83ae08c","resolution":{"observed_at":"2026-08-05T15:27:29.280697Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-08-05T12:07:21.317118Z","title":"An embarrassingly simple approach for llm with strong asr capacity,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.01939","last_updated":"2025-09-02T04:20:12Z","snapshot_observed_at":"2026-08-06T07:12:18.986084Z","submitted_at":"2025-09-02T04:20:12Z","title":"Group Relative Policy Optimization for Speech Recognition","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-05T12:07:21.317118Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2509.01939"},"observation_digest":"sha256:580fbbb810e1c765d256b6abc7c9e7a945947db47622fb9933d8c9d0a0278906","observation_id":"dabc5971-19b3-45e2-8aa3-f0fbed91b423","resolution":{"observed_at":"2026-08-05T12:07:21.317118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-08-05T10:16:04.817224Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.04392","last_updated":"2025-09-04T17:03:58Z","snapshot_observed_at":"2026-08-05T10:16:04.026189Z","submitted_at":"2025-09-04T17:03:58Z","title":"Denoising GER: A Noise-Robust Generative Error Correction with LLM for Speech Recognition","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-05T10:16:04.817224Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2509.04392"},"observation_digest":"sha256:c64a19d3be46b3aeca0012cb1eb3cc1f351148bb48053c5f53e8cfe1c879fa24","observation_id":"16abd9eb-df37-4f74-9542-77af2a53c1b2","resolution":{"observed_at":"2026-08-05T10:16:04.817224Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-08-05T13:53:23.962333Z","title":"An embarrassingly simple approach for llm with strong asr capacity,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.04473","last_updated":"2025-08-29T22:38:16Z","snapshot_observed_at":"2026-08-05T13:53:22.566210Z","submitted_at":"2025-08-29T22:38:16Z","title":"SpeechLLM: Unified Speech and Language Model for Enhanced Multi-Task Understanding in Low Resource Settings","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-05T13:53:23.962333Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2509.04473"},"observation_digest":"sha256:bec1deef8fa93fccddf8ecd22ee77326021ebc90e1ec0c7371e75e5dffaf1828","observation_id":"0a2508d2-110c-4be6-afcc-a424821966ce","resolution":{"observed_at":"2026-08-05T13:53:23.962333Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-08-04T11:29:34.295228Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.04593","last_updated":"2026-06-08T05:49:30Z","snapshot_observed_at":"2026-08-06T08:55:28.794308Z","submitted_at":"2025-10-06T08:47:38Z","title":"UniVoice: Unifying Autoregressive ASR and Flow-Matching based TTS with Large Language Models","version":3},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-04T11:29:34.295228Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2510.04593"},"observation_digest":"sha256:de299d70fbcc7926116015246f900d84b0ea1af54e3d1abf3816edacb0be6d32","observation_id":"f0f964f5-1a8a-4393-b402-e7f00e95610b","resolution":{"observed_at":"2026-08-04T11:29:34.295228Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2601.20898","last_updated":"2026-01-28T09:50:34Z","snapshot_observed_at":"2026-07-06T22:43:21.411174Z","submitted_at":"2026-01-28T09:50:34Z","title":"Reducing Prompt Sensitivity in LLM-based Speech Recognition Through Learnable Projection","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-16T10:37:56.416600Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2601.20898"},"observation_digest":"sha256:b80fd676c7ed459034ec1d516d9e2c6485c9074dc176cc8b930325e31d60a7a2","observation_id":"a0b59a83-9cb8-49fa-a2a4-4b2c8699e090","resolution":{"observed_at":"2026-05-16T10:40:51.425344Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2603.15045","last_updated":"2026-07-07T07:07:08Z","snapshot_observed_at":"2026-08-01T21:49:24.005109Z","submitted_at":"2026-03-16T09:57:06Z","title":"LLMs and Speech: Integration vs. Combination","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-15T10:41:22.138517Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2603.15045"},"observation_digest":"sha256:391305cba27c0891e3ce006315ef8346d076226c17d2a7f2b186062103c6d5cf","observation_id":"ffd01fc1-cdd3-4e8c-b2a2-fd18b9a8ff10","resolution":{"observed_at":"2026-05-15T10:45:28.356069Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-14T20:46:17.286576Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.15045","last_updated":"2026-07-07T07:07:08Z","snapshot_observed_at":"2026-08-01T21:49:24.005109Z","submitted_at":"2026-03-16T09:57:06Z","title":"LLMs and Speech: Integration vs. Combination","version":3},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-07-14T20:46:17.286576Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2603.15045"},"observation_digest":"sha256:8a85e9459fbb9e962696d192cc25036fd86ccae75c7e784259c33ca10e725557","observation_id":"21e856e1-d7e0-4d88-bdde-df5b4978a8f5","resolution":{"observed_at":"2026-07-14T20:46:17.286576Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2604.06487","last_updated":"2026-04-07T21:41:16Z","snapshot_observed_at":"2026-08-01T19:23:38.360811Z","submitted_at":"2026-04-07T21:41:16Z","title":"Closing the Speech-Text Gap with Limited Audio for Effective Domain Adaptation in LLM-Based ASR","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T18:41:31.431567Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2604.06487"},"observation_digest":"sha256:920688a5ecb5e81fc573205af2574605ea57159564435824f5befa5150afecbc","observation_id":"8fd1d5bc-f9d8-4245-ac25-5860d4d827e1","resolution":{"observed_at":"2026-05-11T00:05:52.056588Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2604.09121","last_updated":"2026-04-14T06:45:50Z","snapshot_observed_at":"2026-08-02T20:05:01.635841Z","submitted_at":"2026-04-10T09:02:42Z","title":"Interactive ASR: Towards Human-Like Interaction and Semantic Coherence Evaluation for Agentic Speech Recognition","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-10T18:22:08.670559Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2604.09121"},"observation_digest":"sha256:7f67cee7a6e09e956b9a2ffa4c114c94ddfc609c146051de71dae3562ad157a2","observation_id":"63ac30c2-c9f2-4acb-97ca-3450bc6c544e","resolution":{"observed_at":"2026-05-11T00:45:49.116917Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2604.09332","last_updated":"2026-04-10T14:00:24Z","snapshot_observed_at":"2026-07-06T22:58:12.847181Z","submitted_at":"2026-04-10T14:00:24Z","title":"Phonemes vs. Projectors: An Investigation of Speech-Language Interfaces for LLM-based ASR","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-10T16:33:32.012025Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2604.09332"},"observation_digest":"sha256:0ea9166d18fa11e3f3d0be1efe6b8dd04d513503d0a9d8abd9ff3f12ded4ca15","observation_id":"a6796af1-5374-4c57-8f2f-e507b4237054","resolution":{"observed_at":"2026-05-11T08:40:59.471763Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2604.11269","last_updated":"2026-04-13T10:23:58Z","snapshot_observed_at":"2026-07-06T22:59:40.900521Z","submitted_at":"2026-04-13T10:23:58Z","title":"Speaker Attributed Automatic Speech Recognition Using Speech Aware LLMS","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-10T15:08:56.412939Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2604.11269"},"observation_digest":"sha256:3ad52173e544a919ac6fd97f7da5c12fbd10179cc15b3fa5a5d1660d64981264","observation_id":"96265c51-0a44-4540-9c46-f7f1196ab634","resolution":{"observed_at":"2026-05-11T11:11:00.884611Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2604.12398","last_updated":"2026-04-14T07:33:44Z","snapshot_observed_at":"2026-07-06T23:00:37.144068Z","submitted_at":"2026-04-14T07:33:44Z","title":"Contextual Biasing for ASR in Speech LLM with Common Word Cues and Bias Word Position Prediction","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-10T14:35:20.515352Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2604.12398"},"observation_digest":"sha256:2f35efff5601f532359a5f6dc382601625ea420de098f78fec8f9b20f466b9e3","observation_id":"d5c21b94-693d-4df0-ae81-e480800009b4","resolution":{"observed_at":"2026-05-10T14:35:33.952747Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2604.22817","last_updated":"2026-04-14T20:56:24Z","snapshot_observed_at":"2026-07-06T23:09:10.050398Z","submitted_at":"2026-04-14T20:56:24Z","title":"In-Sync: Adaptation of Speech Aware Large Language Models for ASR with Word Level Timestamp Predictions","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-10T13:25:50.524448Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2604.22817"},"observation_digest":"sha256:a2bd8b5b1e037f0bcd9dae4535e784039be8b2cff11e8c3a9d42e310927adc15","observation_id":"4f734acb-6cfd-4bdc-a430-ccea693940fc","resolution":{"observed_at":"2026-05-10T13:35:26.718348Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2605.14340","last_updated":"2026-06-25T04:41:51Z","snapshot_observed_at":"2026-08-04T15:57:05.501637Z","submitted_at":"2026-05-14T04:04:03Z","title":"Refining Pseudo-Audio Prompts with Speech-Text Alignment for Text-Only Domain Adaptation in LLM-Based ASR","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-15T02:20:53.100809Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2605.14340"},"observation_digest":"sha256:ff094ccb7580c70fce38eca31a3c0080e84cc7824c918dfd8391e0ada1f02474","observation_id":"1cc8376e-1167-47d1-84d9-45cd2b614353","resolution":{"observed_at":"2026-05-15T02:23:31.982546Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2605.14340","last_updated":"2026-06-25T04:41:51Z","snapshot_observed_at":"2026-08-04T15:57:05.501637Z","submitted_at":"2026-05-14T04:04:03Z","title":"Refining Pseudo-Audio Prompts with Speech-Text Alignment for Text-Only Domain Adaptation in LLM-Based ASR","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-30T20:42:10.489048Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2605.14340"},"observation_digest":"sha256:5352c412d79512d4c0f901a362c6949bfb9be036bda0485e8b405ee3d42ec2a8","observation_id":"c0f63a5a-61d1-4360-9057-207586259bc6","resolution":{"observed_at":"2026-06-30T20:45:03.653925Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2606.02400","last_updated":"2026-06-02T17:20:11Z","snapshot_observed_at":"2026-08-04T13:42:09.326463Z","submitted_at":"2026-06-01T15:47:01Z","title":"SoulX-Transcriber: A Robust End-to-End Framework for Multi-Speaker Speech Transcription","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-28T12:37:52.271987Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2606.02400"},"observation_digest":"sha256:4e6397f95e6d9f6df32b7df43d66ac32d81dccd677b02b91d55731321178b3a4","observation_id":"bdd076fe-6dd0-445d-b3f1-f57e0a6feeb6","resolution":{"observed_at":"2026-07-02T01:06:24.528096Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2606.08486","last_updated":"2026-06-07T07:15:34Z","snapshot_observed_at":"2026-07-06T23:47:57.376596Z","submitted_at":"2026-06-07T07:15:34Z","title":"TRADE: Transducer-Augmented Decoder for Speech LLM","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-06-27T18:40:19.688550Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2606.08486"},"observation_digest":"sha256:b3d7f1f1d817d514649e363a1c972e75dec6ccb6e664e05a728ac2f71f866b04","observation_id":"d4277d40-1e2b-4807-87b2-982b3b0785da","resolution":{"observed_at":"2026-07-02T22:47:25.702381Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2606.09966","last_updated":"2026-06-08T16:29:59Z","snapshot_observed_at":"2026-07-06T23:49:12.754871Z","submitted_at":"2026-06-08T16:29:59Z","title":"RespiraMFM: A Multimodal Foundation Model with Contrastive Audio-Language Alignment for Respiratory Disease Identification","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-06-27T15:05:31.609683Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2606.09966"},"observation_digest":"sha256:e3b241c90bfdc6eb88dcb5e5808e75099282808dba77b6adc6c20383bf89f728","observation_id":"632df31e-02d9-4c56-b09d-cab846e7c658","resolution":{"observed_at":"2026-07-03T03:37:35.640798Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2606.10368","last_updated":"2026-06-09T03:27:30Z","snapshot_observed_at":"2026-07-06T23:49:37.345711Z","submitted_at":"2026-06-09T03:27:30Z","title":"Speech Meets ELF: Audio Conditional Continuous-Target Diffusion for Speech Recognition and Translation","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-06-27T12:04:50.483329Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2606.10368"},"observation_digest":"sha256:c26234bf1d8fa4a4bf343d2a4e848016ec9bf89fb53bd38212aad4ddfb13cd3e","observation_id":"ea183f7d-918b-4e36-970f-7b983777c7f2","resolution":{"observed_at":"2026-07-03T07:27:44.951000Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2606.10439","last_updated":"2026-06-09T05:35:31Z","snapshot_observed_at":"2026-08-06T05:02:13.845939Z","submitted_at":"2026-06-09T05:35:31Z","title":"Enhancing Multilingual LLM-based ASR with Mixture of Experts and Dynamic Downsampling","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-27T11:56:08.586936Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2606.10439"},"observation_digest":"sha256:b0607e1731c5e4bfebda84e4b2a539928da2d47280adff867d1e59df6ab915b8","observation_id":"d7aacf95-318f-4cbf-965b-9d006b84d17b","resolution":{"observed_at":"2026-07-03T07:37:45.564372Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2606.10454","last_updated":"2026-06-09T06:02:31Z","snapshot_observed_at":"2026-08-03T09:50:36.542776Z","submitted_at":"2026-06-09T06:02:31Z","title":"Entropy-Aware Domain-Routed Mixture-of-Experts Speech-LLM Framework: A Case Study of Multi-Domain Child-Adult ASR","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-27T12:05:48.855615Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2606.10454"},"observation_digest":"sha256:6ea250561efaed17a2c7d53802c92fe22a39b3ffe3c64d6e5aff6d08a6114807","observation_id":"cfb87c54-0c80-4817-b9a6-c43cd55448eb","resolution":{"observed_at":"2026-07-03T07:27:44.895442Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2606.24123","last_updated":"2026-06-23T04:14:07Z","snapshot_observed_at":"2026-08-06T09:18:37.300889Z","submitted_at":"2026-06-23T04:14:07Z","title":"Aligning MusicLLM with Emotion using Instruction Tuning and Feedback-Driven Alignment","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-06-25T22:56:12.155077Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2606.24123"},"observation_digest":"sha256:06b5de74ff7dd9bd18f59dc32859606fe269e0d0a13d5f8fc5f718b8ad22598d","observation_id":"94579e9f-e5a5-4763-b6ad-d074958d75a7","resolution":{"observed_at":"2026-07-04T18:20:04.579470Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2606.25444","last_updated":"2026-06-24T06:15:18Z","snapshot_observed_at":"2026-08-02T19:50:32.355214Z","submitted_at":"2026-06-24T06:15:18Z","title":"Does Translation-Enhanced Speech Encoder Pre-training Affect Speech LLMs?","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-06-25T20:03:10.858349Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2606.25444"},"observation_digest":"sha256:95c2cf11cd70189f73843f589030d505e44150d9a0672c9e1f86f0a47cbb5c19","observation_id":"bdccff6e-7a42-4bff-9dbd-f250118945e4","resolution":{"observed_at":"2026-07-04T20:30:08.201818Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-12T07:12:19.297723Z","title":"SLAM-ASR: An Embarrassingly Sim- ple Approach for LLM with Strong ASR Capacity,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.02763","last_updated":"2026-07-02T21:00:52Z","snapshot_observed_at":"2026-08-06T11:33:55.365451Z","submitted_at":"2026-07-02T21:00:52Z","title":"LuxSQA: Ask Me in Luxembourgish with TTS-Augmented Spoken Question Answering","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-12T07:12:19.297723Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2607.02763"},"observation_digest":"sha256:3f2b8345e7e9f3f695b52dc515d85cc9bd595de8f20bcb6158b9d6c5d02eeaf1","observation_id":"0867094c-0434-40cb-9141-7303a8ed75e6","resolution":{"observed_at":"2026-07-12T07:12:19.297723Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2607.06827","last_updated":"2026-07-07T21:45:04Z","snapshot_observed_at":"2026-07-11T23:18:39.332726Z","submitted_at":"2026-07-07T21:45:04Z","title":"Compress the Cache, Not the Speech Embedding: KV Compression for Efficient Speech LLMs","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-07-10T20:30:11.127007Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2607.06827"},"observation_digest":"sha256:68d6c18492737cf47b1181f42d14fce89b020cdcb582e6bf6cf94e8ccb06a8c5","observation_id":"3f7e0a52-5d58-43fa-926d-e60fa96128a2","resolution":{"observed_at":"2026-07-10T20:37:34.441426Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":"2402.08846","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-07-10T20:37:34.440026Z","title":"An embarrassingly simple approach for LLM with strong ASR capacity","venue":"cs.CL","work_id":"d0cf5313-bc88-4f8f-9a2c-9012a25a93dd","year":2024},"citing_paper":{"arxiv_id":"2607.08409","last_updated":"2026-07-09T12:34:56Z","snapshot_observed_at":"2026-08-06T00:03:36.736051Z","submitted_at":"2026-07-09T12:34:56Z","title":"When Synthetic Speech Is All You Have: Better Call GRPO","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-10T07:53:41.580597Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2607.08409"},"observation_digest":"sha256:5767d2c7a33e6e8bf35bd39837c299a308c6a2ec193ac7b06275f20375a7af4c","observation_id":"c72269b7-e68e-40fd-9e73-4195f71fe3d5","resolution":{"observed_at":"2026-07-10T07:56:57.809471Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.08846","snapshot_observed_at":"2026-08-01T05:50:28.705395Z","title":"An embarrassingly sim- ple approach for LLM with strong ASR capacity,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.22100","last_updated":"2026-07-24T08:52:44Z","snapshot_observed_at":"2026-08-02T17:58:26.013216Z","submitted_at":"2026-07-24T08:52:44Z","title":"MEUSLI: a Multilingual Projector for LLM-based ASR and Beyond","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-01T05:50:28.705395Z"},"links":{"cited_paper":"/paper/2402.08846","citing_paper":"/paper/2607.22100"},"observation_digest":"sha256:6eb3f98ff1bc920a8dfc1f6537ac8bb59cdbc2aeed462850e7267ab44f926e84","observation_id":"809360d3-6ffb-4409-a6c6-00ac8e76d07b","resolution":{"observed_at":"2026-08-01T05:50:28.705395Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2402.08846/citation-record","integrity":"/paper/2402.08846/integrity","json":"/paper/2402.08846/citation-record.json","paper":"/paper/2402.08846"},"outbound":[],"paper":{"arxiv_id":"2402.08846","last_updated":"2024-02-13T23:25:04Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-06T09:11:01.697195Z","submitted_at":"2024-02-13T23:25:04Z","title":"An Embarrassingly Simple Approach for LLM with Strong ASR Capacity"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"thesis":"As of 6 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 32 inbound Pith citation observations for arXiv:2402.08846."}