{"as_of":"2026-08-08T12:27:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:3dfa7c4dea09435f5fe9e4424d94988f032455dc6053f1fce1c1cb8e864b7106","coverage":[{"denominator":30,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":30,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T05:06:37.433193Z","state":"measured"},{"denominator":30,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":30,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.08836/citation-record","integrity":"/paper/2506.08836/integrity","json":"/paper/2506.08836/citation-record.json","paper":"/paper/2506.08836"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.319054Z","title":"ss\" instead of","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.319054Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:3c69957759b2e8a93e199619faf52c27103a8b1cff9f2ce37af1e4336552bb48","observation_id":"6b5c2e59-d554-4e47-bc31-1350d214273e","resolution":{"observed_at":"2026-08-07T05:06:37.319054Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.114492Z","title":"O’Reilly Media Inc","venue":null,"work_id":"98377288-0e08-404b-a41e-6787c9df3259","year":2009},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.324007Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:a534be123682ef9023d4cd8f8c7219eff99caaec480d37f322469c743da6d11e","observation_id":"1d8532df-3e00-45d9-a91a-fd544949a730","resolution":{"observed_at":"2026-08-07T05:06:38.118282Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1604.06174","last_updated":"2016-04-22T19:21:36Z","snapshot_observed_at":"2026-08-08T09:03:25.135475Z","submitted_at":"2016-04-21T04:15:27Z","title":"Training Deep Nets with Sublinear Memory Cost","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1604.06174","snapshot_observed_at":"2026-08-07T05:06:37.328092Z","title":"https://doi.org/10.48550/arXiv.1604.06174, http: //arxiv.org/abs/1604.06174, arXiv:1604.06174 [cs] version: 2","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.328092Z"},"links":{"cited_paper":"/paper/1604.06174","citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:b594dc5a4675830f4aeda79b06b2a34d2ad86f4c1285b11b7489c58727177c84","observation_id":"9aa1ceef-7b5c-4708-a268-ad9995cafd0a","resolution":{"observed_at":"2026-08-07T05:06:37.328092Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.102118Z","title":null,"venue":null,"work_id":"8d1018f5-fcb7-4105-a235-788f562f26ec","year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.332473Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:3f60ad6f4742fc7b85f6bdc74cd29a3ef89122043e39800aa57e012eb26a9738","observation_id":"867c8dcf-1392-4925-b852-b04b1019e69e","resolution":{"observed_at":"2026-08-07T05:06:38.105913Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.10299","last_updated":"2023-09-19T04:04:14Z","snapshot_observed_at":"2026-07-06T16:20:21.258653Z","submitted_at":"2023-09-19T04:04:14Z","title":"Using fine-tuning and min lookahead beam search to improve Whisper","version":1},"cited_work":{"arxiv_id":"2309.10299","doi":null,"metadata_source":"pith","pith_arxiv_id":"2309.10299","snapshot_observed_at":"2026-08-07T05:06:37.982734Z","title":"Using fine-tuning and min lookahead beam search to improve Whisper","venue":"eess.AS","work_id":"e606724e-a4e8-4e81-8cf3-947ef596731e","year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.336417Z"},"links":{"cited_paper":"/paper/2309.10299","citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:764fb48db7096d523f4a5f35aff449f778e78328bb9d411db07b70b68495acfb","observation_id":"ae460027-1a19-42de-9574-d34a26ae44db","resolution":{"observed_at":"2026-08-07T05:06:37.987225Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2103.11401","last_updated":"2021-03-21T14:00:09Z","snapshot_observed_at":"2026-07-06T10:51:51.998728Z","submitted_at":"2021-03-21T14:00:09Z","title":"SwissDial: Parallel Multidialectal Corpus of Spoken Swiss German","version":1},"cited_work":{"arxiv_id":"2103.11401","doi":null,"metadata_source":"pith","pith_arxiv_id":"2103.11401","snapshot_observed_at":"2026-08-07T05:06:37.963006Z","title":"SwissDial: Parallel Multidialectal Corpus of Spoken Swiss German","venue":"cs.CL","work_id":"5d8ff65d-0d69-4a14-8e3f-279e664c1577","year":2021},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.340733Z"},"links":{"cited_paper":"/paper/2103.11401","citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:dc9891684ea82b88083a72a337d877ddeb86c4d5594465484777ada09d72960c","observation_id":"eaaf9627-300f-4f8c-84f7-75cd3496724c","resolution":{"observed_at":"2026-08-07T05:06:37.968157Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.345312Z","title":"In: Scherrer, Y., Jauhiainen, T., Ljubešić, N., Zampieri, M., Nakov, P., Tiedemann, J","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.345312Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:d6c8142150f927cf20b1b10b41907f623433d92ba0194041fe0aff133025926c","observation_id":"fbc4e6e2-6eef-43eb-8b3c-0dfafc04ca60","resolution":{"observed_at":"2026-08-07T05:06:37.345312Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.349048Z","title":"In: ICASSP 2024 - 2024 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.349048Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:edf627776e45fd56f67190a95c6a8cc0122074d0956fac4e99d3de73b4ba1697","observation_id":"a283b7b9-f3cf-4152-9e6c-0cf4b0c498d4","resolution":{"observed_at":"2026-08-07T05:06:37.349048Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.089074Z","title":null,"venue":null,"work_id":"3e4c7e0a-1dbb-457d-98db-e07961fc83db","year":2022},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.352822Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:abbf89c058b8ff981a77f8a4e3f12953302640af53c06daa425fa980d8a530d3","observation_id":"6b580da3-d246-4983-ba21-cb2d20fdd5ed","resolution":{"observed_at":"2026-08-07T05:06:38.093299Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.356410Z","title":"In: Interspeech 2020","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.356410Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:47d990fb917c0caf6e525410013fcc659e3878a4446bd4bf00382a8db313a5fe","observation_id":"37b6b99c-72ca-4d93-875c-9971a476ee8a","resolution":{"observed_at":"2026-08-07T05:06:37.356410Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.10429","last_updated":"2025-06-14T16:46:01Z","snapshot_observed_at":"2026-07-06T19:16:10.416481Z","submitted_at":"2024-09-16T16:04:16Z","title":"SMILE: Speech Meta In-Context Learning for Low-Resource Language Automatic Speech Recognition","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.10429","snapshot_observed_at":"2026-08-07T05:06:37.360152Z","title":"https://doi.org/10.48550/arXiv","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.360152Z"},"links":{"cited_paper":"/paper/2409.10429","citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:aa5ab2851c6b8a4c6825c9f7432cc7bcd67790e5128e22457251f5c0812e92e8","observation_id":"6ab666c8-9a4a-44cc-a65d-22212e507764","resolution":{"observed_at":"2026-08-07T05:06:37.360152Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.364239Z","title":"EURASIP Journal on Audio, Speech, and Music Process- ing 2024(1), 29 (Jun 2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.364239Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:14e9a0a3dcaa11ab5a730118e435059ea1de33f1d4c4ca71ba040241247e4025","observation_id":"8ea9578c-aecc-4f2f-ad7a-815e594633a2","resolution":{"observed_at":"2026-08-07T05:06:37.364239Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.075349Z","title":null,"venue":null,"work_id":"36ca8b22-2430-46d2-9ec3-cd5aec053ed0","year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.367955Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:bd54338d84a8d710c42fa10677232477b7c88922a849b50ec327e641350108a9","observation_id":"e5f2f6c4-ebb4-4807-bc86-097330cfa496","resolution":{"observed_at":"2026-08-07T05:06:38.079769Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.061351Z","title":"In: 7th Inter- national Conference on Learning Representations, ICLR 2019, New Orleans, LA, USA, May 6-9, 2019","venue":null,"work_id":"4676d0f1-4ebf-4fa1-a17c-24a5ce11963c","year":2019},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.371450Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:ebd4613408958d366c63cc1c2765f21e24c3451396d5454e706972fa567b9468","observation_id":"53d79b48-2038-4aa2-8a40-f4952e7b80f2","resolution":{"observed_at":"2026-08-07T05:06:38.065513Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.21437/slate.2023-20","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.551655Z","title":"In: 9th Workshop on Speech and Language Tech- nology in Education (SLaTE)","venue":null,"work_id":"ab2d12d3-da34-472c-bf4e-cccf82ea1a95","year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.374904Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:181d63d6bdb46cff7ee56670cc3499b7ee70ce87f100d6ca56a8e3f7ffc2067b","observation_id":"136106d7-1175-4976-86f6-5497f7947566","resolution":{"observed_at":"2026-08-07T05:06:37.555540Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2023.findings-emnlp.1018","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.538094Z","title":"In: Bouamor, H., Pino, J., Bali, K","venue":null,"work_id":"b4eaa086-6b93-4e70-b704-a138aae132c0","year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.378461Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:c5be4423665fc00f2735993748ea44d44be35fd7ef39ce3357be73f4cfd71103","observation_id":"3cd49efa-28ac-4cd1-985b-198c398a2c55","resolution":{"observed_at":"2026-08-07T05:06:37.542081Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.382189Z","title":"In: Isabelle, P., Charniak, E., Lin, D","venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.382189Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:4e3cd3470816dd556820dd9adc8ce880d965f0ce71a2da8d6ebbc4816b56a8df","observation_id":"f5a0ff9b-dc99-417d-920e-991dbdda1a9d","resolution":{"observed_at":"2026-08-07T05:06:37.382189Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.04573","last_updated":"2024-11-07T09:57:57Z","snapshot_observed_at":"2026-07-06T19:46:37.095394Z","submitted_at":"2024-11-07T09:57:57Z","title":"Multistage Fine-tuning Strategies for Automatic Speech Recognition in Low-resource Languages","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.04573","snapshot_observed_at":"2026-08-07T05:06:37.385874Z","title":"https://doi.org/10.48550/arXiv.2411.04573, http://arxiv.org/abs/ 2411.04573, arXiv:2411.04573 [cs]","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.385874Z"},"links":{"cited_paper":"/paper/2411.04573","citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:c40981ec4b3d4e59256e55a9fcff3c7defe9d594edcf5f4fcd72045931c28238","observation_id":"04d042c8-45a9-47f2-b339-274358dae5e0","resolution":{"observed_at":"2026-08-07T05:06:37.385874Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.21437/interspeech.2024-734","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.511253Z","title":"In: Interspeech 2024","venue":null,"work_id":"3bf67a5f-a42b-4e12-8eab-ed5f9b43283f","year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.389869Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:55085281646e78b253930b1487d4ef95dfe6b9583c5d4356332fed4ed1fbf2fb","observation_id":"3e3958dc-469f-4e10-8d31-1021ef23885b","resolution":{"observed_at":"2026-08-07T05:06:37.515017Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2023.a","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.496470Z","title":"In: Rogers, A., Boyd-Graber, J., Okazaki, N","venue":null,"work_id":"fdfcab14-7297-40d1-a669-71b8c8d27df1","year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.393485Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:4273b150390f3158413e9840ffbe6c5c98b7a247c5079f67a6e7f6385993b8e1","observation_id":"8f9a8d39-d8c7-4485-a10f-79acd3ad1df1","resolution":{"observed_at":"2026-08-07T05:06:37.501868Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.047209Z","title":"In: Calzolari, N., Béchet, F., Blache, P., Choukri, K., Cieri, C., Declerck, T., Goggi, S., Isahara, H., Maegaard, B., Mariani, J., Mazo, H., Odijk, J., Piperidis, S","venue":null,"work_id":"35bb95eb-db47-4c66-8a5a-9de8508a75ac","year":2022},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.397644Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:4e967f743e1246dff41edf2e8aea4bc270726201c95c5cc4c0fd173e011c16a1","observation_id":"86e2eb03-195b-417b-a424-ba7d03c4439e","resolution":{"observed_at":"2026-08-07T05:06:38.051401Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.033832Z","title":"In: Proceedings of the Swiss Text Analytics Conference 2021","venue":null,"work_id":"c1c23743-b37f-4291-8f83-81c2c59208e2","year":2021},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.401819Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:4f6400ddb03135f87dab4290511387ed411e0a5678d38bbba6e0b0cda5772c85","observation_id":"10733f40-0781-4d46-80bd-6e61a9ac3530","resolution":{"observed_at":"2026-08-07T05:06:38.037955Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.405885Z","title":"In: Interspeech 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.405885Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:7551bf24fd600d9df903f747dc5e1fec6bd5cf3525ffc8dc5673720f64d8fb0a","observation_id":"ed2de395-598e-49b1-85ad-c6c732b404a4","resolution":{"observed_at":"2026-08-07T05:06:37.405885Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.410060Z","title":"In: Proceedings of the 40th International Conference on Machine Learning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.410060Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:bd19bb66eb5d439843c0f3e29be94097feb49905fbbab3b9d360e0ac10e6d3b7","observation_id":"cd581552-c1b9-4d91-9ae2-5d0a3203d0b2","resolution":{"observed_at":"2026-08-07T05:06:37.410060Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2207.00412","last_updated":"2022-11-14T10:35:45Z","snapshot_observed_at":"2026-08-03T23:17:30.052131Z","submitted_at":"2022-07-01T13:43:06Z","title":"Swiss German Speech to Text system evaluation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2207.00412","snapshot_observed_at":"2026-08-07T05:06:37.413778Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.413778Z"},"links":{"cited_paper":"/paper/2207.00412","citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:4c722baebc4f1a6ab17dbefdbed081ac7b89da6a4c83cda1918a9007e148f41c","observation_id":"7a64d7fc-545e-4c06-a1ae-ce2a4d4bad54","resolution":{"observed_at":"2026-08-07T05:06:37.413778Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.020430Z","title":"In: Ghorbel, H., Sokhn, M., Cieliebak, M., Hür- limann, M., de Salis, E., Guerne, J","venue":null,"work_id":"9ab1a0ee-53ae-4d97-a2bf-aa5937959633","year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.418138Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:b6986048e90fe855fffe9c184a9f6aa9dfc75c0ef9d9b8f127309a7921c5b7a7","observation_id":"1f520525-bc8b-4dca-8b04-3691dff1830a","resolution":{"observed_at":"2026-08-07T05:06:38.024910Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.006941Z","title":null,"venue":null,"work_id":"e7dc1afc-19ac-4a56-a144-e5154ab85cfc","year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.421916Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:5015f6f5efb85a1e8a30986775e67df56b1188f1e8648b415f699146e268bb81","observation_id":"31a1a4ff-47f5-40c7-be05-36745d931d07","resolution":{"observed_at":"2026-08-07T05:06:38.011178Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15726","last_updated":"2025-04-22T12:09:39Z","snapshot_observed_at":"2026-07-06T20:10:46.746601Z","submitted_at":"2024-12-20T09:49:02Z","title":"Fine-tuning Whisper on Low-Resource Languages for Real-World Applications","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.15726","snapshot_observed_at":"2026-08-07T05:06:37.425701Z","title":"https://doi.org/10.48550/arXiv.2412.15726, http://arxiv.org/abs/ 2412.15726, arXiv:2412.15726 [cs]","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.425701Z"},"links":{"cited_paper":"/paper/2412.15726","citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:2a9a4f4ebbfa313807c7219080af5361b6ec2abc4025a412e3cb1f4e0761c0cb","observation_id":"c243bc8e-1f80-44e3-a30b-0f469af91b36","resolution":{"observed_at":"2026-08-07T05:06:37.425701Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.429598Z","title":"In: Advances in Neural In- formation Processing Systems","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.429598Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:d25e8a47153a58602ab563b31ab38f1a2967b0a03ef94a2ed10c9f52d5b57278","observation_id":"8490111f-f9a8-43c2-829f-c5ec3883ba90","resolution":{"observed_at":"2026-08-07T05:06:37.429598Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.433193Z","title":"In: Liu, Q., Schlangen, D","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.433193Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:07a900a94ee84f965c7339a8f3aa3c0df182bb6766d8ed828f11270eed69700f","observation_id":"eafbbb55-e5ca-4f16-9b79-d0df68d0e743","resolution":{"observed_at":"2026-08-07T05:06:37.433193Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T18:03:22.334480Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech"},"reference_resolution":{"displayed":30,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":19,"verified_exact":6,"verified_fuzzy":5},"total_outbound_references":30},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 30 of 30 outbound references and 0 inbound Pith citation observations for arXiv:2506.08836."}