{"as_of":"2026-08-07T19:13:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:548099ab9d062e2080de2fb11d0dbda663c21648a63d299db72354714ce94c6c","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":27,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":27,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":27,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":27,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:42:05.292113Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-10T20:37:34.442695Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-08-07T14:42:05.292113Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.17841","last_updated":"2025-05-23T12:55:33Z","snapshot_observed_at":"2026-08-07T14:37:32.901978Z","submitted_at":"2025-05-23T12:55:33Z","title":"TEDI: Trustworthy and Ethical Dataset Indicators to Analyze and Compare Dataset Documentation","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T14:42:05.292113Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2505.17841"},"observation_digest":"sha256:5586a7e07c68c1374e12eac437311768476cc6c199a72ab0915d28ea88046432","observation_id":"b4d8a558-de93-4a17-a722-b3b46e2cda41","resolution":{"observed_at":"2026-08-07T14:42:05.292113Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-08-07T12:35:25.945196Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.24561","last_updated":"2025-05-30T13:16:08Z","snapshot_observed_at":"2026-08-07T17:35:29.046794Z","submitted_at":"2025-05-30T13:16:08Z","title":"Improving Language and Modality Transfer in Translation by Character-level Modeling","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:25.945196Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2505.24561"},"observation_digest":"sha256:af40f74c433c9b1829c76c5b018619dea2c5e89c2ff8d336790f90c301d21128","observation_id":"8831570c-9f43-4ccf-bf00-09910a2d9c68","resolution":{"observed_at":"2026-08-07T12:35:25.945196Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-08-07T11:37:29.411137Z","title":"Fleurs: Few-shot learning evaluation of universal representations of speech","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.01934","last_updated":"2025-06-02T17:53:10Z","snapshot_observed_at":"2026-08-07T12:27:48.889414Z","submitted_at":"2025-06-02T17:53:10Z","title":"RoboEgo System Card: An Omnimodal Model with Native Full Duplexity","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:29.411137Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2506.01934"},"observation_digest":"sha256:bb6d5cc9947f7c0155f8e64738d13efa02e1fff51b980a7da8841bc0198c48bf","observation_id":"7ee4bfd6-b237-46f0-82fe-4747c7b343b2","resolution":{"observed_at":"2026-08-07T11:37:29.411137Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-08-06T15:41:34.664729Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.15272","last_updated":"2025-07-21T06:20:27Z","snapshot_observed_at":"2026-08-07T12:27:49.904339Z","submitted_at":"2025-07-21T06:20:27Z","title":"A2TTS: TTS for Low Resource Indian Languages","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-06T15:41:34.664729Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2507.15272"},"observation_digest":"sha256:b6169ef73c25b2e6b42243c8047c6ad3d0b2490198dcef1f0d70afe3a1fd35d6","observation_id":"c2895bc3-dd03-4a37-8cf1-82d9040065fc","resolution":{"observed_at":"2026-08-06T15:41:34.664729Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-08-06T15:14:29.325515Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.16875","last_updated":"2025-07-22T09:38:30Z","snapshot_observed_at":"2026-08-07T12:27:37.859873Z","submitted_at":"2025-07-22T09:38:30Z","title":"Technical report: Impact of Duration Prediction on Speaker-specific TTS for Indian Languages","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-06T15:14:29.325515Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2507.16875"},"observation_digest":"sha256:583ceb6f690dfb34239e27ae09cb8361111e9510a1ddd09f4a0dd5c45affbfbf","observation_id":"3c674f5f-94d6-4e13-8644-be859c9f748b","resolution":{"observed_at":"2026-08-06T15:14:29.325515Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-08-05T16:54:32.879235Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2508.17494","last_updated":"2025-08-24T19:07:59Z","snapshot_observed_at":"2026-08-07T12:27:49.384036Z","submitted_at":"2025-08-24T19:07:59Z","title":"Improving French Synthetic Speech Quality via SSML Prosody Control","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-05T16:54:32.879235Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2508.17494"},"observation_digest":"sha256:4cac314d4678ab8ae890e77dc0f51763e265afb26f74ddc6e054dd1761615833","observation_id":"4b6359f2-6a4b-4e50-976b-e496ca09749d","resolution":{"observed_at":"2026-08-05T16:54:32.879235Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":"2205.12446","doi":null,"metadata_source":"pith","pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-10T20:37:34.442695Z","title":"FLEURS: Few-shot learning evaluation of universal representations of speech","venue":"cs.CL","work_id":"591017b1-e3a5-4f06-9d81-998d68d87365","year":2022},"citing_paper":{"arxiv_id":"2512.01512","last_updated":"2026-04-13T14:21:57Z","snapshot_observed_at":"2026-07-30T01:31:44.078208Z","submitted_at":"2025-12-01T10:39:12Z","title":"MCAT: Scaling Many-to-Many Speech-to-Text Translation with MLLMs to 70 Languages","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-17T03:15:04.685150Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2512.01512"},"observation_digest":"sha256:7e58b14e2cdc8a9328c2cbd5b0d367873a0fc805d25c1bd471e4c0e3457bbf7f","observation_id":"de6847ad-3c01-4237-aed6-52f5a23364f3","resolution":{"observed_at":"2026-05-17T03:18:57.198237Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":"2205.12446","doi":null,"metadata_source":"pith","pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-10T20:37:34.442695Z","title":"FLEURS: Few-shot learning evaluation of universal representations of speech","venue":"cs.CL","work_id":"591017b1-e3a5-4f06-9d81-998d68d87365","year":2022},"citing_paper":{"arxiv_id":"2512.16378","last_updated":"2026-04-25T20:42:51Z","snapshot_observed_at":"2026-08-02T18:57:30.783881Z","submitted_at":"2025-12-18T10:21:14Z","title":"Hearing to Translate: The Effectiveness of Speech Modality Integration into LLMs","version":4},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-05-16T21:49:21.785096Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2512.16378"},"observation_digest":"sha256:2bc6cc309580482c2e4d2c3cf1b9d7d36790dfeb0410383c17e056aa45c9a044","observation_id":"f47c68f1-7d4f-4085-afb9-b4dd2c8aedce","resolution":{"observed_at":"2026-05-16T21:51:17.797958Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-08-03T12:01:59.543761Z","title":"Fleurs: Few-shot learning evaluation of universal representations of speech","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2601.06199","last_updated":"2026-06-01T03:39:22Z","snapshot_observed_at":"2026-08-06T01:59:30.134013Z","submitted_at":"2026-01-08T07:46:03Z","title":"FastSLM: Hierarchical Temporal Abstraction for Efficient Long-Form Speech Adaptation","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-03T12:01:59.543761Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2601.06199"},"observation_digest":"sha256:5a765a4d750f92524477a5c9bf47df2aac496d9ad2a327ca8db799e1b3ab62b5","observation_id":"98e7d417-f76a-4920-b990-0765c66a7a23","resolution":{"observed_at":"2026-08-03T12:01:59.543761Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-15T00:03:31.986628Z","title":"Fleurs: Few-shot learning evaluation of universal representations of speech,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2603.09725","last_updated":"2026-07-09T13:30:49Z","snapshot_observed_at":"2026-08-07T12:27:46.774534Z","submitted_at":"2026-03-10T14:32:12Z","title":"A Semi-spontaneous Dutch Speech Dataset for Speech Enhancement and Speech Recognition","version":2},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-07-15T00:03:31.986628Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2603.09725"},"observation_digest":"sha256:4d213cf10161f1871036136484dd688075589591de8c15fd5eddfb53fee0773d","observation_id":"a359ee39-045e-4391-90ac-d2c9dfef6e8f","resolution":{"observed_at":"2026-07-15T00:03:31.986628Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":"2205.12446","doi":null,"metadata_source":"pith","pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-10T20:37:34.442695Z","title":"FLEURS: Few-shot learning evaluation of universal representations of speech","venue":"cs.CL","work_id":"591017b1-e3a5-4f06-9d81-998d68d87365","year":2022},"citing_paper":{"arxiv_id":"2604.04598","last_updated":"2026-04-06T11:23:42Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-06T11:23:42Z","title":"Benchmarking Multilingual Speech Models on Pashto: Zero-Shot ASR, Script Failure, and Cross-Domain Evaluation","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-10T19:44:30.762851Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2604.04598"},"observation_digest":"sha256:efac41232d83ceba007534aaaae028f10a24e2b1ee849ebc9da12398dd990152","observation_id":"3c5020b0-ccdd-4bc9-ae71-87a0209c184a","resolution":{"observed_at":"2026-05-10T22:35:48.735609Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":"2205.12446","doi":null,"metadata_source":"pith","pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-10T20:37:34.442695Z","title":"FLEURS: Few-shot learning evaluation of universal representations of speech","venue":"cs.CL","work_id":"591017b1-e3a5-4f06-9d81-998d68d87365","year":2022},"citing_paper":{"arxiv_id":"2604.10736","last_updated":"2026-04-16T21:22:01Z","snapshot_observed_at":"2026-07-06T22:59:18.756089Z","submitted_at":"2026-04-12T17:17:54Z","title":"BlasBench: An Open Benchmark for Irish Speech Recognition","version":2},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-05-10T15:53:54.092426Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2604.10736"},"observation_digest":"sha256:3b304fc9b35e946b0cdedcb7a27335dcb87efa311accb0a7925d2b85db6e2573","observation_id":"f2e50f6c-090e-44e6-a174-a2753e2fba7c","resolution":{"observed_at":"2026-05-11T09:41:01.686680Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":"2205.12446","doi":null,"metadata_source":"pith","pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-10T20:37:34.442695Z","title":"FLEURS: Few-shot learning evaluation of universal representations of speech","venue":"cs.CL","work_id":"591017b1-e3a5-4f06-9d81-998d68d87365","year":2022},"citing_paper":{"arxiv_id":"2604.19221","last_updated":"2026-04-30T07:45:08Z","snapshot_observed_at":"2026-07-06T23:05:53.378585Z","submitted_at":"2026-04-21T08:24:55Z","title":"UAF: A Unified Audio Front-end LLM for Full-Duplex Speech Interaction","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-10T03:05:48.624864Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2604.19221"},"observation_digest":"sha256:86c0c18911f0cac4df7a9d2619efbfc834aaeb87d0f22417298e274aa86c5bb7","observation_id":"0b670554-ff01-4049-983c-a55df6dee226","resolution":{"observed_at":"2026-05-11T12:46:02.907523Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":"2205.12446","doi":null,"metadata_source":"pith","pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-10T20:37:34.442695Z","title":"FLEURS: Few-shot learning evaluation of universal representations of speech","venue":"cs.CL","work_id":"591017b1-e3a5-4f06-9d81-998d68d87365","year":2022},"citing_paper":{"arxiv_id":"2604.19565","last_updated":"2026-04-21T15:18:10Z","snapshot_observed_at":"2026-08-03T00:43:32.245770Z","submitted_at":"2026-04-21T15:18:10Z","title":"Detecting Hallucinations in SpeechLLMs at Inference Time Using Attention Maps","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-05-10T01:56:44.270045Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2604.19565"},"observation_digest":"sha256:5f159f403ef426e10ca62788dc0048f214a5dab36846cde6105e2f61ee20e874","observation_id":"f41d2b04-5582-4c5f-baab-9aa025c9e7c3","resolution":{"observed_at":"2026-05-11T13:21:04.999089Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":"2205.12446","doi":null,"metadata_source":"pith","pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-10T20:37:34.442695Z","title":"FLEURS: Few-shot learning evaluation of universal representations of speech","venue":"cs.CL","work_id":"591017b1-e3a5-4f06-9d81-998d68d87365","year":2022},"citing_paper":{"arxiv_id":"2605.13087","last_updated":"2026-06-29T10:37:57Z","snapshot_observed_at":"2026-07-06T23:24:42.799888Z","submitted_at":"2026-05-13T06:55:55Z","title":"Vividh-ASR: A Complexity-Tiered Benchmark and Optimization Dynamics for Robust Indic Speech Recognition","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-14T19:37:21.180532Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2605.13087"},"observation_digest":"sha256:58f62d89118dd0eab9466d1563155567880e608cefb789bf9619ff1318e30637","observation_id":"7675bf7d-1c01-4393-920c-9b7726f32a76","resolution":{"observed_at":"2026-05-14T19:37:51.725042Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":"2205.12446","doi":null,"metadata_source":"pith","pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-10T20:37:34.442695Z","title":"FLEURS: Few-shot learning evaluation of universal representations of speech","venue":"cs.CL","work_id":"591017b1-e3a5-4f06-9d81-998d68d87365","year":2022},"citing_paper":{"arxiv_id":"2605.13087","last_updated":"2026-06-29T10:37:57Z","snapshot_observed_at":"2026-07-06T23:24:42.799888Z","submitted_at":"2026-05-13T06:55:55Z","title":"Vividh-ASR: A Complexity-Tiered Benchmark and Optimization Dynamics for Robust Indic Speech Recognition","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-06-30T21:54:06.330240Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2605.13087"},"observation_digest":"sha256:c85101954c13adef5f9147b17c9772994e6ea88652c1f6148a19a3624b00b776","observation_id":"f68da92d-ff51-4373-9c11-d06c92d07b0f","resolution":{"observed_at":"2026-06-30T21:55:05.379459Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":"2205.12446","doi":null,"metadata_source":"pith","pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-10T20:37:34.442695Z","title":"FLEURS: Few-shot learning evaluation of universal representations of speech","venue":"cs.CL","work_id":"591017b1-e3a5-4f06-9d81-998d68d87365","year":2022},"citing_paper":{"arxiv_id":"2605.23463","last_updated":"2026-05-22T10:24:50Z","snapshot_observed_at":"2026-08-02T21:50:12.344912Z","submitted_at":"2026-05-22T10:24:50Z","title":"StepAudio 2.5 Technical Report","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-25T02:52:22.610397Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2605.23463"},"observation_digest":"sha256:6cbbfec0b8e15662716a0cd7e90d4870360579601309987921459bea0133d9b6","observation_id":"49a49c25-9c27-466e-ae3e-3b83d2481bb9","resolution":{"observed_at":"2026-05-25T02:55:16.447086Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-13T08:22:44.347941Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2605.23912","last_updated":"2026-04-08T23:43:46Z","snapshot_observed_at":"2026-08-03T04:32:52.637811Z","submitted_at":"2026-04-08T23:43:46Z","title":"Raon-Speech Technical Report","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-07-13T08:22:44.347941Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2605.23912"},"observation_digest":"sha256:d672514d0ee681603d22d9daa41527bbfb31d1eb139e91ed7fa3c4fab23d4a2e","observation_id":"ba3153f9-1fff-43e4-b91e-ca22e7eeda19","resolution":{"observed_at":"2026-07-13T08:22:44.347941Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":"2205.12446","doi":null,"metadata_source":"pith","pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-10T20:37:34.442695Z","title":"FLEURS: Few-shot learning evaluation of universal representations of speech","venue":"cs.CL","work_id":"591017b1-e3a5-4f06-9d81-998d68d87365","year":2022},"citing_paper":{"arxiv_id":"2605.25596","last_updated":"2026-08-02T16:42:39Z","snapshot_observed_at":"2026-08-07T12:27:50.426761Z","submitted_at":"2026-05-25T08:47:33Z","title":"Multilingual Phonological Feature Recognition with Self-Supervised Speech Models","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-06-29T21:35:03.961889Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2605.25596"},"observation_digest":"sha256:c7c0a22a0b093c6d352f89d4e24690f0f3eecd7319084be664d0da2ab1fc6852","observation_id":"499d8a2b-4a8d-4b5e-8845-cc7f2a431d92","resolution":{"observed_at":"2026-06-29T21:43:59.729604Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-08-04T05:02:46.329258Z","title":"Fleurs: Few-shot learning evaluation of universal representations of speech,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2605.25596","last_updated":"2026-08-02T16:42:39Z","snapshot_observed_at":"2026-08-07T12:27:50.426761Z","submitted_at":"2026-05-25T08:47:33Z","title":"Multilingual Phonological Feature Recognition with Self-Supervised Speech Models","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-04T05:02:46.329258Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2605.25596"},"observation_digest":"sha256:de5fa6103c2cc9eca71129112eb46350dca009fca70d5a7068f705d4ecb22787","observation_id":"aede9838-eb4c-4e07-8eda-5847a892c127","resolution":{"observed_at":"2026-08-04T05:02:46.329258Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":"2205.12446","doi":null,"metadata_source":"pith","pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-10T20:37:34.442695Z","title":"FLEURS: Few-shot learning evaluation of universal representations of speech","venue":"cs.CL","work_id":"591017b1-e3a5-4f06-9d81-998d68d87365","year":2022},"citing_paper":{"arxiv_id":"2605.28211","last_updated":"2026-05-27T09:30:36Z","snapshot_observed_at":"2026-07-06T23:37:50.157617Z","submitted_at":"2026-05-27T09:30:36Z","title":"When Helpful Context Leaks: Privacy Risks in Domain-Adapted ASR","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-29T12:35:08.845406Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2605.28211"},"observation_digest":"sha256:c4fb8bcd02650282e9d38d0b6e706d73b668671dabc756ab029f3e6009d4e425","observation_id":"336777c9-5e30-488c-a644-23d88ad7bc59","resolution":{"observed_at":"2026-06-29T12:43:25.969323Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":"2205.12446","doi":null,"metadata_source":"pith","pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-10T20:37:34.442695Z","title":"FLEURS: Few-shot learning evaluation of universal representations of speech","venue":"cs.CL","work_id":"591017b1-e3a5-4f06-9d81-998d68d87365","year":2022},"citing_paper":{"arxiv_id":"2606.09335","last_updated":"2026-06-08T11:03:25Z","snapshot_observed_at":"2026-08-07T05:59:53.128999Z","submitted_at":"2026-06-08T11:03:25Z","title":"Factors affecting ASR performance: A study using state of the art ASR models in Indic Languages","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-06-27T15:10:03.282212Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2606.09335"},"observation_digest":"sha256:20ab554ff4f085275c4c645dbcd787377a184cba14423f158d6ac668d4227984","observation_id":"88655144-e28f-404b-8213-18b620ead68a","resolution":{"observed_at":"2026-07-03T03:37:35.500051Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-11T09:35:44.381743Z","title":"FLEURS: Few- shot learning evaluation of universal representations of speech,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.05051","last_updated":"2026-07-06T13:25:33Z","snapshot_observed_at":"2026-08-06T16:22:46.672614Z","submitted_at":"2026-07-06T13:25:33Z","title":"Listen, Think, Transcribe: Continuous Latent Test-Time Scaling for ASR","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-07-11T09:35:44.381743Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2607.05051"},"observation_digest":"sha256:176b20227f02b74d2b835f6d4d130203441a070174d8ee534b5a3068e7557d4d","observation_id":"bab88ca9-b577-4721-8e2b-fbe6da7cfdfa","resolution":{"observed_at":"2026-07-11T09:35:44.381743Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":"2205.12446","doi":null,"metadata_source":"pith","pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-07-10T20:37:34.442695Z","title":"FLEURS: Few-shot learning evaluation of universal representations of speech","venue":"cs.CL","work_id":"591017b1-e3a5-4f06-9d81-998d68d87365","year":2022},"citing_paper":{"arxiv_id":"2607.06827","last_updated":"2026-07-07T21:45:04Z","snapshot_observed_at":"2026-07-11T23:18:39.332726Z","submitted_at":"2026-07-07T21:45:04Z","title":"Compress the Cache, Not the Speech Embedding: KV Compression for Efficient Speech LLMs","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-07-10T20:30:11.127007Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2607.06827"},"observation_digest":"sha256:b18dad4697f3a6057e3aa325d60b4a0bdc61f2c31244af03cd20dc2d94d638c2","observation_id":"fb2b9101-cb2b-4880-9cb4-6c6957330305","resolution":{"observed_at":"2026-07-10T20:37:34.444265Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-08-01T17:45:41.785420Z","title":"arXiv preprint arXiv:2205.12446 , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.17544","last_updated":"2026-07-20T04:42:07Z","snapshot_observed_at":"2026-08-07T11:10:34.915152Z","submitted_at":"2026-07-20T04:42:07Z","title":"X-Translator: A Real-Time Multilingual Speaker-Aware Speech-to-Speech Translation System","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-01T17:45:41.785420Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2607.17544"},"observation_digest":"sha256:972ba96e3543ab310853aa37871ae1c0012ed63c4100c2231d47ba5aa5047e15","observation_id":"b9f8c728-53b6-43b0-8ad9-6951928f0dd8","resolution":{"observed_at":"2026-08-01T17:45:41.785420Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-08-01T05:50:31.979267Z","title":"Fleurs: Few-shot learning evaluation of universal representations of speech,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.22100","last_updated":"2026-07-24T08:52:44Z","snapshot_observed_at":"2026-08-06T21:56:48.867318Z","submitted_at":"2026-07-24T08:52:44Z","title":"MEUSLI: a Multilingual Projector for LLM-based ASR and Beyond","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-01T05:50:31.979267Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2607.22100"},"observation_digest":"sha256:878872c904f9de8318480abe52bc728240a25017413287b3dbc857455ae2c683","observation_id":"c1eec3a5-1a28-4ebe-9e99-2a11e7f5e83a","resolution":{"observed_at":"2026-08-01T05:50:31.979267Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12446","snapshot_observed_at":"2026-08-03T10:13:54.959703Z","title":"arXiv preprint arXiv:2205.12446 (2022)","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.29279","last_updated":"2026-07-31T10:50:14Z","snapshot_observed_at":"2026-08-06T04:02:53.771607Z","submitted_at":"2026-07-31T10:50:14Z","title":"ParaASR: Multi-Token Prediction for Fast and Long-Context LLM-Based Speech Recognition","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-03T10:13:54.959703Z"},"links":{"cited_paper":"/paper/2205.12446","citing_paper":"/paper/2607.29279"},"observation_digest":"sha256:b07f37e73f1635cd3093501cd81edf175a449aca64a399bce7949871994d92f9","observation_id":"0a00891e-d3f5-4e73-b27d-8578017b231d","resolution":{"observed_at":"2026-08-03T10:13:54.959703Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2205.12446/citation-record","integrity":"/paper/2205.12446/integrity","json":"/paper/2205.12446/citation-record.json","paper":"/paper/2205.12446"},"outbound":[],"paper":{"arxiv_id":"2205.12446","last_updated":"2022-05-25T02:29:03Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-07-06T13:13:37.143547Z","submitted_at":"2022-05-25T02:29:03Z","title":"FLEURS: Few-shot Learning Evaluation of Universal Representations of Speech"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 27 inbound Pith citation observations for arXiv:2205.12446."}