{"as_of":"2026-08-19T06:38:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:d7fd2324b4987c965a077776e0d19f8edbef25b965c611c738f3023eeaa5ce27","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":5,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":5,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-19T06:32:44.657259+00:00","state":"measured"},{"denominator":5,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":5,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:51:16.525195Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T08:06:48.129903Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2308.10390","last_updated":"2024-04-18T08:13:58Z","snapshot_observed_at":"2026-08-16T15:07:09.479698Z","submitted_at":"2023-08-20T23:47:23Z","title":"LibriSQA: A Novel Dataset and Framework for Spoken Question Answering with Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.10390","snapshot_observed_at":"2026-08-07T14:51:16.525195Z","title":"Lib- risqa: Pioneering free-form and open-ended spoken question an- swering with a novel dataset and framework,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.17417","last_updated":"2025-05-23T03:05:47Z","snapshot_observed_at":"2026-08-17T02:24:20.568852Z","submitted_at":"2025-05-23T03:05:47Z","title":"Speechless: Speech Instruction Training Without Speech for Low Resource Languages","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T14:51:16.525195Z"},"links":{"cited_paper":"/paper/2308.10390","citing_paper":"/paper/2505.17417"},"observation_digest":"sha256:6b063276e46c95ed1a752f6a3bf0f0d67049db59b1ccb0175a0c28c3a0bf7ded","observation_id":"a6c1edec-a87e-4e7f-a463-d1e943a859ce","resolution":{"observed_at":"2026-08-07T14:51:16.525195Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.10390","last_updated":"2024-04-18T08:13:58Z","snapshot_observed_at":"2026-08-16T15:07:09.479698Z","submitted_at":"2023-08-20T23:47:23Z","title":"LibriSQA: A Novel Dataset and Framework for Spoken Question Answering with Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.10390","snapshot_observed_at":"2026-08-07T14:33:48.867451Z","title":"Lib- risqa: Pioneering free-form and open-ended spoken question an- swering with a novel dataset and framework,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.18517","last_updated":"2025-05-24T05:28:22Z","snapshot_observed_at":"2026-08-13T03:07:48.461288Z","submitted_at":"2025-05-24T05:28:22Z","title":"LiSTEN: Learning Soft Token Embeddings for Neural Audio LLMs","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T14:33:48.867451Z"},"links":{"cited_paper":"/paper/2308.10390","citing_paper":"/paper/2505.18517"},"observation_digest":"sha256:16de61800da69515d9ab1f8cb1e9c06158db560f045ebb0eb5b0ac4e022e68bb","observation_id":"ecdd04f8-aec5-413c-9845-52f45ed73967","resolution":{"observed_at":"2026-08-07T14:33:48.867451Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.10390","last_updated":"2024-04-18T08:13:58Z","snapshot_observed_at":"2026-08-16T15:07:09.479698Z","submitted_at":"2023-08-20T23:47:23Z","title":"LibriSQA: A Novel Dataset and Framework for Spoken Question Answering with Large Language Models","version":4},"cited_work":{"arxiv_id":"2308.10390","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2308.10390","snapshot_observed_at":"2026-07-02T08:06:48.129903Z","title":"Yaowei Zheng, Richong Zhang, Junhao Zhang, Yanhan Ye, Zheyan Luo, Zhangchi Feng, and Yongqiang Ma","venue":null,"work_id":"35a39988-1087-4017-80a2-c18105bbbce0","year":2023},"citing_paper":{"arxiv_id":"2507.08128","last_updated":"2025-07-28T22:53:43Z","snapshot_observed_at":"2026-08-17T22:20:20.496099Z","submitted_at":"2025-07-10T19:40:21Z","title":"Audio Flamingo 3: Advancing Audio Intelligence with Fully Open Large Audio Language Models","version":2},"reference_index":121,"source":"pdf_text","source_observed_at":"2026-05-15T03:42:44.523919Z"},"links":{"cited_paper":"/paper/2308.10390","citing_paper":"/paper/2507.08128"},"observation_digest":"sha256:ba88e189e828efb84f5e256a17f9e6cd3a024d01f82250217b3860abc4b63fb3","observation_id":"bc4ec414-e4a8-4080-b61a-b7f241eff03f","resolution":{"observed_at":"2026-05-15T03:42:44.731170Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.10390","last_updated":"2024-04-18T08:13:58Z","snapshot_observed_at":"2026-08-16T15:07:09.479698Z","submitted_at":"2023-08-20T23:47:23Z","title":"LibriSQA: A Novel Dataset and Framework for Spoken Question Answering with Large Language Models","version":4},"cited_work":{"arxiv_id":"2308.10390","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2308.10390","snapshot_observed_at":"2026-07-02T08:06:48.129903Z","title":"Yaowei Zheng, Richong Zhang, Junhao Zhang, Yanhan Ye, Zheyan Luo, Zhangchi Feng, and Yongqiang Ma","venue":null,"work_id":"35a39988-1087-4017-80a2-c18105bbbce0","year":2023},"citing_paper":{"arxiv_id":"2606.04730","last_updated":"2026-06-03T11:13:37Z","snapshot_observed_at":"2026-08-16T11:18:55.350023Z","submitted_at":"2026-06-03T11:13:37Z","title":"Multilingual Long-Form Speech Instruction Following: KIT's Submission to IWSLT 2026","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-06-28T06:22:10.945176Z"},"links":{"cited_paper":"/paper/2308.10390","citing_paper":"/paper/2606.04730"},"observation_digest":"sha256:7442b16d54fc1ae15c94b6099020cf0f908c8d073cdb05351b4e4863be8bbb6c","observation_id":"3064d83b-d41e-474f-bcc2-7a646b8e3d93","resolution":{"observed_at":"2026-07-02T08:06:48.131615Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.10390","last_updated":"2024-04-18T08:13:58Z","snapshot_observed_at":"2026-08-16T15:07:09.479698Z","submitted_at":"2023-08-20T23:47:23Z","title":"LibriSQA: A Novel Dataset and Framework for Spoken Question Answering with Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.10390","snapshot_observed_at":"2026-08-04T19:02:41.528122Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.01881","last_updated":"2026-08-03T08:24:54Z","snapshot_observed_at":"2026-08-16T18:27:51.029961Z","submitted_at":"2026-08-03T08:24:54Z","title":"Hear, Invoke, and Understand: A Skill-Calling Multimodal Agent for Large Audio Language Models","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-04T19:02:41.528122Z"},"links":{"cited_paper":"/paper/2308.10390","citing_paper":"/paper/2608.01881"},"observation_digest":"sha256:0af091c96b0ebf6dbb1eee5118b0f0ef2be085c062b65b397f050da6b3985709","observation_id":"9c752120-f0a1-4d5f-b55f-f14617d2bbb0","resolution":{"observed_at":"2026-08-04T19:02:41.528122Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2308.10390/citation-record","integrity":"/paper/2308.10390/integrity","json":"/paper/2308.10390/citation-record.json","paper":"/paper/2308.10390"},"outbound":[],"paper":{"arxiv_id":"2308.10390","last_updated":"2024-04-18T08:13:58Z","latest_version":4,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-16T15:07:09.479698Z","submitted_at":"2023-08-20T23:47:23Z","title":"LibriSQA: A Novel Dataset and Framework for Spoken Question Answering with Large Language Models"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"thesis":"As of 19 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 5 inbound Pith citation observations for arXiv:2308.10390."}