{"as_of":"2026-08-07T13:18:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:b3d831c8e737b5d354ef0ca77df698bba7a81bf8165fb48352a4848bc90db14d","coverage":[{"denominator":13,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":13,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-10T09:24:25.616750Z","state":"measured"},{"denominator":13,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":13,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2604.14806/citation-record","integrity":"/paper/2604.14806/integrity","json":"/paper/2604.14806/citation-record.json","paper":"/paper/2604.14806"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.06261","last_updated":"2025-12-19T14:25:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-07-07T17:36:04Z","title":"Gemini 2.5: Pushing the Frontier with Advanced Reasoning, Multimodality, Long Context, and Next Generation Agentic Capabilities","version":6},"cited_work":{"arxiv_id":"2507.06261","doi":"10.48550/arxiv.2503.19","metadata_source":"pith","pith_arxiv_id":"2507.06261","snapshot_observed_at":"2026-07-11T03:17:51.364436Z","title":"Gemini 2.5: Pushing the Frontier with Advanced Reasoning, Multimodality, Long Context, and Next Generation Agentic Capabilities","venue":"cs.CL","work_id":"008df105-2fdd-45d8-857a-8e35868aecb6","year":2025},"citing_paper":{"arxiv_id":"2604.14806","last_updated":"2026-04-16T09:30:13Z","snapshot_observed_at":"2026-08-01T04:51:16.142077Z","submitted_at":"2026-04-16T09:30:13Z","title":"Listen, Pause, and Reason: Toward Perception-Grounded Hybrid Reasoning for Audio Understanding","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-10T09:24:25.616750Z"},"links":{"cited_paper":"/paper/2507.06261","citing_paper":"/paper/2604.14806"},"observation_digest":"sha256:259b51696c5e744770799a5fb1638bdcdbef9e39a079f6c0b447e4992d583832","observation_id":"9149bfdd-0f17-4827-ac6a-fef947fa6f44","resolution":{"observed_at":"2026-05-10T09:28:38.899988Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11768","last_updated":"2024-06-17T17:31:01Z","snapshot_observed_at":"2026-08-01T20:00:05.737836Z","submitted_at":"2024-06-17T17:31:01Z","title":"GAMA: A Large Audio-Language Model with Advanced Audio Understanding and Complex Reasoning Abilities","version":1},"cited_work":{"arxiv_id":"2406.11768","doi":"10.48550/arxiv.2406.11768","metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11768","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"GAMA: A large audio- language model with advanced audio understanding and complex rea- soning abilities","venue":"arXiv (Cornell University)","work_id":"f407a4af-a90d-4687-b15e-70573cc84a5d","year":2024},"citing_paper":{"arxiv_id":"2604.14806","last_updated":"2026-04-16T09:30:13Z","snapshot_observed_at":"2026-08-01T04:51:16.142077Z","submitted_at":"2026-04-16T09:30:13Z","title":"Listen, Pause, and Reason: Toward Perception-Grounded Hybrid Reasoning for Audio Understanding","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-10T09:24:25.616750Z"},"links":{"cited_paper":"/paper/2406.11768","citing_paper":"/paper/2604.14806"},"observation_digest":"sha256:575f119fdf8d6abb279624bef70f85d3ecdcfa2a9f2d7dbd4f28b88d367698fe","observation_id":"e8da2d69-0366-4eef-81b8-3841b3aeff5e","resolution":{"observed_at":"2026-05-10T09:28:38.897129Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.19168","last_updated":"2024-10-24T21:20:10Z","snapshot_observed_at":"2026-07-06T19:39:23.839070Z","submitted_at":"2024-10-24T21:20:10Z","title":"MMAU: A Massive Multi-Task Audio Understanding and Reasoning Benchmark","version":1},"cited_work":{"arxiv_id":"2410.19168","doi":"10.48550/arxiv.2410.19168","metadata_source":"pith","pith_arxiv_id":"2410.19168","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"MMAU: A Massive Multi-Task Audio Understanding and Reasoning Benchmark","venue":"eess.AS","work_id":"e60f85db-636c-4830-af85-5d31ebc74a1b","year":2024},"citing_paper":{"arxiv_id":"2604.14806","last_updated":"2026-04-16T09:30:13Z","snapshot_observed_at":"2026-08-01T04:51:16.142077Z","submitted_at":"2026-04-16T09:30:13Z","title":"Listen, Pause, and Reason: Toward Perception-Grounded Hybrid Reasoning for Audio Understanding","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-10T09:24:25.616750Z"},"links":{"cited_paper":"/paper/2410.19168","citing_paper":"/paper/2604.14806"},"observation_digest":"sha256:dfde1621d297d56be4d88f31ba3ab094fe09a2e51814c96d45560fd9529d278c","observation_id":"4e9c569d-cd0f-4967-ad7e-fe587a7c3829","resolution":{"observed_at":"2026-05-13T14:04:55.589197Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1510.08484","last_updated":"2015-10-28T20:59:04Z","snapshot_observed_at":"2026-07-06T04:34:36.474437Z","submitted_at":"2015-10-28T20:59:04Z","title":"MUSAN: A Music, Speech, and Noise Corpus","version":1},"cited_work":{"arxiv_id":"1510.08484","doi":null,"metadata_source":"pith","pith_arxiv_id":"1510.08484","snapshot_observed_at":"2026-07-10T12:57:07.598373Z","title":"MUSAN: A Music, Speech, and Noise Corpus","venue":"cs.SD","work_id":"7c604702-578b-4f91-9cc6-f8aa7dbe6d26","year":2015},"citing_paper":{"arxiv_id":"2604.14806","last_updated":"2026-04-16T09:30:13Z","snapshot_observed_at":"2026-08-01T04:51:16.142077Z","submitted_at":"2026-04-16T09:30:13Z","title":"Listen, Pause, and Reason: Toward Perception-Grounded Hybrid Reasoning for Audio Understanding","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T09:24:25.616750Z"},"links":{"cited_paper":"/paper/1510.08484","citing_paper":"/paper/2604.14806"},"observation_digest":"sha256:337020fb83bd1d6681e14f23daf5e0d1a7edbe9bd972100254425065d58d0076","observation_id":"5cf35d7b-74ef-472d-ab47-04611e1d885f","resolution":{"observed_at":"2026-05-10T09:28:38.902876Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.18897","last_updated":"2025-07-25T02:44:30Z","snapshot_observed_at":"2026-08-06T14:28:34.982929Z","submitted_at":"2025-07-25T02:44:30Z","title":"HH-Codec: High Compression High-fidelity Discrete Neural Codec for Spoken Language Modeling","version":1},"cited_work":{"arxiv_id":"2507.18897","doi":"10.48550/arxiv.2507.18897","metadata_source":"arxiv_reference","pith_arxiv_id":"2507.18897","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"InICASSP 2021-2021 IEEE International Conference on Acous- tics, Speech and Signal Processing (ICASSP), pages 606–610","venue":"ArXiv.org","work_id":"40938cf7-d6bd-47f1-a2cd-45e37adda59c","year":2025},"citing_paper":{"arxiv_id":"2604.14806","last_updated":"2026-04-16T09:30:13Z","snapshot_observed_at":"2026-08-01T04:51:16.142077Z","submitted_at":"2026-04-16T09:30:13Z","title":"Listen, Pause, and Reason: Toward Perception-Grounded Hybrid Reasoning for Audio Understanding","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-10T09:24:25.616750Z"},"links":{"cited_paper":"/paper/2507.18897","citing_paper":"/paper/2604.14806"},"observation_digest":"sha256:87a2c3a327d802f810eaadcde304566167cf62c43e3cee22bfea11192519e9fe","observation_id":"ec4ecc8f-4c80-40da-9c5e-df10243578cd","resolution":{"observed_at":"2026-05-10T09:28:38.891667Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.13837","last_updated":"2025-11-24T06:11:04Z","snapshot_observed_at":"2026-07-06T21:11:34.701779Z","submitted_at":"2025-04-18T17:59:56Z","title":"Does Reinforcement Learning Really Incentivize Reasoning Capacity in LLMs Beyond the Base Model?","version":5},"cited_work":{"arxiv_id":"2504.13837","doi":"10.48550/arxiv.2504.13837","metadata_source":"pith","pith_arxiv_id":"2504.13837","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Does Reinforcement Learning Really Incentivize Reasoning Capacity in LLMs Beyond the Base Model?","venue":"cs.AI","work_id":"d854765a-e664-41c0-8655-21c4bf2e0cc4","year":2025},"citing_paper":{"arxiv_id":"2604.14806","last_updated":"2026-04-16T09:30:13Z","snapshot_observed_at":"2026-08-01T04:51:16.142077Z","submitted_at":"2026-04-16T09:30:13Z","title":"Listen, Pause, and Reason: Toward Perception-Grounded Hybrid Reasoning for Audio Understanding","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-10T09:24:25.616750Z"},"links":{"cited_paper":"/paper/2504.13837","citing_paper":"/paper/2604.14806"},"observation_digest":"sha256:4097ac03a60137d03746b396e2040414d2eb8a1b7d3305cc0035b76a36e8f948","observation_id":"0c920ddc-1ba4-41fd-a051-9f5dc541cc44","resolution":{"observed_at":"2026-05-10T09:28:38.888145Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-09T10:48:42.002839+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-09T10:48:42.002839+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"- Repetition or redundant phrasing that should be removed or marked clearly","venue":null,"work_id":"ff7ec883-3d9f-40e2-b2e3-3bfe247dd521","year":null},"citing_paper":{"arxiv_id":"2604.14806","last_updated":"2026-04-16T09:30:13Z","snapshot_observed_at":"2026-08-01T04:51:16.142077Z","submitted_at":"2026-04-16T09:30:13Z","title":"Listen, Pause, and Reason: Toward Perception-Grounded Hybrid Reasoning for Audio Understanding","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-10T09:24:25.616750Z"},"links":{"citing_paper":"/paper/2604.14806"},"observation_digest":"sha256:25b1e0235ddfc363b566e7414ced032ff83eb10d6c9b42f00899631f9830de60","observation_id":"da545adb-c160-4ea1-8504-b1544a622f74","resolution":{"observed_at":"2026-05-20T17:33:48.005950Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"- The corrections or adjustments needed (without referencing or leaking the gold standard answer text)","venue":null,"work_id":"a8055ffd-1261-408c-aee9-8a39a38d7fe4","year":null},"citing_paper":{"arxiv_id":"2604.14806","last_updated":"2026-04-16T09:30:13Z","snapshot_observed_at":"2026-08-01T04:51:16.142077Z","submitted_at":"2026-04-16T09:30:13Z","title":"Listen, Pause, and Reason: Toward Perception-Grounded Hybrid Reasoning for Audio Understanding","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T09:24:25.616750Z"},"links":{"citing_paper":"/paper/2604.14806"},"observation_digest":"sha256:6c285252643b3e569f44361c1b3e4851d62724a62b6dc2544ded180d0f245be1","observation_id":"673fa253-de4f-41f9-a84c-b0ce832deb97","resolution":{"observed_at":"2026-05-20T17:33:47.992400Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Ignore Output","venue":null,"work_id":"8a38a89b-3c80-46e4-962a-9fd27111c2b2","year":2025},"citing_paper":{"arxiv_id":"2604.14806","last_updated":"2026-04-16T09:30:13Z","snapshot_observed_at":"2026-08-01T04:51:16.142077Z","submitted_at":"2026-04-16T09:30:13Z","title":"Listen, Pause, and Reason: Toward Perception-Grounded Hybrid Reasoning for Audio Understanding","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-10T09:24:25.616750Z"},"links":{"citing_paper":"/paper/2604.14806"},"observation_digest":"sha256:8e5dfa98ed7546d6a47cbed180f3ed8a1b6f47aa814d84b445b5519ed8118642","observation_id":"c4a70b80-81c0-40b7-89ed-84d26f39ba07","resolution":{"observed_at":"2026-05-20T17:33:48.008235Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"f9b25746-5077-494a-82da-de26cfdba850","year":null},"citing_paper":{"arxiv_id":"2604.14806","last_updated":"2026-04-16T09:30:13Z","snapshot_observed_at":"2026-08-01T04:51:16.142077Z","submitted_at":"2026-04-16T09:30:13Z","title":"Listen, Pause, and Reason: Toward Perception-Grounded Hybrid Reasoning for Audio Understanding","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-10T09:24:25.616750Z"},"links":{"citing_paper":"/paper/2604.14806"},"observation_digest":"sha256:d3d13d2e59a64ce0d7352704cb16c68c273a75bedc1439b5e56243e00c87abb3","observation_id":"f726788e-93c3-41b6-852a-26315b314acc","resolution":{"observed_at":"2026-05-20T17:33:47.995105Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"28c50161-186f-452e-8ef2-be4859cfc96d","year":null},"citing_paper":{"arxiv_id":"2604.14806","last_updated":"2026-04-16T09:30:13Z","snapshot_observed_at":"2026-08-01T04:51:16.142077Z","submitted_at":"2026-04-16T09:30:13Z","title":"Listen, Pause, and Reason: Toward Perception-Grounded Hybrid Reasoning for Audio Understanding","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-10T09:24:25.616750Z"},"links":{"citing_paper":"/paper/2604.14806"},"observation_digest":"sha256:41e22503d2cf16b531addb7f6f2b7b89050c7fa23f97e04fc5622a155f105022","observation_id":"18b04251-0ea8-4f0b-8b99-a18dd9ac0fea","resolution":{"observed_at":"2026-05-20T17:33:48.003187Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"- (b) A chair: Similarly, a chair might require beveling, but it isn’t their primary focus","venue":null,"work_id":"cd2d8ff4-96e4-44b0-a737-83bcb95209d0","year":null},"citing_paper":{"arxiv_id":"2604.14806","last_updated":"2026-04-16T09:30:13Z","snapshot_observed_at":"2026-08-01T04:51:16.142077Z","submitted_at":"2026-04-16T09:30:13Z","title":"Listen, Pause, and Reason: Toward Perception-Grounded Hybrid Reasoning for Audio Understanding","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-10T09:24:25.616750Z"},"links":{"citing_paper":"/paper/2604.14806"},"observation_digest":"sha256:954ed0f9ce4642dc1b92737727d9c44fc0d647c5606da69216a6d650dc538187","observation_id":"40132a34-589a-43d7-a9d6-6ba2b9a4f4a5","resolution":{"observed_at":"2026-05-20T17:33:48.000296Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"piece A vs. B","venue":null,"work_id":"d1fe4d4e-953c-42a2-943c-398159187884","year":2024},"citing_paper":{"arxiv_id":"2604.14806","last_updated":"2026-04-16T09:30:13Z","snapshot_observed_at":"2026-08-01T04:51:16.142077Z","submitted_at":"2026-04-16T09:30:13Z","title":"Listen, Pause, and Reason: Toward Perception-Grounded Hybrid Reasoning for Audio Understanding","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-10T09:24:25.616750Z"},"links":{"citing_paper":"/paper/2604.14806"},"observation_digest":"sha256:8f4c5a44683766f0f008a9e830b5d109d195cb0b220d31a261ade73698f76c1d","observation_id":"c0d04274-27a3-4c31-8dee-5f0a37fffe64","resolution":{"observed_at":"2026-05-20T17:33:47.989468Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2604.14806","last_updated":"2026-04-16T09:30:13Z","latest_version":1,"primary_category":"cs.SD","snapshot_observed_at":"2026-08-01T04:51:16.142077Z","submitted_at":"2026-04-16T09:30:13Z","title":"Listen, Pause, and Reason: Toward Perception-Grounded Hybrid Reasoning for Audio Understanding"},"reference_resolution":{"displayed":13,"state_counts":{"malformed_identifier":0,"metadata_mismatch":5,"parse_uncertain":0,"unresolved":2,"verified_exact":1,"verified_fuzzy":5},"total_outbound_references":13},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 13 of 13 outbound references and 0 inbound Pith citation observations for arXiv:2604.14806."}