{"as_of":"2026-08-09T06:50:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:bb7dc2e593525377c3e3d41b23f268f99524e9764f9db5eb27a2d635060af160","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":14,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":14,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":14,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":14,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T00:33:19.577773Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":1,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-07T00:33:19.577773Z","title":"Llama-omni2: Llm- based real-time spoken chatbot with autoregressive streaming speech synthesis, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.13642","last_updated":"2025-06-22T07:56:58Z","snapshot_observed_at":"2026-08-09T02:59:40.714095Z","submitted_at":"2025-06-16T16:06:45Z","title":"Stream-Omni: Simultaneous Multimodal Interactions with Large Language-Vision-Speech Model","version":2},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-07T00:33:19.577773Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2506.13642"},"observation_digest":"sha256:f8280067f1f4d26f5f02138b97405fd18a7c4d9155f1d1c442ef7dd9ea992536","observation_id":"f5e67540-2a28-4247-9b5c-61fabd5ceea5","resolution":{"observed_at":"2026-08-07T00:33:19.577773Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T15:52:43.998471Z","title":"Llama-omni2: Llm- based real-time spoken chatbot with autoregressive streaming speech synthesis,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.00078","last_updated":"2025-08-26T20:40:24Z","snapshot_observed_at":"2026-08-05T15:52:43.760838Z","submitted_at":"2025-08-26T20:40:24Z","title":"ChipChat: Low-Latency Cascaded Conversational Agent in MLX","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-05T15:52:43.998471Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2509.00078"},"observation_digest":"sha256:4f66400e7182d8bcbff969c5ee21a0e71d3037f3c6fba07fa49e856b6900fff0","observation_id":"dc1e401d-6272-4516-afe5-6481c35deecb","resolution":{"observed_at":"2026-08-05T15:52:43.998471Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2509.22220","last_updated":"2026-04-13T11:56:11Z","snapshot_observed_at":"2026-08-02T12:47:50.534468Z","submitted_at":"2025-09-26T11:32:51Z","title":"StableToken: A Noise-Robust Semantic Speech Tokenizer for Resilient SpeechLLMs","version":2},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-05-18T12:57:04.450462Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2509.22220"},"observation_digest":"sha256:fd05d197c217d0aa51ea8298c085588d8f45d1695b14e0674397d84104ebd477","observation_id":"1497d69c-0b13-4a0c-8c71-a5d4043ee17c","resolution":{"observed_at":"2026-05-18T13:01:24.365147Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2510.09592","last_updated":"2026-05-10T15:21:35Z","snapshot_observed_at":"2026-07-06T22:32:22.953630Z","submitted_at":"2025-10-10T17:50:59Z","title":"Mind-Paced Speaking: A Dual-Brain Approach to Real-Time Reasoning in Spoken Language Models","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-18T07:43:23.913399Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2510.09592"},"observation_digest":"sha256:79f25cc9dd441c298bd005db7da59031805e1fa89a435ffcb2dcced09caced2b","observation_id":"d503e0da-a3f1-44fe-a927-eed086ae7943","resolution":{"observed_at":"2026-05-18T07:46:03.589866Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-03T12:25:01.710490Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.03190","last_updated":"2026-07-23T19:26:16Z","snapshot_observed_at":"2026-08-03T12:24:58.228438Z","submitted_at":"2026-01-06T17:10:48Z","title":"Maximizing Local Entropy Where It Matters: Prefix-Aware Localized LLM Unlearning","version":4},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-03T12:25:01.710490Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2601.03190"},"observation_digest":"sha256:87269f95ff333a8e268bb640b2d1b6ad7ca748f63bdb4ec6c5d5347e4df8d89b","observation_id":"0b2bf6c6-73b0-46cb-9d81-f1357ac89c8f","resolution":{"observed_at":"2026-08-03T12:25:01.710490Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2604.13804","last_updated":"2026-04-15T12:39:03Z","snapshot_observed_at":"2026-07-06T23:01:41.643335Z","submitted_at":"2026-04-15T12:39:03Z","title":"Character Beyond Speech: Leveraging Role-Playing Evaluation in Audio Large Language Models via Reinforcement Learning","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-10T14:22:25.660785Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2604.13804"},"observation_digest":"sha256:2dc9d19e361bb8426087700b59597ba83aa615ea1bf0aa478b203ea726bef8df","observation_id":"2c54e787-bd4b-49e5-8bd5-f205b4baa4a4","resolution":{"observed_at":"2026-05-10T14:25:30.047240Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2604.19300","last_updated":"2026-04-21T10:05:28Z","snapshot_observed_at":"2026-08-03T21:49:52.344628Z","submitted_at":"2026-04-21T10:05:28Z","title":"HalluAudio: A Comprehensive Benchmark for Hallucination Detection in Large Audio-Language Models","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T01:34:54.375266Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2604.19300"},"observation_digest":"sha256:644d8492ef500de8aa692b71873a4fa4d72c2ee7069c30ce47e4cfe0553f8860","observation_id":"5ff634e1-5961-4926-8e3d-76138e3aa974","resolution":{"observed_at":"2026-05-11T13:31:03.386871Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2604.20842","last_updated":"2026-04-22T17:59:58Z","snapshot_observed_at":"2026-08-04T16:34:40.989674Z","submitted_at":"2026-04-22T17:59:58Z","title":"SpeechParaling-Bench: A Comprehensive Benchmark for Paralinguistic-Aware Speech Generation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-10T00:39:04.303837Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2604.20842"},"observation_digest":"sha256:8fc0f9e75b062918115a677b2bbc3163c084459c9b1df820e6b197a8b8ad957d","observation_id":"6c489d98-1bb1-4ec1-8352-cbfaa2e5f95b","resolution":{"observed_at":"2026-05-10T00:39:48.480242Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2606.01016","last_updated":"2026-05-31T05:13:32Z","snapshot_observed_at":"2026-07-06T23:41:39.172310Z","submitted_at":"2026-05-31T05:13:32Z","title":"PolySpeech-100: A Large-Scale Benchmark for Speech Understanding Across 100+ Languages and Dialects","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-06-28T17:44:07.669223Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2606.01016"},"observation_digest":"sha256:44215d4b1915c9523d3d4105767c66fe3576ca62dd7b98877b2290262a84cdc3","observation_id":"9d99baf3-d54d-4dcc-af2b-3f8950b489ac","resolution":{"observed_at":"2026-06-28T17:52:26.993555Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2606.05121","last_updated":"2026-06-03T17:26:11Z","snapshot_observed_at":"2026-08-02T17:56:34.317060Z","submitted_at":"2026-06-03T17:26:11Z","title":"Audio Interaction Model","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-28T04:57:05.062465Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2606.05121"},"observation_digest":"sha256:28acc2fb61833da69d6fba745d6c217f4d71d24cca3d211c00939c1ef0d6e923","observation_id":"023d51a6-739e-4ff4-b00a-dabae30b93a1","resolution":{"observed_at":"2026-07-02T10:46:52.382243Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2606.07433","last_updated":"2026-06-05T16:29:13Z","snapshot_observed_at":"2026-08-01T21:05:06.439607Z","submitted_at":"2026-06-05T16:29:13Z","title":"Watch, Remember, Reason: Human-View Video Understanding with MLLMs","version":1},"reference_index":127,"source":"pdf_text","source_observed_at":"2026-06-27T22:00:28.350003Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2606.07433"},"observation_digest":"sha256:135af2fa62e599ec6441208b4c28a3a396b9af401570b5edb4831fe5e85a2360","observation_id":"0f7c03a4-2614-492e-96a3-577d3af0b60c","resolution":{"observed_at":"2026-07-02T17:27:14.747442Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2606.25444","last_updated":"2026-06-24T06:15:18Z","snapshot_observed_at":"2026-08-02T19:50:32.355214Z","submitted_at":"2026-06-24T06:15:18Z","title":"Does Translation-Enhanced Speech Encoder Pre-training Affect Speech LLMs?","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-06-25T20:03:10.858349Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2606.25444"},"observation_digest":"sha256:ae11b5c17428cfa1b8a75358dac877795ac6ceb6e5b2164cd83c51ef9c03e456","observation_id":"1347e7a1-5d37-48c1-9d6b-1b492fd9f58d","resolution":{"observed_at":"2026-07-04T20:30:08.189345Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2606.30944","last_updated":"2026-06-29T21:55:18Z","snapshot_observed_at":"2026-08-07T21:37:48.909123Z","submitted_at":"2026-06-29T21:55:18Z","title":"Preserving Speech-to-Text LLM Capabilities in Speech-to-Speech Generation","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-07-01T01:01:24.536821Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2606.30944"},"observation_digest":"sha256:2c6e4e9fca829f34fa950ccff9e1b5e00936453de39e134cde2f1232c218543d","observation_id":"b9d06f2e-9fcc-459b-87e8-153712bfeec4","resolution":{"observed_at":"2026-07-01T13:05:45.816779Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-01T07:12:17.669270Z","title":"arXiv preprint arXiv:2505.02625 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.21550","last_updated":"2026-07-23T17:35:20Z","snapshot_observed_at":"2026-08-08T08:48:36.880078Z","submitted_at":"2026-07-23T17:35:20Z","title":"X$^3$-OPD: Distilling Reasoning into Large Audio-Language Models via On-Policy Alignment","version":1},"reference_index":139,"source":"arxiv_source","source_observed_at":"2026-08-01T07:12:17.669270Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2607.21550"},"observation_digest":"sha256:d5e77ea54262b55369f553366d7e9e9e60de2884d394de3ee034d451262e1571","observation_id":"93186e2e-187a-4d75-a08c-e04e9e56982b","resolution":{"observed_at":"2026-08-01T07:12:17.669270Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2505.02625/citation-record","integrity":"/paper/2505.02625/integrity","json":"/paper/2505.02625/citation-record.json","paper":"/paper/2505.02625"},"outbound":[],"paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 14 inbound Pith citation observations for arXiv:2505.02625."}