{"as_of":"2026-08-08T07:00:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:02a0616cdbddc156086b00f5cf3fff45fb126698edcd92b3d0d82de717503a8b","coverage":[{"denominator":6,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":6,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T19:14:11.370246Z","state":"measured"},{"denominator":6,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":6,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2507.06116/citation-record","integrity":"/paper/2507.06116/integrity","json":"/paper/2507.06116/citation-record.json","paper":"/paper/2507.06116"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:14:11.786260Z","title":"wav2vec 2.0: A framework for self- supervised learning of speech representations","venue":null,"work_id":"b7228910-6ede-49b9-8bc2-f989462310e0","year":2020},"citing_paper":{"arxiv_id":"2507.06116","last_updated":"2025-07-08T16:00:13Z","snapshot_observed_at":"2026-08-06T19:08:17.275644Z","submitted_at":"2025-07-08T16:00:13Z","title":"Speech Quality Assessment Model Based on Mixture of Experts: System-Level Performance Enhancement and Utterance-Level Challenge Analysis","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T19:14:10.817559Z"},"links":{"citing_paper":"/paper/2507.06116"},"observation_digest":"sha256:6379d363d90cf68c23afadce03a889d48aedd592a0639e77f33b1eb81d536b04","observation_id":"1dafe0bb-dc99-4971-b496-c82e5ddeda11","resolution":{"observed_at":"2026-08-06T19:14:11.877014Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:14:11.562134Z","title":"Generalization ability of mos prediction networks","venue":null,"work_id":"56ca0104-c9e9-4818-8afa-b3408284dc30","year":2022},"citing_paper":{"arxiv_id":"2507.06116","last_updated":"2025-07-08T16:00:13Z","snapshot_observed_at":"2026-08-06T19:08:17.275644Z","submitted_at":"2025-07-08T16:00:13Z","title":"Speech Quality Assessment Model Based on Mixture of Experts: System-Level Performance Enhancement and Utterance-Level Challenge Analysis","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T19:14:10.906772Z"},"links":{"citing_paper":"/paper/2507.06116"},"observation_digest":"sha256:83f3387c0c73201d96736baa8d6eadce2854c8757b08151149984ee54a375218","observation_id":"5b087b82-4a04-492e-a311-ebe5617bcdd5","resolution":{"observed_at":"2026-08-06T19:14:11.657018Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.05407","last_updated":"2024-07-09T07:42:51Z","snapshot_observed_at":"2026-07-06T18:42:34.958119Z","submitted_at":"2024-07-07T15:16:19Z","title":"CosyVoice: A Scalable Multilingual Zero-shot Text-to-speech Synthesizer based on Supervised Semantic Tokens","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.05407","snapshot_observed_at":"2026-08-06T19:14:10.969660Z","title":"Cosyvoice: A scalable multilingual zero-shot text-to- speech synthesizer based on supervised semantic tokens","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.06116","last_updated":"2025-07-08T16:00:13Z","snapshot_observed_at":"2026-08-06T19:08:17.275644Z","submitted_at":"2025-07-08T16:00:13Z","title":"Speech Quality Assessment Model Based on Mixture of Experts: System-Level Performance Enhancement and Utterance-Level Challenge Analysis","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T19:14:10.969660Z"},"links":{"cited_paper":"/paper/2407.05407","citing_paper":"/paper/2507.06116"},"observation_digest":"sha256:0746508e85f07fa5fc731ea7cd77a8949243b38e1ae4adaf8c1cd41242356139","observation_id":"589760ff-a27e-4e54-9a80-86546b8e52c0","resolution":{"observed_at":"2026-08-06T19:14:10.969660Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.17589","last_updated":"2025-05-27T07:48:34Z","snapshot_observed_at":"2026-07-06T21:29:06.781083Z","submitted_at":"2025-05-23T07:55:21Z","title":"CosyVoice 3: Towards In-the-wild Speech Generation via Scaling-up and Post-training","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.17589","snapshot_observed_at":"2026-08-06T19:14:11.092792Z","title":"Cosyvoice 3: Towards in-the-wild speech generation via scaling-up and post-training","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.06116","last_updated":"2025-07-08T16:00:13Z","snapshot_observed_at":"2026-08-06T19:08:17.275644Z","submitted_at":"2025-07-08T16:00:13Z","title":"Speech Quality Assessment Model Based on Mixture of Experts: System-Level Performance Enhancement and Utterance-Level Challenge Analysis","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T19:14:11.092792Z"},"links":{"cited_paper":"/paper/2505.17589","citing_paper":"/paper/2507.06116"},"observation_digest":"sha256:55442abdc90ba58142d77f2c059834e8fa10dd138e8ec5fb8cfcd084781711c8","observation_id":"feecbaed-154d-41c9-9cef-86146d1a9a3e","resolution":{"observed_at":"2026-08-06T19:14:11.092792Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.03283","last_updated":"2025-04-11T07:36:53Z","snapshot_observed_at":"2026-07-06T19:10:48.519252Z","submitted_at":"2024-09-05T06:48:02Z","title":"FireRedTTS: A Foundation Text-To-Speech Framework for Industry-Level Generative Speech Applications","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.03283","snapshot_observed_at":"2026-08-06T19:14:11.205847Z","title":"Fireredtts: A foundation text-to-speech framework for industry-level generative speech applications","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.06116","last_updated":"2025-07-08T16:00:13Z","snapshot_observed_at":"2026-08-06T19:08:17.275644Z","submitted_at":"2025-07-08T16:00:13Z","title":"Speech Quality Assessment Model Based on Mixture of Experts: System-Level Performance Enhancement and Utterance-Level Challenge Analysis","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T19:14:11.205847Z"},"links":{"cited_paper":"/paper/2409.03283","citing_paper":"/paper/2507.06116"},"observation_digest":"sha256:fe6729557f8ec4719321a56f6cb301f6a0e445343ce695d6737b5acb708b66c7","observation_id":"30f9b39f-def3-4812-afda-b754c9957879","resolution":{"observed_at":"2026-08-06T19:14:11.205847Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:14:11.370246Z","title":"Adaptive mixtures of local experts","venue":null,"work_id":null,"year":1991},"citing_paper":{"arxiv_id":"2507.06116","last_updated":"2025-07-08T16:00:13Z","snapshot_observed_at":"2026-08-06T19:08:17.275644Z","submitted_at":"2025-07-08T16:00:13Z","title":"Speech Quality Assessment Model Based on Mixture of Experts: System-Level Performance Enhancement and Utterance-Level Challenge Analysis","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T19:14:11.370246Z"},"links":{"citing_paper":"/paper/2507.06116"},"observation_digest":"sha256:aedb567c8774390d34c5af45111f512d5a9b4e002b3a98ba3dc8a06e96081f43","observation_id":"e3f33b4f-16bf-4ef7-871b-4e221029ce10","resolution":{"observed_at":"2026-08-06T19:14:11.370246Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2507.06116","last_updated":"2025-07-08T16:00:13Z","latest_version":1,"primary_category":"cs.SD","snapshot_observed_at":"2026-08-06T19:08:17.275644Z","submitted_at":"2025-07-08T16:00:13Z","title":"Speech Quality Assessment Model Based on Mixture of Experts: System-Level Performance Enhancement and Utterance-Level Challenge Analysis"},"reference_resolution":{"displayed":6,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":4,"verified_exact":0,"verified_fuzzy":2},"total_outbound_references":6},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 6 of 6 outbound references and 0 inbound Pith citation observations for arXiv:2507.06116."}