{"as_of":"2026-08-13T12:46:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:5ba298252edd7aac4053991cb40fc91e6710666a52a410b02a6f9e78d7c5e296","coverage":[{"denominator":27,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":27,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-10T12:07:56.969101Z","state":"measured"},{"denominator":28,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":28,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-10T12:07:56.969101Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-05-10T12:10:21.825104Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"cited_work":{"arxiv_id":"2604.13528","doi":null,"metadata_source":"pith","pith_arxiv_id":"2604.13528","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","venue":"eess.AS","work_id":"10344d81-84e6-4255-96ca-e286c67acb08","year":2026},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"cited_paper":"/paper/2604.13528","citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:df1f40102dff3244ba1e927f670776a0ce04e2ac548e2dc5bb2536bad4f6968b","observation_id":"f2378383-e264-472e-883b-31ec387fa8d9","resolution":{"observed_at":"2026-05-10T12:10:21.827452Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2604.13528/citation-record","integrity":"/paper/2604.13528/integrity","json":"/paper/2604.13528/citation-record.json","paper":"/paper/2604.13528"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"cited_work":{"arxiv_id":"2604.13528","doi":null,"metadata_source":"pith","pith_arxiv_id":"2604.13528","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","venue":"eess.AS","work_id":"10344d81-84e6-4255-96ca-e286c67acb08","year":2026},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"cited_paper":"/paper/2604.13528","citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:df1f40102dff3244ba1e927f670776a0ce04e2ac548e2dc5bb2536bad4f6968b","observation_id":"f2378383-e264-472e-883b-31ec387fa8d9","resolution":{"observed_at":"2026-05-10T12:10:21.827452Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Zero-shot GatherMOS Given an input waveformx∈R T , we extract several acous- tic descriptors that summarize temporal, spectral, and percep- tual information","venue":null,"work_id":"de7a11ca-e5ad-4800-8d33-38761ea22fba","year":null},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:8441dd210d1a185e0da951d5465d20c8c71a57eb69b904326b50a33d20857bec","observation_id":"b93204a4-42ee-4abd-81bd-12d9e014daa4","resolution":{"observed_at":"2026-05-19T12:13:07.450226Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Experimental setup The proposed approaches are evaluated on the V oiceBank- DEMAND dataset [15], which is also included in the test evaluation of the V oiceMOS Challenge 2024 [16]","venue":null,"work_id":"f25a34db-b801-49f1-bcb9-2180a7fdbb0c","year":2024},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:f3799787f5fdc2201b678e5207b20b347a7c23322285b38a87e993c52b25e989","observation_id":"ec545162-3080-4a39-8384-7ce234ef6cd4","resolution":{"observed_at":"2026-05-19T12:13:07.452750Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"By leveraging the reasoning capabilities of large language models, GatherMOS integrates these diverse signals to produce more reliable MOS pre- dictions","venue":null,"work_id":"f9cb3c00-ca65-4fe9-8127-4337524a9a98","year":null},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:4b3c1cb7d5cea410db73a1198139bbd240b73fee1be2065032962139a866ee8a","observation_id":"f046e69a-b2a9-4263-9684-6b1edf672019","resolution":{"observed_at":"2026-05-19T12:13:07.485513Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"839d654e-bd20-44a7-846d-44c82cb4019d","year":2013},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:18f02e7c081164bda0438e7637167b7fae0968a5ffacc988ffcf2c4be8b2f637","observation_id":"15d11bbd-f636-4dc1-9bd6-ac1a48edd0e5","resolution":{"observed_at":"2026-05-19T12:13:07.448984Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The Hearing-Aid Speech Quality Index (HASQI) Version 2","venue":null,"work_id":"8216cb76-b032-4fd4-91d8-27719166ec6f","year":2014},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:849fa70643ac2427df434458440a146ef7d202c22751481c4dd018905db3e4ba","observation_id":"5c5a32a8-5bd0-4c6b-94ea-c4996fb3c7c4","resolution":{"observed_at":"2026-05-19T12:13:07.454892Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Perceptual ob- jective listening quality assessment (POLQA), the third generation ITU-T standard for end-to-end speech qual- ity measurement part i—temporal alignment","venue":null,"work_id":"6a5b41e8-b6e6-429f-a84c-0a760988c645","year":2013},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:c3b3b3db8074eb5dc25e5b98fcab34c50f69246c5044da3c439e1977de23a005","observation_id":"8b836792-62c9-4c14-a7bf-f5b641455344","resolution":{"observed_at":"2026-05-19T12:13:07.447806Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"MOSNet: Deep learning-based objective assessment for voice conver- sion","venue":null,"work_id":"e48558f7-60e2-4feb-b79c-2fb8e5b7fd71","year":2019},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:e1ae66fd7a8e3e616cee0502e1863e5fcd804d8ebd526d1615328a73636ab40a","observation_id":"72f74347-0c75-4e19-abc5-8c52cc6853b8","resolution":{"observed_at":"2026-05-19T12:13:07.455099Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Deep learning-based non-intrusive multi- objective speech assessment model with cross-domain features","venue":null,"work_id":"ad672bb4-6cb7-4e03-9e61-eeeff28e6e07","year":2023},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:b1d1f8ddd8e4be956b2c60c76c0b32743079df3e915b3eafa4967029cfb921b9","observation_id":"48a352aa-c9eb-40a2-a5c2-3bb74c439632","resolution":{"observed_at":"2026-05-19T12:13:07.461260Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Self-supervised speech quality estimation and en- hancement using only clean speech","venue":null,"work_id":"7116e21f-b3b1-4960-bb82-09ef54bae63b","year":2024},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:5fad1c59441573af832e599a9b37c9a4e45805e5b08440fd868d524d1dafae97","observation_id":"05cf4dfc-23e8-40bd-bc91-89c9fe66c477","resolution":{"observed_at":"2026-05-19T12:13:07.483020Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Generalization ability of MOS prediction networks","venue":null,"work_id":"a23ccd52-4d9b-4684-89db-415b789fb5b4","year":2022},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:5bca33e84e9c4ec930168ff9d5ca7cc91407a33aae69dc2045e711255de3d3ae","observation_id":"beccd227-6142-4d15-a715-f276297b1361","resolution":{"observed_at":"2026-05-19T12:13:07.447009Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Enabling auditory large language models for automatic speech quality evaluation","venue":null,"work_id":"f11cc1e1-f6cd-4730-9e28-8316fc6552a1","year":2025},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:450bc8ca7ffbf90297b595a142ab02520e7b9dae9eca5b7a3f20ee29be6fafc7","observation_id":"3c78827f-fae6-4291-8f0c-65de74604cb9","resolution":{"observed_at":"2026-05-19T12:13:07.460432Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Audio large language models can be descriptive speech quality evaluators","venue":null,"work_id":"1366be5e-723f-498d-86df-4967721fb68b","year":2025},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:1d2d298ea16593ad2f9e4266b7d5b8a46d6c79096becacfe371e4d37453ad83c","observation_id":"d4da0c24-7136-45e8-87be-2af135d1488b","resolution":{"observed_at":"2026-05-19T12:13:07.462754Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Wav2vec 2.0: A framework for self-supervised learn- ing of speech representations","venue":null,"work_id":"80bb56d6-6108-4005-91a1-3459f03e56b7","year":2020},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:3368364837c0a32005e737f041e14f7b2e142e3719ce11616fbca6fdc4adbc7d","observation_id":"20ca4490-1770-4895-8b33-b45f89265b49","resolution":{"observed_at":"2026-05-19T12:13:07.496647Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Robust speech recogni- tion via large-scale weak supervision","venue":null,"work_id":"aaac27f3-7e4a-4e6d-bebf-6b40c809cf24","year":2023},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:1ae2bd60c2fb4732a4eb39eebfb9b633f4900ce058f0c3b726d2fa86a48f7dba","observation_id":"b0a6d23f-fd4b-4f0d-a92c-90c48a1c6714","resolution":{"observed_at":"2026-05-19T12:13:07.502035Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.23873","last_updated":"2025-03-31T09:23:52Z","snapshot_observed_at":"2026-08-07T16:25:43.332778Z","submitted_at":"2025-03-31T09:23:52Z","title":"Exploring In-Context Learning Capabilities of ChatGPT for Pathological Speech Detection","version":1},"cited_work":{"arxiv_id":"2503.23873","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.23873","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Exploring in-context learning capabilities of ChatGPT for patho- logical speech detection","venue":null,"work_id":"4021c621-b50e-445c-81f6-a1e7e5444575","year":2025},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"cited_paper":"/paper/2503.23873","citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:a27d3ac51bbc75c9739b039399532f979221858e337df9839f5b980e4f07be6b","observation_id":"94c5c0ad-0a2c-41b5-b8ae-06d39c32bd4c","resolution":{"observed_at":"2026-05-10T12:10:21.832107Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A study on zero-shot non-intrusive speech assessment using large language models","venue":null,"work_id":"66a6e021-223d-48a0-b041-90c6e39fde16","year":2025},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:7d4dc205cbeb4759aefe98b734360ed154b6884924a41220d4675f3c4cafba0b","observation_id":"be062fdf-0a2f-4bd1-8192-52dad3ce4bf0","resolution":{"observed_at":"2026-05-19T12:13:07.499024Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"DNSMOS: A non-intrusive perceptual objective speech quality metric to evaluate noise suppressors","venue":null,"work_id":"83285f76-e0ff-4399-a2ce-1ffda721a9e3","year":2021},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:a4b4a30b260dd66d4c647c755e2abe62946696de81b80010f0629276394cfb52","observation_id":"6734d477-7a10-4f65-a3d9-a8d1048519d9","resolution":{"observed_at":"2026-05-19T12:13:07.487917Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Investigating RNN-based speech enhancement methods for noise-robust text-to-speech","venue":null,"work_id":"0d9c775e-878e-4b03-8e0b-754b305767c7","year":2016},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:413b134e5c9766bbcea9182f3e459f3900190eb2a15e14bd30f4c1c20a0cbf8c","observation_id":"a3e4a0f8-85bc-458f-80a9-360a2fa7431b","resolution":{"observed_at":"2026-05-19T12:13:07.494166Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The V oicemos challenge 2024: Beyond speech quality pre- diction","venue":null,"work_id":"d29dc442-8af2-444c-9166-4ee7ddaf916b","year":2024},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:9346c7975f10d08ceea7248d7a504a9cc9b742fa48301f2284c20eb429b59c0a","observation_id":"f8738367-8df3-42c1-96a2-5438b8804045","resolution":{"observed_at":"2026-05-19T12:13:07.478514Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Boosting Self-Supervised Embeddings for Speech Enhancement","venue":null,"work_id":"aeee7222-7d01-4ced-9d25-e1ee902924c0","year":2022},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:c458870c85ad58af4a7e2e50bc47f882b8e45de52b50b3b1f88428015a406f6a","observation_id":"d6d7055e-6b88-47cf-9ca7-21e9ab5c4664","resolution":{"observed_at":"2026-05-19T12:13:07.478236Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"MP-SENet: A Speech Enhancement Model with Parallel Denoising of Magni- tude and Phase Spectra","venue":null,"work_id":"2f363db2-611c-4835-b51b-043612312d71","year":2023},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:03fa6314737b5edf947c9283ef2bcb535680a9e1a64f5f54fb1e74b5d6b971ac","observation_id":"c27fffb1-f7c3-4973-9c8f-b800a0858ca1","resolution":{"observed_at":"2026-05-19T12:13:07.480684Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"CMGAN: Conformer-based Metric GAN for Speech Enhance- ment","venue":null,"work_id":"dd57de46-2983-4392-b4b4-ab04f79ea00e","year":2022},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:02590949400f63736f12bf54dd09374d1d0851909bbd5b9a285063b09684ca6f","observation_id":"f6836bb7-3a8c-442e-b394-a44d1bc3b8c0","resolution":{"observed_at":"2026-05-19T12:13:07.473817Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Real Time Speech Enhancement in the Waveform Domain","venue":null,"work_id":"4519c43b-9daf-4877-a8fb-4f95699286ba","year":2020},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:c2ff9a057cc2eec4fffa2e39289e126f3515e004a4b0d5ef08473103c298bff3","observation_id":"c7086e7e-c440-460d-bce3-90c9dc60bd4a","resolution":{"observed_at":"2026-05-19T12:13:07.480840Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The proof and measurement of associ- ation between two things","venue":null,"work_id":"aa856cd8-47b1-460d-b613-8a2b5f6c5b23","year":1904},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:1239c020edbaa8d4f82411a63693841a660d9553226c0592eaff8e0fef712a00","observation_id":"a3a7be59-edfc-4343-b2af-cc5d8e24f3e2","resolution":{"observed_at":"2026-05-19T12:13:07.485711Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The CHiME-7 UDASE task: Unsupervised domain adaptation for conversational speech enhance- ment","venue":null,"work_id":"64557be1-ff94-428e-bc9c-4947087db769","year":2023},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:241d37ed65d745b7c04ff3139e0b47e27ab96f379d4811a794febd7443797dcd","observation_id":"646e73e8-d170-47e0-9ba9-60d362d2bd18","resolution":{"observed_at":"2026-05-19T12:13:07.471172Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Objective and subjec- tive evaluation of speech enhancement methods in the UDASE task of the 7th CHiME challenge","venue":null,"work_id":"c260b95b-e939-4df5-b5fb-56261930ba6d","year":2025},"citing_paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-10T12:07:56.969101Z"},"links":{"citing_paper":"/paper/2604.13528"},"observation_digest":"sha256:c7bcb5a92e61e8bf3bcf81acee0993a42906fd655ecdca2f9ba85dcc6c01e213","observation_id":"b2241d74-58dc-4929-aa6e-279bc8057aec","resolution":{"observed_at":"2026-05-19T12:13:07.490583Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2604.13528","last_updated":"2026-04-15T06:23:20Z","latest_version":1,"primary_category":"eess.AS","snapshot_observed_at":"2026-07-06T23:01:32.542040Z","submitted_at":"2026-04-15T06:23:20Z","title":"Few-Shot and Pseudo-Label Guided Speech Quality Evaluation with Large Language Models"},"reference_resolution":{"displayed":27,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":1,"verified_exact":1,"verified_fuzzy":24},"total_outbound_references":27},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 27 of 27 outbound references and 1 inbound Pith citation observation for arXiv:2604.13528."}