{"as_of":"2026-08-09T23:08:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:49e0a0be19ab2994c4c532126dd9bbf10834494c8855b95cef504dc34eb6e0c0","coverage":[{"denominator":41,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":41,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T12:30:37.358858Z","state":"measured"},{"denominator":42,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":42,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T12:30:33.191440Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-07T12:30:37.561387Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"cited_work":{"arxiv_id":"2505.24446","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.24446","snapshot_observed_at":"2026-08-07T12:30:37.561387Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","venue":"cs.SD","work_id":"fc2b80d3-0e17-48a4-b3e5-0c5f2daaf58e","year":2025},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:33.191440Z"},"links":{"cited_paper":"/paper/2505.24446","citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:b5eedc7b794300f515b53c218db4243183f81e89ec249e9d221473ddf004ca65","observation_id":"0baad087-ec5b-4a9d-9faf-9347e3748b26","resolution":{"observed_at":"2026-08-07T12:30:37.628141Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2505.24446/citation-record","integrity":"/paper/2505.24446/integrity","json":"/paper/2505.24446/citation-record.json","paper":"/paper/2505.24446"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:40.654341Z","title":null,"venue":null,"work_id":"9c1356a8-88c1-42e8-b31b-477796f706cb","year":2022},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:33.118167Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:2c1c24f9ea3e1fd91cfff704dfd1fbb9d6322eaf5abaffe87708c6e608810441","observation_id":"e086a301-6884-4aed-9048-d373e01098ca","resolution":{"observed_at":"2026-08-07T12:30:40.728451Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"cited_work":{"arxiv_id":"2505.24446","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.24446","snapshot_observed_at":"2026-08-07T12:30:37.561387Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","venue":"cs.SD","work_id":"fc2b80d3-0e17-48a4-b3e5-0c5f2daaf58e","year":2025},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:33.191440Z"},"links":{"cited_paper":"/paper/2505.24446","citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:b5eedc7b794300f515b53c218db4243183f81e89ec249e9d221473ddf004ca65","observation_id":"0baad087-ec5b-4a9d-9faf-9347e3748b26","resolution":{"observed_at":"2026-08-07T12:30:37.628141Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:40.539258Z","title":"GSS+G-SpatialNet","venue":null,"work_id":"8f7bab22-ad48-4f2d-8a28-64b53ce0f7d1","year":null},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:33.328563Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:c971b959e8d4b0dbc3105e4ea5079b6bbaaf41e33cbec7138451d93b5e9ce4e7","observation_id":"88aa6d18-1c61-4945-99d2-8e1dad08c2bf","resolution":{"observed_at":"2026-08-07T12:30:40.585357Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:40.411573Z","title":"Speech enhancement For G-SpatialNet training, we use only the audio data of MISP- Meeting training set","venue":null,"work_id":"6817d34b-f364-4e24-a4a5-57684b1b75a3","year":null},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:33.420778Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:4b00b764ed94b2b48951814a32199b45abbb3da249b706c78a17e3ee8dbf6abe","observation_id":"d0f03dd5-1dc2-49f1-823b-d22f35f89e62","resolution":{"observed_at":"2026-08-07T12:30:40.474508Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:40.247742Z","title":"Speech enhancement: TLS and G-SpatialNet Fig","venue":null,"work_id":"af6895c2-a6a2-4159-867f-c1ededc876e9","year":null},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:33.559942Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:aef0b5532036237b50aa51cfc3715f86370997c8cdc65bd965566907d93a1f53","observation_id":"05ef510e-2257-4d1e-8e01-cc74be98b1ef","resolution":{"observed_at":"2026-08-07T12:30:40.312656Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:40.113592Z","title":"We proposed a novel framework, TLS, which gen- erates high-quality pseudo labels for real-world meeting data, enabling the direct training of SE models in real-world scenar- ios","venue":null,"work_id":"a996f35a-599b-4b07-8ff2-d5632686977d","year":null},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:33.653493Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:0e78e68c53e38409080739b195adb569876627ff9ccef5c2b13696e627d8f9e4","observation_id":"cac23929-d5b6-44ff-8e4b-b60ea6de3646","resolution":{"observed_at":"2026-08-07T12:30:40.183988Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:39.993727Z","title":null,"venue":null,"work_id":"a3fa84ee-8c1f-44c9-80bf-6bf644904807","year":null},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:33.759723Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:4faba54be9335e4a0dbc43076db22aeaecf75d41def1f8839e41d58aab485bce","observation_id":"f75f8690-7494-45b3-8e5a-f0097c7eeb51","resolution":{"observed_at":"2026-08-07T12:30:40.051630Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:34.055974Z","title":"The first multimodal information based speech processing (misp) chal- lenge: Data, tasks, baselines and results,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:34.055974Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:8d41986d6e1af9efaf6a42215007c13962da71029aa3b7881d1bd63faef95d5e","observation_id":"fd00500a-cfd2-444c-bab2-6e34e2750ee0","resolution":{"observed_at":"2026-08-07T12:30:34.055974Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:34.249821Z","title":"The mul- timodal information based speech processing (misp) 2022 chal- lenge: Audio-visual diarization and recognition,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:34.249821Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:3648905d27417ed37364c9315455d850b58f902ddc897258d1f82807233dd26b","observation_id":"0c8b1eca-66c1-4847-b629-37e5c63ae2bb","resolution":{"observed_at":"2026-08-07T12:30:34.249821Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:34.448469Z","title":"The multimodal information based speech processing (misp) 2023 challenge: Audio-visual target speaker extraction,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:34.448469Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:86e4edc91973dbcb2250693921ede28f1e1416f7fb4d76628c4e743293d64587","observation_id":"c89daf59-499e-4c73-8329-05e78a3f7786","resolution":{"observed_at":"2026-08-07T12:30:34.448469Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:34.562630Z","title":"Front-end processing for the chime-5 dinner party scenario,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:34.562630Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:8fd0c21cf150c6d71c4b6da257656c71528f82a02dfaf9c4b268b143295d02a7","observation_id":"77534a2a-7df3-4fd9-8c78-1767047bba86","resolution":{"observed_at":"2026-08-07T12:30:34.562630Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:34.658833Z","title":"Gpu-accelerated guided source separation for meeting transcription,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:34.658833Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:4ca13d487472b0320e3b6568959fd9340fccfc8de2ca34a171d45780f3eac44d","observation_id":"7c2f5306-f406-4c2d-b5ca-2874bc0cd2a4","resolution":{"observed_at":"2026-08-07T12:30:34.658833Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:34.733183Z","title":"Chime-6 challenge: Tackling multispeaker speech recognition for unsegmented recordings,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:34.733183Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:aa323613239096399d71327cb50b47b427371b3a48e7afbd6b9543c1db2a3806","observation_id":"aaea91ec-fddd-410b-ab0e-178c6ed4e591","resolution":{"observed_at":"2026-08-07T12:30:34.733183Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.13734","last_updated":"2023-07-14T09:45:21Z","snapshot_observed_at":"2026-07-06T15:46:06.220522Z","submitted_at":"2023-06-23T18:49:20Z","title":"The CHiME-7 DASR Challenge: Distant Meeting Transcription with Multiple Devices in Diverse Scenarios","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.13734","snapshot_observed_at":"2026-08-07T12:30:34.836786Z","title":"The chime-7 dasr challenge: Distant meeting transcrip- tion with multiple devices in diverse scenarios,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:34.836786Z"},"links":{"cited_paper":"/paper/2306.13734","citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:e952e58f48511f11138f50ba3873f45977e5e5a0bccb314061f4e84c8e321b55","observation_id":"c2bc6e74-be33-4ac6-847c-86e91fb657f1","resolution":{"observed_at":"2026-08-07T12:30:34.836786Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.16447","last_updated":"2024-07-23T12:54:32Z","snapshot_observed_at":"2026-08-03T05:44:10.596840Z","submitted_at":"2024-07-23T12:54:32Z","title":"The CHiME-8 DASR Challenge for Generalizable and Array Agnostic Distant Automatic Speech Recognition and Diarization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.16447","snapshot_observed_at":"2026-08-07T12:30:34.936099Z","title":"The chime- 8 dasr challenge for generalizable and array agnostic distant automatic speech recognition and diarization,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:34.936099Z"},"links":{"cited_paper":"/paper/2407.16447","citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:eea3a3cbb3fa4959e0514502a7a56d66de4b2444ba7879dd9e99ed0c1fec96d4","observation_id":"06536cf9-065f-4b4b-baab-5fcd5ad6c7c2","resolution":{"observed_at":"2026-08-07T12:30:34.936099Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08887","last_updated":"2024-01-16T23:50:26Z","snapshot_observed_at":"2026-07-06T17:16:34.085999Z","submitted_at":"2024-01-16T23:50:26Z","title":"NOTSOFAR-1 Challenge: New Datasets, Baseline, and Tasks for Distant Meeting Transcription","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.08887","snapshot_observed_at":"2026-08-07T12:30:35.016211Z","title":"Notsofar- 1 challenge: New datasets, baseline, and tasks for distant meeting transcription,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:35.016211Z"},"links":{"cited_paper":"/paper/2401.08887","citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:81a77e697d3bc5c58d3bae5c22c841e9286c8cde526a1ceb2874b6e705e8161a","observation_id":"92dcfdc3-af3d-40a7-8c6f-c67dfb21b154","resolution":{"observed_at":"2026-08-07T12:30:35.016211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:35.115350Z","title":"The chime- 7 udase task: Unsupervised domain adaptation for conversational speech enhancement,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:35.115350Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:d6eff5f68e6954a8ac9c7c8c91d3e70255f1ff376df4ecf2b9aee52ae69dba63","observation_id":"0478ca83-f839-49e0-84d8-9a49f65c3d27","resolution":{"observed_at":"2026-08-07T12:30:35.115350Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:39.777923Z","title":"Mixture to mixture: Leveraging close-talk mixtures as weak-supervision for speech separation,","venue":null,"work_id":"689b5619-0073-4ad1-853f-efc1efc207e4","year":2024},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:35.193823Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:f9919b8e0fee91ffa95cf3f714f7eb21fe463cc1fc03767962ee545ec6d957be","observation_id":"5da5c48f-4e35-4350-a93b-3a0c06312abe","resolution":{"observed_at":"2026-08-07T12:30:39.825392Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:39.670095Z","title":"Unssor: unsupervised neural speech separation by leveraging over-determined training mix- tures,","venue":null,"work_id":"21b0b63e-9c48-4f57-a878-3304ad40f3d8","year":2024},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:35.236452Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:b5476ead39f37120c00fca5e5f133acd1004f2bb3d91829f4b8573ec192f8a91","observation_id":"3c1465f9-5da6-49c6-9ba1-6a59dfbc6635","resolution":{"observed_at":"2026-08-07T12:30:39.713295Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:35.308901Z","title":"Paraformer: Fast and accurate parallel transformer for non-autoregressive end-to- end speech recognition,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:35.308901Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:c050ad7a1b7b645aac797c51fe0960cbfe7bd746fe8d27048080e15f16c6379f","observation_id":"1d5b34c6-e3a1-4e74-bd53-114ab4a38863","resolution":{"observed_at":"2026-08-07T12:30:35.308901Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:39.491450Z","title":"Blind speech dereverberation with multi-channel linear prediction based on short time fourier transform representation,","venue":null,"work_id":"d8345467-6087-4027-8758-66c0724e88ac","year":2008},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:35.401751Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:352dff2761d4abb2acec4cd12279c77e879bdee1c5ca866dd0daa83d3c32ad84","observation_id":"01dc8897-efe5-4a11-86b7-106beb3ce534","resolution":{"observed_at":"2026-08-07T12:30:39.554595Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:39.372382Z","title":"Speech dereverberation based on variance-normalized de- layed linear prediction,","venue":null,"work_id":"beb60607-1d44-410f-8f47-dce6dce462f2","year":2010},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:35.477297Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:292e17229735718ad08250348cca1c9a788a30c730c5bc3eadeb2da6fe8f3fb5","observation_id":"89a96892-47a6-4105-b1af-e058696cf403","resolution":{"observed_at":"2026-08-07T12:30:39.414404Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:39.226777Z","title":"Complex angular central gaus- sian mixture model for directional statistics in mask-based mi- crophone array signal processing,","venue":null,"work_id":"a2e50000-1cab-431b-a199-a6a888051484","year":2016},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:35.590528Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:1a257707b4f61bf44640ad4aa6e171a27709a918c8c02d1d86189954e2e22e0c","observation_id":"cc7b0625-2350-4b4d-945f-7b7f8c16a51e","resolution":{"observed_at":"2026-08-07T12:30:39.289492Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:39.012809Z","title":"On optimal frequency- domain multichannel linear filtering for noise reduction,","venue":null,"work_id":"02b6b162-0c45-4303-aa23-3c1543a16ee3","year":2009},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:35.729681Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:f6dee8614f8a0e35d57d1c60252ef181ff592dd698ff9f9139bba4cd3dee8185","observation_id":"3ceb2b44-1b71-4a3f-9c0a-8e04a957d3d3","resolution":{"observed_at":"2026-08-07T12:30:39.158420Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:38.831987Z","title":"Improved mvdr beamforming using single-channel mask prediction networks","venue":null,"work_id":"b5581aff-0b90-4a9c-8400-d457dad81233","year":2016},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:35.806739Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:f0a82d0a1e5176554df898b67e4f056c18eaae450c2a6ebc997f796402e60fc7","observation_id":"1f81d9d0-9d7e-46d6-b37d-7af21374bbda","resolution":{"observed_at":"2026-08-07T12:30:38.925746Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:38.672000Z","title":"Spatialnet: Extensively learning spatial in- formation for multichannel joint speech separation, denoising and dereverberation,","venue":null,"work_id":"d309bd48-8986-447b-9d76-f254dbd37231","year":2024},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:35.928098Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:9b6042d92d18c4c5c9b73566a30762f538cf959e664dd98b8be56509b12a4dce","observation_id":"528a73af-2815-4c12-97be-4f5abee422f4","resolution":{"observed_at":"2026-08-07T12:30:38.760348Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:36.031594Z","title":"The xmuspeech system for audio-visual target speaker extraction in misp 2023 challenge,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:36.031594Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:15a36f6166d32e6eede5f68778d6bde09373b876216bd802d676aa6a9e0b54af","observation_id":"d55be7e7-cb0f-42d6-9ecd-20df419773d7","resolution":{"observed_at":"2026-08-07T12:30:36.031594Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:38.510533Z","title":"Phase- sensitive and recognition-boosted speech separation using deep recurrent neural networks,","venue":null,"work_id":"8426138a-976c-4716-b4fb-6e76758515ca","year":2015},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:36.139783Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:f43cc993f11bb7af3c630bc99083dc494b0f2b5a5799ffec5a2b7e339e9d7dc3","observation_id":"28906618-6b33-4427-925a-765623af4fbf","resolution":{"observed_at":"2026-08-07T12:30:38.556755Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:36.238313Z","title":"The fifth’chime’speech separation and recognition challenge: Dataset, task and baselines,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:36.238313Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:ad2a8d43650998e358c1ca733fec255c4a8ce74c0c654268acaf1f208a4d247b","observation_id":"8db3ae9a-a524-4205-9ca4-1595c0cb0bf2","resolution":{"observed_at":"2026-08-07T12:30:36.238313Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:36.317225Z","title":"M2met: The icassp 2022 multi- channel multi-party meeting transcription challenge,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:36.317225Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:3b458405d8431fb584947a352750367ecb1f82db4293fa00c42c1ba78da11fe7","observation_id":"6f1bb983-7646-4524-9f8a-d811de2fbd73","resolution":{"observed_at":"2026-08-07T12:30:36.317225Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:36.460559Z","title":"Dnsmos p. 835: A non-intrusive perceptual objective speech quality metric to eval- uate noise suppressors,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:36.460559Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:b601a80b7c1df0f82da965eafe5081ffc955d2675bd2d20d664b766589547a4e","observation_id":"bb25eaf6-d172-45a3-ac90-8b72ad48775d","resolution":{"observed_at":"2026-08-07T12:30:36.460559Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:36.557559Z","title":"Robust speech recognition via large-scale weak supervision,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:36.557559Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:cf458276328688696453d69a7ee3cc74cdd2fea9d29d0841f4bbe05bddf7818f","observation_id":"227ce695-6ba4-4329-9fff-2eb6789b77c8","resolution":{"observed_at":"2026-08-07T12:30:36.557559Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:38.322492Z","title":"Robust localization in reverberant rooms,","venue":null,"work_id":"f3b9d6bd-67ee-4cc5-9d6d-a4aace887a6f","year":2001},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:36.662380Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:532c953d88e6ba5ca26d50cde864e0b3d6d268d26219f2e84ca459ba0a3eb9ec","observation_id":"9af971e3-f9ee-4b52-bbd2-63dc77428dbe","resolution":{"observed_at":"2026-08-07T12:30:38.372043Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:36.736817Z","title":"Convolutive predic- tion for monaural speech dereverberation and noisy-reverberant speaker separation,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:36.736817Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:f95df63808174f5100f2d104aaefe047bff8fddc7362a7de993fc32150dcce30","observation_id":"cd07b7f5-23d1-4535-bd07-7f012c7c1029","resolution":{"observed_at":"2026-08-07T12:30:36.736817Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:36.811669Z","title":"Supervised speech separation based on deep learning: An overview,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:36.811669Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:7842990bb91acb03285fa26d7f6b079427aba57daf8e0ee32c079ca35e591b7b","observation_id":"9e0e357c-1846-48bd-9b2c-5366be07f668","resolution":{"observed_at":"2026-08-07T12:30:36.811669Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:38.214320Z","title":"A simultane- ous denoising and dereverberation framework with target decou- pling,","venue":null,"work_id":"5d40fd83-19f1-4931-b788-162c74bab387","year":2021},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:36.903270Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:c45a84641acdb2a20b7146652c0b96ccab9ae49120ad1aea1cd14c668636a3fd","observation_id":"7c2d1981-9fd9-46ca-9842-56f8dec17ffe","resolution":{"observed_at":"2026-08-07T12:30:38.247147Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:38.082181Z","title":"Cif: Continuous integrate-and-fire for end- to-end speech recognition,","venue":null,"work_id":"fc1ae277-24ea-4a0c-bad3-72b6051bcd60","year":2020},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:36.974767Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:b5f9a7cdf1d796c5aa6a66cadcf2153dbc6761c777316cb1a2c96778e4285d28","observation_id":"b021024e-6bd0-41af-95c7-4175d0ef24a4","resolution":{"observed_at":"2026-08-07T12:30:38.151024Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2008.07905","last_updated":"2021-05-13T14:41:40Z","snapshot_observed_at":"2026-07-06T09:48:06.247669Z","submitted_at":"2020-08-18T13:04:03Z","title":"Glancing Transformer for Non-Autoregressive Neural Machine Translation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2008.07905","snapshot_observed_at":"2026-08-07T12:30:37.126709Z","title":"Glancing transformer for non-autoregressive neural machine translation,","venue":null,"work_id":null,"year":2008},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:37.126709Z"},"links":{"cited_paper":"/paper/2008.07905","citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:cfde77f35d82e7ce483942ae043a385dd7b61eb861df8ebb25d9b73bb7f99972","observation_id":"48393f72-73bc-4155-a2d0-9112d4fdac34","resolution":{"observed_at":"2026-08-07T12:30:37.126709Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:37.941275Z","title":"Aishell-4: An open source dataset for speech enhancement, separation, recognition and speaker diarization in conference scenario,","venue":null,"work_id":"fa6f0a38-c1b0-443f-9846-fd9cb84726c2","year":2021},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:37.200851Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:0d5ee8e6f3f90b3cd9093124a540fd82e32f139449b13e042c0c820575306a32","observation_id":"1ef82e78-35fe-4881-91a0-74aa41d679dd","resolution":{"observed_at":"2026-08-07T12:30:38.011386Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:30:37.726943Z","title":"Wenetspeech: A 10000+ hours multi-domain mandarin corpus for speech recognition,","venue":null,"work_id":"5fbd531d-29f6-46e8-a6d5-c8a012b4d0e8","year":null},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:37.271432Z"},"links":{"citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:f8beccfeccbba0e1b6d316c38d119c4c5debae799c260df674d2a8f2337fb70d","observation_id":"36285c20-a811-4f59-bba0-2c2a572fe025","resolution":{"observed_at":"2026-08-07T12:30:37.815144Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2110.03370","last_updated":"2022-02-23T06:42:31Z","snapshot_observed_at":"2026-08-09T03:42:41.612429Z","submitted_at":"2021-10-07T12:05:29Z","title":"WenetSpeech: A 10000+ Hours Multi-domain Mandarin Corpus for Speech Recognition","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.03370","snapshot_observed_at":"2026-08-07T12:30:37.358858Z","title":"Available: https://arxiv.org/abs/2110.03370","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge","version":2},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-07T12:30:37.358858Z"},"links":{"cited_paper":"/paper/2110.03370","citing_paper":"/paper/2505.24446"},"observation_digest":"sha256:51f2332836cf67cb192a5b2d9eaaa3306a5508ec0198b3d8d71848d8baf1b52b","observation_id":"20ba5e4b-8a6c-45a4-b627-74fb13feaee1","resolution":{"observed_at":"2026-08-07T12:30:37.358858Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2505.24446","last_updated":"2025-06-23T07:54:08Z","latest_version":2,"primary_category":"cs.SD","snapshot_observed_at":"2026-08-07T12:19:35.007574Z","submitted_at":"2025-05-30T10:33:54Z","title":"Pseudo Labels-based Neural Speech Enhancement for the AVSR Task in the MISP-Meeting Challenge"},"reference_resolution":{"displayed":41,"state_counts":{"malformed_identifier":1,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":22,"verified_exact":0,"verified_fuzzy":17},"total_outbound_references":41},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 41 of 41 outbound references and 1 inbound Pith citation observation for arXiv:2505.24446."}