{"as_of":"2026-08-08T23:55:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:11927404d87150228f5cdca732fa003fe9814b0dcdef0a5e4175d7f3e797b3f2","coverage":[{"denominator":33,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":33,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T12:55:24.534982Z","state":"measured"},{"denominator":34,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":34,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T12:55:22.270498Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-07T12:55:26.461821Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"cited_work":{"arxiv_id":"2505.23207","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.23207","snapshot_observed_at":"2026-08-07T12:55:26.461821Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","venue":"cs.SD","work_id":"32ae2e0d-6e7e-4430-81ee-f255fb40fcaa","year":2025},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:22.270498Z"},"links":{"cited_paper":"/paper/2505.23207","citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:9e220f9a5d7779ff4bb72fcdd08c72047d026ddf179cd1b5ae88827abc2a3b0d","observation_id":"0c939e92-2816-42d6-a6c3-f0aef865badc","resolution":{"observed_at":"2026-08-07T12:55:26.468070Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2505.23207/citation-record","integrity":"/paper/2505.23207/integrity","json":"/paper/2505.23207/citation-record.json","paper":"/paper/2505.23207"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:27.836144Z","title":"Overlapping regions present unique challenges, as they involve both overlap of speech segments and speaker identities","venue":null,"work_id":"8667a66b-5bfb-46c8-b404-02572f8918a8","year":null},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:22.204152Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:45023bb80aee263edde8c81b8c57d915af30178f1f3ff7f9fc48a6d2a3850441","observation_id":"7d5a2b0e-7db5-476e-8385-2cfddbf7dc21","resolution":{"observed_at":"2026-08-07T12:55:27.881421Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:27.746152Z","title":"We be- gin with an overview of the progressive OSD architecture, fol- lowed by a detailed explanation of the progressive OSD training strategy","venue":null,"work_id":"d563984f-8de8-4c23-9d60-a235a705fa75","year":null},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:22.362092Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:437970d3508f4abb22b023a04c8190b889eae2da9e42736f12e4e932c9a48a8b","observation_id":"01f4d765-3c10-453b-a958-7fd9cc6d5a7e","resolution":{"observed_at":"2026-08-07T12:55:27.776832Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:27.533714Z","title":"Experiments setup In this section, we describe the experimental setup used for our paper, including the datasets and training configurations","venue":null,"work_id":"4edc2d85-1bd4-4b10-acef-f1d4b9c19530","year":null},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:22.540302Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:8f5afe820b363a9cea76d8501636d7719d144c8f69a02ce1f7a173b369b38fbd","observation_id":"d0a1eabd-7eed-407a-8820-ee4890fd8a98","resolution":{"observed_at":"2026-08-07T12:55:27.601092Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:27.428285Z","title":"Ex- perimental results consistently demonstrate the advantages of this approach across different configurations","venue":null,"work_id":"e502c62e-0bb1-45aa-968c-16ae6bfc6e89","year":2025},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:22.623923Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:4b2da2d8f84735f34ec7e70211d1012b00bd3d888a876ee751802cd7f05a634d","observation_id":"2cb4a584-f50c-45a1-a8d2-4053d88c68f4","resolution":{"observed_at":"2026-08-07T12:55:27.480431Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:27.018354Z","title":"Detection of overlapping speech for the purposes of speaker diarization,","venue":null,"work_id":"8837505a-53c7-4d1a-86a6-0803da4d9b2a","year":2019},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:23.056677Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:1b3d72c4abfa65103792c703199a232ee6de6c881267cfe980923dba1b55148f","observation_id":"54bfc9cf-c1a6-41d2-91e2-7490932c5b79","resolution":{"observed_at":"2026-08-07T12:55:27.048409Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1910.11646","last_updated":"2019-10-25T12:22:14Z","snapshot_observed_at":"2026-07-06T08:32:15.011487Z","submitted_at":"2019-10-25T12:22:14Z","title":"Overlap-aware diarization: resegmentation using neural end-to-end overlapped speech detection","version":1},"cited_work":{"arxiv_id":"1910.11646","doi":null,"metadata_source":"pith","pith_arxiv_id":"1910.11646","snapshot_observed_at":"2026-08-07T12:55:26.341121Z","title":"Overlap-aware diarization: resegmentation using neural end-to-end overlapped speech detection","venue":"eess.AS","work_id":"bd900b56-373b-42f9-b853-318cc187d5d8","year":2019},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:23.123922Z"},"links":{"cited_paper":"/paper/1910.11646","citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:31fd7d29b37220586832c84d4468e93180277352dae8dcf5e51bfe4346076e02","observation_id":"19b4013f-3da8-48cb-99b7-59c0466f72c6","resolution":{"observed_at":"2026-08-07T12:55:26.413671Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:27.265215Z","title":"Diarization is hard: Some experiences and lessons learned for the jhu team in the inaugural dihard challenge,","venue":null,"work_id":"65f6ae4e-edee-4cac-9c4a-653a7139d58d","year":2018},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:22.723328Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:9650e7fe880c00b11dcc1b219446c55fe745a8e61f480bc1a9f09857f358edf6","observation_id":"7801a417-4517-4951-915e-a1b063b288fd","resolution":{"observed_at":"2026-08-07T12:55:27.335488Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:27.122584Z","title":"Overlapped speech detection for improved speaker diarization in multiparty meetings,","venue":null,"work_id":"f8a1f43b-9858-48a9-b110-46c30974ebf7","year":2008},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:22.809830Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:41eaf0f15ccd76287927844c43592391f868013f0cfee04e5c0d292023dbbb02","observation_id":"3c60974b-963c-4127-b547-3db2af44b773","resolution":{"observed_at":"2026-08-07T12:55:27.163884Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:27.107489Z","title":"Impact of overlapping speech detection on speaker diarization for broadcast news and debates,","venue":null,"work_id":"2779731a-3d5f-4c25-b917-b6ffb770e37e","year":2013},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:22.911625Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:f39650c7873e77ff4bfcc8e6625396d13a3775e84fd67a962ef521cd1565a0b1","observation_id":"dedf134a-0f63-4b36-ae91-0818acd1973b","resolution":{"observed_at":"2026-08-07T12:55:27.117536Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:27.057315Z","title":"But system for dihard speech diariza- tion challenge 2018,","venue":null,"work_id":"f931d9a7-be4d-42c1-b5ff-9a2264fe6ba8","year":2018},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:22.977314Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:f34148cd26424d06da6b6ef53a6539988fb9fd405a0624a498b3a4c92217651b","observation_id":"7c58cc05-b030-4585-826a-87c52ec2d9c8","resolution":{"observed_at":"2026-08-07T12:55:27.089770Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:26.832254Z","title":"Over- lapped speech detection in broadcast streams using x-vectors","venue":null,"work_id":"b0acbddb-f5af-42cf-a62c-cf5e61f75036","year":2022},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:23.504834Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:fe52004dda6b9266b8a633c1c62243385d36eaeea8c48b93fdabd9147a56b1ad","observation_id":"68d7d2b3-2805-4f86-ab3e-857b402f75db","resolution":{"observed_at":"2026-08-07T12:55:26.882336Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.00819","last_updated":"2024-09-01T19:23:08Z","snapshot_observed_at":"2026-07-06T19:08:57.415038Z","submitted_at":"2024-09-01T19:23:08Z","title":"LibriheavyMix: A 20,000-Hour Dataset for Single-Channel Reverberant Multi-Talker Speech Separation, ASR and Speaker Diarization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.00819","snapshot_observed_at":"2026-08-07T12:55:23.552544Z","title":"Libriheavymix: a 20,000-hour dataset for single-channel reverberant multi-talker speech separation, asr and speaker diarization,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:23.552544Z"},"links":{"cited_paper":"/paper/2409.00819","citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:bd9899540c716baf44ff30986de13959abdf6799ba00200fd5a0d5faadcd9372","observation_id":"25eddd46-20c7-413e-abc3-cbe12ac82788","resolution":{"observed_at":"2026-08-07T12:55:23.552544Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2209.12002","last_updated":"2022-09-24T13:04:52Z","snapshot_observed_at":"2026-08-06T16:59:53.365265Z","submitted_at":"2022-09-24T13:04:52Z","title":"Spatial-aware Speaker Diarization for Multi-channel Multi-party Meeting","version":1},"cited_work":{"arxiv_id":"2209.12002","doi":null,"metadata_source":"pith","pith_arxiv_id":"2209.12002","snapshot_observed_at":"2026-08-07T12:55:26.135381Z","title":"Spatial-aware Speaker Diarization for Multi-channel Multi-party Meeting","venue":"eess.AS","work_id":"986f5ced-653b-436a-9be4-983437031208","year":2022},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:23.180642Z"},"links":{"cited_paper":"/paper/2209.12002","citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:6db82f014a45db8daecbdd13afae46a1ca95ed2fadca01478b266bf52cc9a8a1","observation_id":"aa3cc08e-f03e-4394-bb18-11330dd16de1","resolution":{"observed_at":"2026-08-07T12:55:26.254347Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"cited_work":{"arxiv_id":"2505.23207","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.23207","snapshot_observed_at":"2026-08-07T12:55:26.461821Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","venue":"cs.SD","work_id":"32ae2e0d-6e7e-4430-81ee-f255fb40fcaa","year":2025},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:22.270498Z"},"links":{"cited_paper":"/paper/2505.23207","citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:9e220f9a5d7779ff4bb72fcdd08c72047d026ddf179cd1b5ae88827abc2a3b0d","observation_id":"0c939e92-2816-42d6-a6c3-f0aef865badc","resolution":{"observed_at":"2026-08-07T12:55:26.468070Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2104.04045","last_updated":"2021-06-10T13:51:47Z","snapshot_observed_at":"2026-08-08T15:02:17.377370Z","submitted_at":"2021-04-08T20:38:17Z","title":"End-to-end speaker segmentation for overlap-aware resegmentation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2104.04045","snapshot_observed_at":"2026-08-07T12:55:23.255131Z","title":"End-to-end speaker segmentation for overlap-aware resegmentation,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:23.255131Z"},"links":{"cited_paper":"/paper/2104.04045","citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:94c9fdd44a412726094c67627d796269f9d665b18a25db3bc8fe0442c12c3f3c","observation_id":"57d4a822-6ce5-4784-8828-42d818cbb285","resolution":{"observed_at":"2026-08-07T12:55:23.255131Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.05987","last_updated":"2023-09-07T07:56:10Z","snapshot_observed_at":"2026-07-06T16:05:11.328794Z","submitted_at":"2023-08-11T07:50:41Z","title":"Large-Scale Learning on Overlapped Speech Detection: New Benchmark and New General System","version":3},"cited_work":{"arxiv_id":"2308.05987","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.05987","snapshot_observed_at":"2026-08-07T12:55:25.970505Z","title":"Large-Scale Learning on Overlapped Speech Detection: New Benchmark and New General System","venue":"cs.SD","work_id":"fe053dc3-0730-4f77-92b7-74fead59aee9","year":2023},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:23.323154Z"},"links":{"cited_paper":"/paper/2308.05987","citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:ae6ceaba5ab2ec49ef1413a32e8ddca0000c537385178e659f783d180934da32","observation_id":"9f58ddf8-6df3-40cb-b617-85568ac791f2","resolution":{"observed_at":"2026-08-07T12:55:26.006554Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:26.969942Z","title":"Automatic detection of multi- speaker fragments with high time resolution,","venue":null,"work_id":"536b4e30-41cd-42c0-bfdd-4697e730cdf1","year":2018},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:23.420013Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:9109d86d6458836db9b5c6e9030a2e289769a4c739f0a2c96769dfa1b3395f5a","observation_id":"bf33a67d-5e8c-452d-a8a6-5d0cca8ca36c","resolution":{"observed_at":"2026-08-07T12:55:26.995175Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2109.04049","last_updated":"2021-09-09T06:10:48Z","snapshot_observed_at":"2026-07-06T11:45:52.747643Z","submitted_at":"2021-09-09T06:10:48Z","title":"BeamTransformer: Microphone Array-based Overlapping Speech Detection","version":1},"cited_work":{"arxiv_id":"2109.04049","doi":null,"metadata_source":"pith","pith_arxiv_id":"2109.04049","snapshot_observed_at":"2026-08-07T12:55:25.736785Z","title":"BeamTransformer: Microphone Array-based Overlapping Speech Detection","venue":"cs.SD","work_id":"69d1d80c-acc7-457d-ad88-c31b04e08e5c","year":2021},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:23.981026Z"},"links":{"cited_paper":"/paper/2109.04049","citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:5842dc3b1a6f5716f071823c49997ccda4787ab3a00e68711222be914889c6de","observation_id":"2c72fb69-e9f7-4391-8b94-51c92042e0e1","resolution":{"observed_at":"2026-08-07T12:55:25.924681Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:23.633794Z","title":"Multitask detection of speaker changes, overlapping speech and voice activity using wav2vec 2.0,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:23.633794Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:a8d6cce548cbfa161d80a10103f3e954c9112203029856d02ef7fba632902d9c","observation_id":"8e3967e4-3b72-4537-bb5e-e4e3e5bb4fad","resolution":{"observed_at":"2026-08-07T12:55:23.633794Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:23.697828Z","title":"M2met: The icassp 2022 multi- channel multi-party meeting transcription challenge,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:23.697828Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:7108c87ee365737bef22fd0feb2f2485393b1872118d9911f6d04a12852f7200","observation_id":"ffd7667d-2e9d-4c9a-a948-d991e34033ce","resolution":{"observed_at":"2026-08-07T12:55:23.697828Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:26.727920Z","title":"Recognition and under- standing of meetings the ami and amida projects,","venue":null,"work_id":"e7daa072-5847-45cd-ac65-f72d1e21016d","year":2007},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:23.764082Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:4e85fda646ac5f6f90ae36022ea05e33a50af008eefbe9be7fa4697d7466e1b2","observation_id":"700fc423-6ce2-49c8-98f2-b9d2ece4006b","resolution":{"observed_at":"2026-08-07T12:55:26.773092Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:23.827907Z","title":"A survey on multi-task learning,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:23.827907Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:daa33e68fb27a200cbb643bcd20f18be594c6c483125364f1cd939ff6f5dbe28","observation_id":"af77ebdc-b08f-4a93-bb19-332d31537a5a","resolution":{"observed_at":"2026-08-07T12:55:23.827907Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2104.02878","last_updated":"2021-04-07T03:01:34Z","snapshot_observed_at":"2026-07-06T10:57:08.129258Z","submitted_at":"2021-04-07T03:01:34Z","title":"Three-class Overlapped Speech Detection using a Convolutional Recurrent Neural Network","version":1},"cited_work":{"arxiv_id":"2104.02878","doi":null,"metadata_source":"pith","pith_arxiv_id":"2104.02878","snapshot_observed_at":"2026-08-07T12:55:25.955781Z","title":"Three-class Overlapped Speech Detection using a Convolutional Recurrent Neural Network","venue":"eess.AS","work_id":"face87b5-27fc-4a9e-bf48-92ff1dfed051","year":2021},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:23.910130Z"},"links":{"cited_paper":"/paper/2104.02878","citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:0bb3700092141636f0e5356a195842d26d40cad52630cffb8ea94be31197bb88","observation_id":"2bfca183-8cf7-4478-bc99-e79e6d22bb95","resolution":{"observed_at":"2026-08-07T12:55:25.958145Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:26.472265Z","title":"Target- speaker voice activity detection: A novel approach for multi- speaker diarization in a dinner party scenario,","venue":null,"work_id":"81e8ebe4-46ce-49b2-8168-30d733ca5d68","year":null},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:24.369768Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:87e8fc485a59d20f558b98e25579a80d18fbe47b915ce34efe3a99143eb7988b","observation_id":"6fbb3121-2452-41fd-803a-6ab5455fd730","resolution":{"observed_at":"2026-08-07T12:55:26.474965Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:26.605284Z","title":"A survey on text-dependent and text-independent speaker verification,","venue":null,"work_id":"2dddd383-94bf-4cd2-af03-6867b538eb7d","year":2022},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:24.058217Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:3868fdbcc4069cdfddde7ce3c4ef9bfdf8450cd68fb3b9722662747f95e6c805","observation_id":"dc20b860-6cce-48e7-b612-ff57b515e4b2","resolution":{"observed_at":"2026-08-07T12:55:26.651952Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:27.673146Z","title":"In this paper, we adapt WavLM for OSD tasks","venue":null,"work_id":"98c7c1aa-3a04-428b-b8c8-a3180af1b53e","year":null},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:22.454267Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:00af942615e4077ace2a94519aafc02e1cce86a43df54ff2ed99a7550af29749","observation_id":"7de0ce69-62c6-4636-b13a-852c09fa9e51","resolution":{"observed_at":"2026-08-07T12:55:27.688496Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2104.01466","last_updated":"2021-04-03T19:37:51Z","snapshot_observed_at":"2026-08-03T23:33:42.614126Z","submitted_at":"2021-04-03T19:37:51Z","title":"ECAPA-TDNN Embeddings for Speaker Diarization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2104.01466","snapshot_observed_at":"2026-08-07T12:55:24.123323Z","title":"Ecapa-tdnn embeddings for speaker di- arization,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:24.123323Z"},"links":{"cited_paper":"/paper/2104.01466","citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:c2875d6156ec62bb1fec28a18c3fb1861bcdb2f393f7eb7775876694c94f1235","observation_id":"6794acf1-22c7-4ed7-916c-4b27545f55c5","resolution":{"observed_at":"2026-08-07T12:55:24.123323Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.00332","last_updated":"2023-06-16T09:07:11Z","snapshot_observed_at":"2026-08-05T18:44:00.793642Z","submitted_at":"2023-03-01T08:50:31Z","title":"CAM++: A Fast and Efficient Network for Speaker Verification Using Context-Aware Masking","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.00332","snapshot_observed_at":"2026-08-07T12:55:24.199871Z","title":"Cam++: A fast and efficient network for speaker verification using context- aware masking,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:24.199871Z"},"links":{"cited_paper":"/paper/2303.00332","citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:a07f2f4e1c5753e33b985ebaaf5549d1c47f7e20fc6e56bbca6e3e7f1126955e","observation_id":"e3a807a5-cb05-448e-beb9-10498a10e830","resolution":{"observed_at":"2026-08-07T12:55:24.199871Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:24.261854Z","title":"Wavlm: Large-scale self- supervised pre-training for full stack speech processing,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:24.261854Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:f39a49d1d4fb188e2e19a9549ff489e8e9c4c4c1686a91e25e882051ab0071f9","observation_id":"800c25ff-5116-41db-ba03-bf0c0d9e30a7","resolution":{"observed_at":"2026-08-07T12:55:24.261854Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:55:26.489351Z","title":"Frame-wise and overlap-robust speaker em- beddings for meeting diarization,","venue":null,"work_id":"8878baf7-1a2a-4691-b005-19be132a32d2","year":2023},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:24.319519Z"},"links":{"citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:650f78a44310498ab3ee97886e536e5ad93d67e72c989b20ff42a1a240e24d82","observation_id":"0017e364-824d-4ecd-ab6f-53770d6b2b09","resolution":{"observed_at":"2026-08-07T12:55:26.521430Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.13012","last_updated":"2023-07-24T14:29:21Z","snapshot_observed_at":"2026-07-06T15:57:57.167169Z","submitted_at":"2023-07-24T14:29:21Z","title":"Joint speech and overlap detection: a benchmark over multiple audio setup and speech domains","version":1},"cited_work":{"arxiv_id":"2307.13012","doi":null,"metadata_source":"pith","pith_arxiv_id":"2307.13012","snapshot_observed_at":"2026-08-07T12:55:25.348748Z","title":"Joint speech and overlap detection: a benchmark over multiple audio setup and speech domains","venue":"cs.SD","work_id":"50d6f856-dc7e-44b0-b1fa-79053b278201","year":2023},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:24.417081Z"},"links":{"cited_paper":"/paper/2307.13012","citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:fce47fcde94b04e3d278db53e2beb3cfe320b2c74080ff35f3d401bf98b9e315","observation_id":"064fbc58-6a89-4bdc-8922-2bf9f78607b6","resolution":{"observed_at":"2026-08-07T12:55:25.549707Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.08552","last_updated":"2024-09-13T06:11:11Z","snapshot_observed_at":"2026-08-06T17:39:08.220185Z","submitted_at":"2024-09-13T06:11:11Z","title":"Unified Audio Event Detection","version":1},"cited_work":{"arxiv_id":"2409.08552","doi":null,"metadata_source":"pith","pith_arxiv_id":"2409.08552","snapshot_observed_at":"2026-08-07T12:55:24.968697Z","title":"Unified Audio Event Detection","venue":"eess.AS","work_id":"096ae29f-0019-4f12-ab80-b18e9fb41bdd","year":2024},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:24.457100Z"},"links":{"cited_paper":"/paper/2409.08552","citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:38f2876df7e722fb5bc65897ece3e2972124b956c10caaf6989cf76963a6fe6f","observation_id":"012e5df1-d63a-49c9-842e-ae7e0e7ec4e8","resolution":{"observed_at":"2026-08-07T12:55:25.125436Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.01960","last_updated":"2024-03-04T11:57:32Z","snapshot_observed_at":"2026-08-06T09:21:17.823914Z","submitted_at":"2024-03-04T11:57:32Z","title":"A robust audio deepfake detection system via multi-view feature","version":1},"cited_work":{"arxiv_id":"2403.01960","doi":null,"metadata_source":"pith","pith_arxiv_id":"2403.01960","snapshot_observed_at":"2026-08-07T12:55:24.615890Z","title":"A robust audio deepfake detection system via multi-view feature","venue":"cs.SD","work_id":"87ffefc2-5e43-446b-8bc7-7fc261324489","year":2024},"citing_paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T12:55:24.534982Z"},"links":{"cited_paper":"/paper/2403.01960","citing_paper":"/paper/2505.23207"},"observation_digest":"sha256:dd9e803549f917b8b341ad48bb8e27a533cd6d72c522dcb8a481f244c6d343b8","observation_id":"7dc939a6-a656-4182-aa5a-07550c43ef72","resolution":{"observed_at":"2026-08-07T12:55:24.757257Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2505.23207","last_updated":"2025-05-29T07:47:48Z","latest_version":1,"primary_category":"cs.SD","snapshot_observed_at":"2026-08-08T12:28:17.887912Z","submitted_at":"2025-05-29T07:47:48Z","title":"Towards Robust Overlapping Speech Detection: A Speaker-Aware Progressive Approach Using WavLM"},"reference_resolution":{"displayed":33,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":8,"verified_exact":9,"verified_fuzzy":16},"total_outbound_references":33},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 33 of 33 outbound references and 1 inbound Pith citation observation for arXiv:2505.23207."}