{"as_of":"2026-08-09T22:27:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:33769d5c6a84ff3eaffd885819268dec6abf3e27f19ebaa6dcc049e6e707f3fb","coverage":[{"denominator":21,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":21,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T22:36:53.842471Z","state":"measured"},{"denominator":21,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":21,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.21174/citation-record","integrity":"/paper/2506.21174/integrity","json":"/paper/2506.21174/citation-record.json","paper":"/paper/2506.21174"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:57.577427Z","title":"This complex task requires a system to identify active sound classes (audio tagging) and to isolate their corresponding anechoic source signals accurately","venue":null,"work_id":"f4a66d28-9b9a-4e54-a93f-f16519db3633","year":2025},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:51.984029Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:1e2d4645489784f50289c5584a55157f8db3ba2990fce36de96eaac5b2bafa8d","observation_id":"594845b1-4df0-45ac-9b2f-5fbbad5e9c55","resolution":{"observed_at":"2026-08-06T22:36:57.666390Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:57.354890Z","title":"This section outlines the com- position of the official dataset and the specific data curation and augmentation steps to improve the performance of the model","venue":null,"work_id":"42782da9-ea99-4d84-8e5c-8dbeb4eaa5ea","year":2025},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:52.076241Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:04c942fe005c386c4cb8801388870516761bdd73027e0dc10da008f2febd4218","observation_id":"6d07eef9-ab98-4f10-b40b-044f7210a1ab","resolution":{"observed_at":"2026-08-06T22:36:57.489728Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:56.932540Z","title":null,"venue":null,"work_id":"ee408885-da62-46c2-a403-4448136ec8ac","year":null},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:52.269462Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:5269543765ab8447c4c39e34a695e1c39493c9dc1d99a0c9b5549bb64ac7f835","observation_id":"9aef203b-bec4-4fff-b5d5-01e5669690f0","resolution":{"observed_at":"2026-08-06T22:36:57.034367Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:56.531372Z","title":"Model training The audio-tagging and source-separation models were trained in- dependently","venue":null,"work_id":"bfc06118-4042-45e1-ac7e-28530c2dcece","year":2025},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:52.564116Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:fce63a34fc7669ae86533431e7d180b1e3f37f23b1cfa392af580442acdf96d0","observation_id":"df393d6e-d508-49e4-83be-ef658276eaea","resolution":{"observed_at":"2026-08-06T22:36:56.632808Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:56.349495Z","title":"The proposed strategy combined additional audio feature input (spectral roll-off and chroma), a dataset refinement process, and an agent-based er- ror correction system","venue":null,"work_id":"8e8275ff-c192-4021-95be-aca57c1b461f","year":2025},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:52.670206Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:ba4cf57c23f4b4893b0478fac39b0d835a328a3d94fbc350f4828d4ff9b209b6","observation_id":"7d65e086-be41-4beb-b043-31c4a2c1be09","resolution":{"observed_at":"2026-08-06T22:36:56.421479Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:55.393887Z","title":"SpatialScaper: A library to simulate and augment soundscapes for sound event localization and detection in re- alistic rooms,","venue":null,"work_id":"9a1c1e24-7823-4ae2-b3ef-0b5b4d640a7d","year":2022},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:53.260325Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:c5a82538dab85b35bc5d4bd9d095c8fd3c790899454a4f59af6daadbe508267e","observation_id":"b8cb2d5b-f9bb-406d-8a2e-200ac2f866fd","resolution":{"observed_at":"2026-08-06T22:36:55.479863Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:55.170842Z","title":"FSD50K: An open dataset of human -labeled sound events,","venue":null,"work_id":"fa542f6c-5d88-41b9-9305-f10e0b822fa4","year":2022},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:53.336148Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:e744e42b3236ce7b66b132d8bf1ae1c0279b743985a25376527190d4190c7de2","observation_id":"c64136c5-da74-48c4-865f-95c227379066","resolution":{"observed_at":"2026-08-06T22:36:55.270723Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.10676","last_updated":"2025-06-12T13:07:54Z","snapshot_observed_at":"2026-08-07T04:18:06.150385Z","submitted_at":"2025-06-12T13:07:54Z","title":"Description and Discussion on DCASE 2025 Challenge Task 4: Spatial Semantic Segmentation of Sound Scenes","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.10676","snapshot_observed_at":"2026-08-06T22:36:52.750273Z","title":"Description and Discussion on DCASE 2025 Challenge Task 4: Spatial Semantic Segmentation of Sound Scenes,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:52.750273Z"},"links":{"cited_paper":"/paper/2506.10676","citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:8ea5018675cd57dba55c4e5c6e8aab9b11e3ab6741b0036e5182d35df86e28bc","observation_id":"7fe2ea60-7d33-4e04-b7c3-bff7316d9dae","resolution":{"observed_at":"2026-08-06T22:36:52.750273Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.22088","last_updated":"2025-06-09T08:23:48Z","snapshot_observed_at":"2026-08-07T16:31:17.639140Z","submitted_at":"2025-03-28T02:08:58Z","title":"Baseline Systems and Evaluation Metrics for Spatial Semantic Segmentation of Sound Scenes","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.22088","snapshot_observed_at":"2026-08-06T22:36:52.846269Z","title":"Baseline systems and evaluation metrics for spatial semantic segmentation of sound scenes,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:52.846269Z"},"links":{"cited_paper":"/paper/2503.22088","citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:5cafb1652453ed0d0e4875880a6f649a935fa3a929ea1312da7e97228e386ed7","observation_id":"7020296b-b2a9-42cf-b4d9-ecbe937a319f","resolution":{"observed_at":"2026-08-06T22:36:52.846269Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:56.750488Z","title":null,"venue":null,"work_id":"7fc5bf11-4f15-4e3b-baeb-f9ca59c508b2","year":2025},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:52.410628Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:219a6622454bf994d16aebce585fc436096520bc92de8570690d5d64c5870959","observation_id":"a98befa5-111e-4ae3-83a4-b9bf9eb68ec6","resolution":{"observed_at":"2026-08-06T22:36:56.829497Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:57.143444Z","title":"All audio mixes of 10 seconds long each were generated at a sampling rate of 32 kHz","venue":null,"work_id":"f59629c8-43be-4f50-b651-e7a0d3fa9153","year":null},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:52.165942Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:eb67d8e12b21abd4c3dcc55660055eb915edf81fd59e9bafdb2d5eceae4f0c57","observation_id":"3cfc95c4-4be1-4441-81b5-2830196e37f8","resolution":{"observed_at":"2026-08-06T22:36:57.244612Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:56.128720Z","title":"Classification in the presence of label noise: A survey,","venue":null,"work_id":"16102391-1b5e-4359-af7b-c961feca44a4","year":2014},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:52.958954Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:3610b0013da12553e6534838c14bafd501d8850eae40bea0902de68dd2918c77","observation_id":"ea2b1f06-cf0c-40a7-9cb1-7769dbc84a58","resolution":{"observed_at":"2026-08-06T22:36:56.231767Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:55.845587Z","title":"Acoustic classification and seg- mentation using modified spectral roll-off and variance-based features","venue":null,"work_id":"df5c8129-abd4-4a02-b190-f7baae952a74","year":2013},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:53.071778Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:6a11063a65fefd35fa2c5e7b3bfae3aaf8d638348e07e960c50512e995dcdc62","observation_id":"8f659957-e00e-4b53-9744-bb82fbce96c8","resolution":{"observed_at":"2026-08-06T22:36:55.984683Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:55.623189Z","title":"Making chorma features more robust to timber changes","venue":null,"work_id":"e8fa6566-ec72-48d4-9d34-91389419a316","year":2009},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:53.153606Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:bc961ec1b801d5b65380a17a7fc9044b38f46fadd9363c4de76a3a3c2006d786","observation_id":"9f57c45c-423a-40b3-9400-ff259b0e0f05","resolution":{"observed_at":"2026-08-06T22:36:55.727487Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:55.001779Z","title":"EARS: An anechoic fullband speech dataset benchmarked for speech enhancement and dereverbera- tion,","venue":null,"work_id":"0a379ffb-09f3-482b-ac38-6458b84163af","year":2022},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:53.424366Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:02159bfb65d732936ce1cfc4ef34b3388e179eaa098fda30be461afa28399556","observation_id":"2b769b2d-e693-43e8-a747-fdbea6bcf71f","resolution":{"observed_at":"2026-08-06T22:36:55.076894Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:54.853787Z","title":"Echo-aware adaptation of sound event localization and detection in unknown envi- ronments,","venue":null,"work_id":"c8196a45-eeee-4b3c-84cd-061950ef0680","year":2022},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:53.529845Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:818e9e26efcfdb52c0643f9b5b34854818c999f37480136ba9cf70a199b471c9","observation_id":"8c429450-61fb-43f4-bc84-13cd888635b0","resolution":{"observed_at":"2026-08-06T22:36:54.896198Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:54.715476Z","title":"ESC: Dataset for environmental sound classi- fication,","venue":null,"work_id":"c5d03ac7-86b9-4c23-ad5d-264b4837da39","year":2015},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:53.614307Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:d10458be287f4012dcb1cfa91ff0afddab65645b9def7d358779a5132f9572dd","observation_id":"ad108375-86e0-4399-a657-6c546a9420cc","resolution":{"observed_at":"2026-08-06T22:36:54.800161Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:54.519957Z","title":"13, 2025)","venue":null,"work_id":"b2ca5b51-5506-41d5-b4dc-a666345a3c92","year":2025},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:53.715954Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:6d1607cf146d323c674b67efce199666ef841c34b03f95d12b6b15639c89479f","observation_id":"8be1995d-3557-4b11-9a06-cda8c31233ba","resolution":{"observed_at":"2026-08-06T22:36:54.613294Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:53.792270Z","title":"Audio set: An ontology and human- labeled dataset for audio events,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:53.792270Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:3ff9da7e739dcac80c8734e3d298454c575a8ad9c28aa5ee6af4fdeacb124f45","observation_id":"a93be642-2af9-4dd4-a2f8-3f88dbf1fad2","resolution":{"observed_at":"2026-08-06T22:36:53.792270Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:54.207168Z","title":"Masked modeling duo: Towards a universal audio pre-training fram ework,","venue":null,"work_id":"56622180-e69f-4ac0-aa9f-853287ad5ea0","year":2024},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:53.838897Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:d25d8c8a9e3cb58db2848c2705d78f3f643becb93c5e40265b62b90fd3844c06","observation_id":"8abe9694-6c7f-4345-a470-a39c478645ad","resolution":{"observed_at":"2026-08-06T22:36:54.357859Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:36:54.015181Z","title":"13, 2025)","venue":null,"work_id":"836030de-3323-49eb-8043-791204f6976c","year":2025},"citing_paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T22:36:53.842471Z"},"links":{"citing_paper":"/paper/2506.21174"},"observation_digest":"sha256:8a3c3ea2f584d5bda7f3602da49cf68b162cf00c009975d2a2681594bee8fc08","observation_id":"5b41a0eb-11a9-4b2a-826a-8d0bb5f47fb4","resolution":{"observed_at":"2026-08-06T22:36:54.087029Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.21174","last_updated":"2025-06-26T12:27:52Z","latest_version":1,"primary_category":"eess.AS","snapshot_observed_at":"2026-08-06T22:30:07.344322Z","submitted_at":"2025-06-26T12:27:52Z","title":"Performance improvement of spatial semantic segmentation with enriched audio features and agent-based error correction for DCASE 2025 Challenge Task 4"},"reference_resolution":{"displayed":21,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":5,"verified_exact":0,"verified_fuzzy":15},"total_outbound_references":21},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 21 of 21 outbound references and 0 inbound Pith citation observations for arXiv:2506.21174."}