{"as_of":"2026-08-13T12:48:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:ad8a0b2ffff46253061c14e55259a2d9236d2e87fe72b91a1d6670f5b1480fcb","coverage":[{"denominator":48,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":48,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T21:14:15.382375Z","state":"measured"},{"denominator":49,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":49,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T21:14:12.433222Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-04T21:14:15.628368Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"cited_work":{"arxiv_id":"2509.08173","doi":null,"metadata_source":"pith","pith_arxiv_id":"2509.08173","snapshot_observed_at":"2026-08-04T21:14:15.628368Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","venue":"eess.AS","work_id":"d35b9b16-fe39-468b-a2bd-80e8fa66af79","year":2025},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:12.433222Z"},"links":{"cited_paper":"/paper/2509.08173","citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:206ea2d52f0911abadcd4d907078055f25264b72da0a49d27f7fb01771229b74","observation_id":"e9dfd7fa-8bc1-456a-9c70-599f69ef689e","resolution":{"observed_at":"2026-08-04T21:14:15.681500Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2509.08173/citation-record","integrity":"/paper/2509.08173/integrity","json":"/paper/2509.08173/citation-record.json","paper":"/paper/2509.08173"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.269408Z","title":null,"venue":null,"work_id":"7d307d9d-c6cc-4bba-bd63-99078a88fb3b","year":null},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:12.271450Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:6e7bc212f66e15088df9880ea647b293d3dc048edc3d011f52733312f4724b5e","observation_id":"e5c0ae98-f2d8-4c03-8a93-3b59e89286ad","resolution":{"observed_at":"2026-08-04T21:14:18.274398Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"cited_work":{"arxiv_id":"2509.08173","doi":null,"metadata_source":"pith","pith_arxiv_id":"2509.08173","snapshot_observed_at":"2026-08-04T21:14:15.628368Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","venue":"eess.AS","work_id":"d35b9b16-fe39-468b-a2bd-80e8fa66af79","year":2025},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:12.433222Z"},"links":{"cited_paper":"/paper/2509.08173","citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:206ea2d52f0911abadcd4d907078055f25264b72da0a49d27f7fb01771229b74","observation_id":"e9dfd7fa-8bc1-456a-9c70-599f69ef689e","resolution":{"observed_at":"2026-08-04T21:14:15.681500Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.241831Z","title":"Attribute Recognition and Knowledge Integration The proposed bottom-up framework is illustrated in Figure 1","venue":null,"work_id":"9e2b954c-3460-4961-9295-6f8782d752e0","year":null},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:12.482305Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:a110287ecd3e3bd9c11f7d46373d1ec8459097ea9e047c925f0900c5aa2647e2","observation_id":"e049eea2-f200-40c5-ad8f-56946dd76739","resolution":{"observed_at":"2026-08-04T21:14:18.246637Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.256068Z","title":null,"venue":null,"work_id":"4fd0916f-5921-4aef-9a0c-5f1c70a02c24","year":null},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:12.368974Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:3f3f84335d53705f3893cbc2bb514264878f5098265354f525cf8e26147b56d7","observation_id":"97795ed0-5a65-4ff4-a32f-b9cfeb12a9ab","resolution":{"observed_at":"2026-08-04T21:14:18.260725Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.227660Z","title":"basic5000","venue":null,"work_id":"30b886eb-6faa-48a8-b78b-4ff29c1e892b","year":null},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:12.547785Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:41bf13c4ffbf994f482225a0e264375d231fc6448fac8c850f82cdfa705043a4","observation_id":"6ba4e3bc-6e83-4dd9-9b97-ee80a824301f","resolution":{"observed_at":"2026-08-04T21:14:18.232151Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.113153Z","title":"Speech recognition by machines and humans,","venue":null,"work_id":"755e8fd1-8f28-4d78-9139-5a0c3675e7ad","year":1997},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.136563Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:28fea0533117f626168afa71d682d69a96de857c2d1e8dfd4ccacaac7e003c3c","observation_id":"30dae91a-7cb7-4d3f-a6a2-7c6afb97bfa0","resolution":{"observed_at":"2026-08-04T21:14:18.117690Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.197775Z","title":null,"venue":null,"work_id":"5a70932f-d8b1-4e62-89e5-90c63c66b7cc","year":null},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:12.694374Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:6774f6e65969e77681b61e4a2872375c0f5eaffbc307adc57d5f1cda133aca22","observation_id":"47765219-fd24-475b-8517-7be0568c4b43","resolution":{"observed_at":"2026-08-04T21:14:18.202976Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.183056Z","title":"Continuous speech recognition by statistical methods,","venue":null,"work_id":"1ee8fa4b-bafb-4ce8-925a-3bbec22d6191","year":1976},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:12.776136Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:aacdb1efa74a9b71cbfbb11fe41afe1115b52d3c2b16db99b1041e9fec162242","observation_id":"ebf116fd-b126-401a-917d-3da8161ac9a8","resolution":{"observed_at":"2026-08-04T21:14:18.187774Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.169239Z","title":"The kaldi speech recognition toolkit,","venue":null,"work_id":"1fe19cc6-ac3b-4f97-bfb3-90e3630306c6","year":2011},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:12.878798Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:a285957a8c427cb4ff9450537dd97f59f11c549f44301a3d04ba8ba6738fe8ef","observation_id":"f8ca8fbf-c283-4815-b70e-10f240c2bbc5","resolution":{"observed_at":"2026-08-04T21:14:18.174279Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.155351Z","title":"Large-vocabulary speaker-independent con- tinuous speech recognition using hmm,","venue":null,"work_id":"bcbd47b5-c2b2-4ffe-8072-0acabe6bad30","year":1988},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:12.948656Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:5489601c99eb5e0a39e2c823d6520518f861280d94b89c73c98931e0e5296cb2","observation_id":"4bc0ba0a-f0ae-4287-9dde-a4c2b6730c12","resolution":{"observed_at":"2026-08-04T21:14:18.159896Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.141542Z","title":"An information- extraction approach to speech processing: Analysis, detection, verifi- cation, and recognition,","venue":null,"work_id":"02bc1378-befe-4b66-9e43-68e05dcf7ba1","year":2013},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.044039Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:4f45ea8dd19dd0d1e66bb8652b155daeb6d915edcad8d58a807acc673e9ede1b","observation_id":"5c358c93-1c48-4aa4-af43-f26eb0cccaed","resolution":{"observed_at":"2026-08-04T21:14:18.146102Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.126957Z","title":"Allen,How do Humans Process and Recognize Speech?, Springer US, Boston, MA, 1995","venue":null,"work_id":"dc3b5aea-5fe0-4e0f-be6c-5dc273cca0df","year":1995},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.088873Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:a0b122424763fd06fdf56d049f8482bb075e2bfc360d8addaaf17332032e08e8","observation_id":"b55f293d-050d-4467-a9b0-58622b132628","resolution":{"observed_at":"2026-08-04T21:14:18.131682Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.015685Z","title":"Combining articulatory and acoustic information for speech recognition in noisy and reverberant environments,","venue":null,"work_id":"98a87114-a13b-42c8-af5c-5b2f74ce5187","year":1998},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.510509Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:3158ac49357bfe0e85db8ffa099805e71093252bd8eac3a417eb6f4e9ba4a759","observation_id":"e4e042fc-f1cd-48b8-b9c6-eb62218f1dfc","resolution":{"observed_at":"2026-08-04T21:14:18.019722Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.098959Z","title":null,"venue":null,"work_id":"0f74943c-c9c0-4bfa-b550-b83a3c0a1cf8","year":1968},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.215478Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:9a6657d72857f87b6c5b537a72117c8e7066bb6170c6d5649a402a9e2bc81947","observation_id":"bfa46a26-a5be-474c-83ad-0846300d7cdb","resolution":{"observed_at":"2026-08-04T21:14:18.103873Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.084873Z","title":"The geometry of phonological features,","venue":null,"work_id":"e16bfe83-e67a-478e-963a-80d5b86dbae5","year":1985},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.244522Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:421f215ce5ce582657b970c700935b79d6e9684e6760ce3c79f398b9c213ccae","observation_id":"29117e98-c383-4fd4-8433-520b5278ec99","resolution":{"observed_at":"2026-08-04T21:14:18.089009Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.070847Z","title":"Hybrid ctc- attention based end-to-end speech recognition using subword units,","venue":null,"work_id":"2a9f0aab-c66b-4315-8096-a736d11092c4","year":2018},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.304437Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:b6eb950ec290c038127f10964fc80bf9f86e844f4d0828b06ab7515a61d6b57e","observation_id":"56a7eb3c-706f-487b-b58c-b62c9d73840d","resolution":{"observed_at":"2026-08-04T21:14:18.075152Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.056979Z","title":"ESPnet: End-to-end speech processing toolkit,","venue":null,"work_id":"5532d695-cda1-41ea-9bdf-5b3bde28c153","year":2018},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.359889Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:f8827e1b2d959c22b2d9693870d51de567e07f9573523bc31fb40dfabd83813f","observation_id":"c73eca58-d0d8-4732-9914-78059185e28e","resolution":{"observed_at":"2026-08-04T21:14:18.061216Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.042888Z","title":"Hybrid ctc/attention architecture for end-to-end speech recognition,","venue":null,"work_id":"bc6794c9-0e17-41ca-bbbd-365d7d2aed3e","year":2017},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.406638Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:8eb38be4f3abf0c068d7e3f36ca3a823ecb1e1aac7e4ab0958418e5eeacafdda","observation_id":"7ae045aa-56f6-43e0-a58f-3a964f59d719","resolution":{"observed_at":"2026-08-04T21:14:18.047846Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.029036Z","title":"Robust speech recognition via large- scale weak supervision,","venue":null,"work_id":"bf6cd9ca-5b8b-41d5-b66a-fb802ba5a51a","year":2022},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.452804Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:cb93170768a9765f67a37da94320c7c1b206313f418e3af9788ed25324730a16","observation_id":"e192ca46-598f-4591-b27f-f96af3b9507c","resolution":{"observed_at":"2026-08-04T21:14:18.033479Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.923763Z","title":"Syllable-based large vocabulary continuous speech recogni- tion,","venue":null,"work_id":"295cf90b-f460-4720-996e-0f6172075c12","year":2001},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.916894Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:4a3f2c4733e84fdf3fff72239a748e930811aaf65cbeafac030ab450f755027d","observation_id":"ed184326-b178-4ef7-ae5d-5de0a2c913a6","resolution":{"observed_at":"2026-08-04T21:14:17.928092Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.001860Z","title":"dissertation, Carnegie Mellon University, Pitts- burgh, PA, USA, 1992","venue":null,"work_id":"8fb066d3-4ade-4b4c-8a1e-5fbd6b2a3ba2","year":1992},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.549163Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:db8a9f50171488579490c279228a9a2be0c4d9db87429b8d2cfe0479484f21fa","observation_id":"1289e771-f6fd-4a4c-b959-3bf68b0e0f09","resolution":{"observed_at":"2026-08-04T21:14:18.006192Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.988665Z","title":"Language-universal speech attributes modeling for zero-shot multilin- gual spoken keyword recognition,","venue":null,"work_id":"6ec8076c-0623-4002-a210-d77930dd1bed","year":2025},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.599452Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:62246cea494a7935e3ee0fe06d9fc02c1250fb201f4e758318ec56dc958ca950","observation_id":"a9c17f3c-fafa-472b-9c1b-034a57494ad2","resolution":{"observed_at":"2026-08-04T21:14:17.992674Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.976366Z","title":"Detection-based asr in the au- tomatic speech attribute transcription project,","venue":null,"work_id":"613ffe33-492e-4854-be62-618676163c9a","year":2007},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.673820Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:d9ee61882ab2e9da0c6d40f78686c0586d5029c631963ddf1507dfc334b9f6d4","observation_id":"80bb1f80-2166-41a3-a713-4fa64b968a37","resolution":{"observed_at":"2026-08-04T21:14:17.980325Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.963232Z","title":"A flexible stream architecture for asr using articulatory features,","venue":null,"work_id":"a412b380-ca32-4f9c-afd0-fd461ef24bb1","year":2002},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.714767Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:28e29e8effc1ca4bbe7e1fbab54a826c4d4c33ceca4c69c975ee12a157fec2c7","observation_id":"03989220-93b6-42bd-9283-c29283cbf881","resolution":{"observed_at":"2026-08-04T21:14:17.967045Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.950880Z","title":"dissertation, Massachusetts Institute of Tech- nology, Cambridge, MA, USA, 1996","venue":null,"work_id":"0e00a6ec-2447-48b0-92d5-3258266e2ad4","year":1996},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.799104Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:da686dd1ce8b32ba1b6d1b4087aceaad2400942b52fd2180af084ac0528ee427","observation_id":"ba7c39c6-494c-4924-b674-65c54e432489","resolution":{"observed_at":"2026-08-04T21:14:17.954888Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.937052Z","title":"An event-based acoustic-phonetic approach to speech segmentation and e-set recogni- tion,","venue":null,"work_id":"0a75472f-cfaf-449f-b2ba-8303f18dbdc1","year":2002},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:13.856497Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:eaac0ca3337fcd4458e69d274d0f1581b983fe8e1c5b21997c2690e815cbeaed","observation_id":"25f8fb4b-f3ea-4c2f-855b-efc69f210710","resolution":{"observed_at":"2026-08-04T21:14:17.941509Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.846050Z","title":"What makes a word: Learning base units in Japanese for speech recognition,","venue":null,"work_id":"c9b99af7-599a-462e-a12c-097fee163700","year":1997},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:14.424643Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:eda362436f759c33bda7ee7f6b3056de1e1d3284470c2da58a33f61caef60222","observation_id":"b79c6b63-1f57-4a5e-8f52-b2e390cc7f09","resolution":{"observed_at":"2026-08-04T21:14:17.849985Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.910875Z","title":"Context-dependent syllable acoustic model for continuous chinese speech recognition,","venue":null,"work_id":"bb4123a8-cab5-4547-8ea3-43e4555637ca","year":2007},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:14.014998Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:70013db5be5b589ffe36da1105474b8bf654a0fe003fbaad40bfdb990aaa2fdc","observation_id":"03858724-e8d2-47ea-939f-df3d8f1cf2d8","resolution":{"observed_at":"2026-08-04T21:14:17.914990Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.897376Z","title":"Syllable-based acoustic modeling with ctc-smbr-lstm,","venue":null,"work_id":"dceafb0b-e508-4f30-9692-831637037cee","year":2017},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:14.098492Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:ad5b093114fbab8dfc423d5ca570357a6677bafd3c05f9563e5753b2e88e285b","observation_id":"d2de5e83-a667-4a55-8962-a73f53603e64","resolution":{"observed_at":"2026-08-04T21:14:17.901323Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1805.06239","last_updated":"2018-05-18T12:46:54Z","snapshot_observed_at":"2026-07-06T06:39:26.409401Z","submitted_at":"2018-05-16T10:43:47Z","title":"A Comparison of Modeling Units in Sequence-to-Sequence Speech Recognition with the Transformer on Mandarin Chinese","version":2},"cited_work":{"arxiv_id":"1805.06239","doi":null,"metadata_source":"pith","pith_arxiv_id":"1805.06239","snapshot_observed_at":"2026-08-04T21:14:15.493292Z","title":"A Comparison of Modeling Units in Sequence-to-Sequence Speech Recognition with the Transformer on Mandarin Chinese","venue":"eess.AS","work_id":"61193539-74dd-403d-8fa2-7028705dc34a","year":2018},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:14.160151Z"},"links":{"cited_paper":"/paper/1805.06239","citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:fed7f9a1c78798358186324fbf80335c11c7d4c046d51422e2af7c6a8723a574","observation_id":"d5971335-de77-49c2-8466-6ce6bef2c373","resolution":{"observed_at":"2026-08-04T21:14:15.533771Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.884143Z","title":"Syllable-based sequence-to-sequence speech recognition with the transformer in man- darin chinese,","venue":null,"work_id":"4e0e0fde-9960-4b40-b7d0-1bcae1d8bde5","year":2018},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:14.244578Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:38a6c829eca72f3710203152fb90f8359ca04c900c486bab1009e39dfa1f0808","observation_id":"3a8637f4-e988-49fe-9d85-23cd830e3937","resolution":{"observed_at":"2026-08-04T21:14:17.888229Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.871179Z","title":"Decoupling recognition and transcription in mandarin asr,","venue":null,"work_id":"bcbef55f-a533-4a41-b7c0-bfb6c2d7b807","year":2021},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:14.311524Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:41c2d0f982dc564e1149702083f8c6d5ba90478ddeecd360f49e520843734ecc","observation_id":"b003f7be-d80b-4b59-a61b-65f8a20c3510","resolution":{"observed_at":"2026-08-04T21:14:17.875150Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.858827Z","title":"The mora and syllable structure in japanese: Evi- dence from speech errors,","venue":null,"work_id":"329e858b-ca3b-449c-8a86-eeaf1f1b8d32","year":1989},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:14.385159Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:c95c2d1aaeaa21d525281fdfe6422bc4231eccf2f7642cf2f7f97901556d7186","observation_id":"e5d5b57d-1451-4dfc-b6ea-489b6153fbfb","resolution":{"observed_at":"2026-08-04T21:14:17.862796Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:16.835279Z","title":"Akamatsu,Japanese phonetics : theory and practice / Tsu- tomu Akamatsu, LINCOM studies in Asian linguistics ; 3","venue":null,"work_id":"0b8f64a7-82f9-4b9f-b663-00105d57f65b","year":1997},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:14.944995Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:fc3539c7b0de726b5ef4b570e905976b89dc93e3b5198f02f8534857e3a3f0e8","observation_id":"a00cee17-f750-4f6a-9854-f9a92e33986c","resolution":{"observed_at":"2026-08-04T21:14:16.997440Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.832785Z","title":"Syllable recognition us- ing syllable-segment statistics and syllable-based hmm,","venue":null,"work_id":"f6cc065a-7aba-47d1-8a18-eaaf7daed6c8","year":2002},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:14.490977Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:81c74cfd8b0683b932cf82a589f037ad332645710556887d2703a17cc04d4dac","observation_id":"b94b1492-422d-4d0b-9d5c-a04168fb3cb0","resolution":{"observed_at":"2026-08-04T21:14:17.837041Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.819471Z","title":"Compari- son of syllable-based and phoneme-based dnn-hmm in japanese speech recognition,","venue":null,"work_id":"333a85b8-61aa-49e5-b43b-af8c9513fa34","year":2014},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:14.571698Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:79910ca5a7a59a0e03f5456d8e61b3bb7dab071db0d182ffbba4f55de1986628","observation_id":"3b7f773f-61ee-45f7-aeb1-ba0bde1fa039","resolution":{"observed_at":"2026-08-04T21:14:17.823654Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.806320Z","title":"Wavlm: Large-scale self-supervised pre-training for full stack speech processing,","venue":null,"work_id":"ac8f787d-7d6d-47ff-b1f3-1a6f02948ed4","year":2022},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:14.629978Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:908870a9b3c72adc03270a297d34a06b5139716405963fe9beb9d1cb0edc8a17","observation_id":"1afdcdb8-812d-4f26-bdfe-e364eed41ed1","resolution":{"observed_at":"2026-08-04T21:14:17.810621Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.737744Z","title":"Fant,Speech Sounds and Features, The MIT Press, 1973","venue":null,"work_id":"60654052-8a3b-4e5d-8126-4d230e312b36","year":1973},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:14.712276Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:954667dc8cadb4fcfba3af331b7fb5ad7d2d5fae3d6a7b2bb6d268f8ff343644","observation_id":"1fc5285c-e264-495b-a59b-e392712be29e","resolution":{"observed_at":"2026-08-04T21:14:17.793815Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.633272Z","title":"Ladefoged and S.F","venue":null,"work_id":"bfe782ad-5a37-4887-9769-6a18a97e7df4","year":2012},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:14.745851Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:bbd46a51f04318d0a76d10d6b92bab0eed2db591bca44c5aea313c0d29c1d929","observation_id":"d5a7fc4f-89ce-4c69-8d92-2fc91e91ecdd","resolution":{"observed_at":"2026-08-04T21:14:17.683076Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:18.212735Z","title":"To support decoding with CTC model, we trained separate KenLM language models tailored to each modeling unit","venue":null,"work_id":"8b7a7f44-8fa2-4684-bcb4-2e0509ae4145","year":null},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:12.624087Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:b1893a42be6d78719588eace285b7bb61e99d79c230ef063e1671be977e288e2","observation_id":"ae59aec8-fcb1-4bc5-928a-d9a6377236d7","resolution":{"observed_at":"2026-08-04T21:14:18.218047Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:17.258284Z","title":"Modeling linguistic fea- tures in speech recognition,","venue":null,"work_id":"8699712d-ad0d-4b50-b333-f64889a1dfb1","year":2003},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:14.861086Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:12f136453b4e4742233aedabe64822d0b0f2252956b0303c3a86d93b93f05303","observation_id":"dfb80fda-59ee-4ce8-a97c-e7d1722a2893","resolution":{"observed_at":"2026-08-04T21:14:17.463595Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:16.534598Z","title":"Acoustic cues of the stop voicing contrast in mod- ern tokyo japanese,","venue":null,"work_id":"01ac3892-cce5-4093-9415-4672a4aada53","year":2018},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:14.987060Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:73b868d519424e64729b52ca7fc593ed28d5336641cee5a43375a38dfa698a5d","observation_id":"eb0edb80-68e6-441f-b3c1-d970e386b300","resolution":{"observed_at":"2026-08-04T21:14:16.652414Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:16.401389Z","title":"Syllable-based acoustic modeling for japanese spontaneous speech recognition,","venue":null,"work_id":"811c3063-73d3-4207-9ace-562ac82696b8","year":2003},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:15.069186Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:1f31fd75db3f22c9b8f7f3c0fc73f86fa6837a912f8c7db88a8d124ef6c37501","observation_id":"a970f853-17bc-488b-914b-ac9f865889b1","resolution":{"observed_at":"2026-08-04T21:14:16.476687Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:16.212950Z","title":"Aishell- 1: An open-source mandarin speech corpus and a speech recognition baseline,","venue":null,"work_id":"f78ddca8-9425-4d3f-8e69-816710616f4e","year":2017},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:15.102927Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:e2ebb5041e8b2ec628b0c0fafb3278020734fbec68df093a97dadd171ebb6cb0","observation_id":"2329b416-c6e2-4b47-a9ba-7663f0922ff4","resolution":{"observed_at":"2026-08-04T21:14:16.289241Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1711.00354","last_updated":"2017-10-28T05:28:01Z","snapshot_observed_at":"2026-08-10T14:44:27.962270Z","submitted_at":"2017-10-28T05:28:01Z","title":"JSUT corpus: free large-scale Japanese speech corpus for end-to-end speech synthesis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1711.00354","snapshot_observed_at":"2026-08-04T21:14:15.181943Z","title":"JSUT corpus: free large-scale japanese speech corpus for end-to-end speech synthesis,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:15.181943Z"},"links":{"cited_paper":"/paper/1711.00354","citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:0e0ee66b6e9167b6bfc541ed72428c73cff00de592342e873244957e4c56123c","observation_id":"4554dadd-6888-4231-b273-01e35ad8b690","resolution":{"observed_at":"2026-08-04T21:14:15.181943Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:16.074431Z","title":"Atten- tion is all you need,","venue":null,"work_id":"b9023565-8ba4-4edb-9c8b-d5a7f5b6c349","year":2017},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:15.244907Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:935f83783c55963e5948d72be9af3eb88c508b5b22bbf3c8d79f2392e9b2027b","observation_id":"ff822c4a-456e-4a59-aa1e-eb8869cec52c","resolution":{"observed_at":"2026-08-04T21:14:16.152917Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:15.932748Z","title":"Decoupled weight decay regular- ization,","venue":null,"work_id":"d5b2c9ff-c47a-45f5-aa5e-95b8648ae427","year":2019},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:15.306161Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:674d1b59e0eb588420701b50598069f88e8016989d97c37d42d9caa9495ea9d4","observation_id":"8661e1ac-fd37-4922-ab76-4868086bd10d","resolution":{"observed_at":"2026-08-04T21:14:16.013203Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:14:15.744999Z","title":"Mls: A large-scale multilingual dataset for speech research,","venue":null,"work_id":"18fc9df4-2be3-4b90-a1ef-7de0a5f8d8cd","year":2020},"citing_paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-04T21:14:15.382375Z"},"links":{"citing_paper":"/paper/2509.08173"},"observation_digest":"sha256:38c9ae58ef93c32aa9566798a5976c2583da22f55a000b21f77c13c0b37436c9","observation_id":"16060b58-cf3e-42c6-ab77-49e6808e6d6c","resolution":{"observed_at":"2026-08-04T21:14:15.825217Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2509.08173","last_updated":"2025-09-09T22:20:38Z","latest_version":1,"primary_category":"eess.AS","snapshot_observed_at":"2026-08-07T14:19:49.286136Z","submitted_at":"2025-09-09T22:20:38Z","title":"A Bottom-up Framework with Language-universal Speech Attribute Modeling for Syllable-based ASR"},"reference_resolution":{"displayed":48,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":5,"verified_exact":2,"verified_fuzzy":41},"total_outbound_references":48},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 48 of 48 outbound references and 1 inbound Pith citation observation for arXiv:2509.08173."}