{"as_of":"2026-08-08T23:44:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:a925d0a2863101edc2b815312b767f0d030f481575de306173d35bba90b1f3ee","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":22,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":22,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":22,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":22,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T19:20:50.875290Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T08:29:41.902212Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2501.09747","last_updated":"2025-01-16T18:57:04Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-16T18:57:04Z","title":"FAST: Efficient Action Tokenization for Vision-Language-Action Models","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-05-11T08:52:31.686474Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2501.09747"},"observation_digest":"sha256:383b96e846a253e0361efc2f46b0eb3683d85264586401c897a786fcccdd598b","observation_id":"79955a72-0acc-47f7-a02a-8a740e4e8185","resolution":{"observed_at":"2026-05-11T08:52:32.118854Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-08-06T19:20:50.875290Z","title":"Soundstream: An end-to-end neural audio codec,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.06040","last_updated":"2025-07-08T14:41:42Z","snapshot_observed_at":"2026-08-07T22:25:38.892483Z","submitted_at":"2025-07-08T14:41:42Z","title":"EdgeCodec: Onboard Lightweight High Fidelity Neural Compressor with Residual Vector Quantization","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T19:20:50.875290Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2507.06040"},"observation_digest":"sha256:2bc61ea62a5ac881a975a560db14753ad3eb10f09233afcffd5165228412d3d8","observation_id":"47a07caf-a475-4e70-85fe-b4cde9d82706","resolution":{"observed_at":"2026-08-06T19:20:50.875290Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-08-06T18:28:30.286712Z","title":"Zeghidour, A","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.08236","last_updated":"2025-07-11T00:47:08Z","snapshot_observed_at":"2026-08-06T18:20:54.021233Z","submitted_at":"2025-07-11T00:47:08Z","title":"Distilling Spectrograms into Tokens: Fast and Lightweight Bioacoustic Classification for BirdCLEF+ 2025","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T18:28:30.286712Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2507.08236"},"observation_digest":"sha256:36f3c9d755c19b10bb230761d9c74d6641aeace384ea9911062e209701554cb8","observation_id":"c2fefea6-c362-4b77-a9d6-e5da99af7fa5","resolution":{"observed_at":"2026-08-06T18:28:30.286712Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-08-05T14:36:37.718921Z","title":"SoundStream: An End-to-End Neural Audio Codec,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.21153","last_updated":"2025-08-28T18:38:42Z","snapshot_observed_at":"2026-08-07T14:49:09.772290Z","submitted_at":"2025-08-28T18:38:42Z","title":"WaveLLDM: Design and Development of a Lightweight Latent Diffusion Model for Speech Enhancement and Restoration","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-05T14:36:37.718921Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2508.21153"},"observation_digest":"sha256:e494372ea4778855b4ab7ccd68e91e4723872b51fc4197c678be5955fb777b76","observation_id":"f8aa1a5a-b4ef-4032-9bd8-cfdf23931ef9","resolution":{"observed_at":"2026-08-05T14:36:37.718921Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-08-05T11:41:03.517998Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2509.02349","last_updated":"2025-09-04T14:25:57Z","snapshot_observed_at":"2026-08-08T05:41:53.428522Z","submitted_at":"2025-09-02T14:15:22Z","title":"AudioCodecBench: A Comprehensive Benchmark for Audio Codec Evaluation","version":2},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-05T11:41:03.517998Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2509.02349"},"observation_digest":"sha256:ff04ca48298cc1e6a5e6630ffba8ed5db4f8d19dc136cf81c82b9fa347c3c00c","observation_id":"f7a8ed48-8e2f-4b77-985e-251722f8d049","resolution":{"observed_at":"2026-08-05T11:41:03.517998Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-08-05T11:29:01.660239Z","title":"Zeghidour, A","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2509.02771","last_updated":"2025-09-02T19:20:06Z","snapshot_observed_at":"2026-08-07T13:23:05.616950Z","submitted_at":"2025-09-02T19:20:06Z","title":"Analysis of Speaker Verification Performance Trade-offs with Neural Audio Codec Transmission","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-05T11:29:01.660239Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2509.02771"},"observation_digest":"sha256:131d5b8b4bea6f8dd7c975ab4b29159a920c3b459f210f8136d17936449d2792","observation_id":"a76137b1-16cc-4475-9fc3-80d93e4a3337","resolution":{"observed_at":"2026-08-05T11:29:01.660239Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2604.01929","last_updated":"2026-04-29T10:54:09Z","snapshot_observed_at":"2026-07-06T22:51:40.181565Z","submitted_at":"2026-04-02T11:49:00Z","title":"Woosh: A Sound Effects Foundation Model","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-13T20:51:08.144573Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2604.01929"},"observation_digest":"sha256:c094fe54a05fd81073eb10d5177e36a7b780ceeb9c8e399c1b48ef473d3fdb87","observation_id":"797383dd-4533-42af-9127-d130d7dcafc7","resolution":{"observed_at":"2026-05-13T20:53:15.789605Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2604.12965","last_updated":"2026-04-14T16:59:03Z","snapshot_observed_at":"2026-08-06T12:38:28.688700Z","submitted_at":"2026-04-14T16:59:03Z","title":"Efficient Retrieval Scaling with Hierarchical Indexing for Large Scale Recommendation","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-05-10T14:25:53.909672Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2604.12965"},"observation_digest":"sha256:683dba5e8ab4698cacad45a10ca9c96bf67928d076b1336148a12a3f5de3ae77","observation_id":"2b0019ce-021a-4bee-a712-34ed89a03939","resolution":{"observed_at":"2026-05-11T11:36:01.029236Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2605.13789","last_updated":"2026-05-14T01:44:35Z","snapshot_observed_at":"2026-07-30T00:30:13.289148Z","submitted_at":"2026-05-13T17:08:41Z","title":"ENSEMBITS: an alphabet of protein conformational ensembles","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-14T19:11:29.935218Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2605.13789"},"observation_digest":"sha256:b23525d0945d95ee9d99923067f908ec66e8212d903dd4e417e6b70c2a620e05","observation_id":"cd7b13a5-b159-4f4f-b0c1-75ee70cb3f54","resolution":{"observed_at":"2026-05-14T19:12:50.731283Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2605.13789","last_updated":"2026-05-14T01:44:35Z","snapshot_observed_at":"2026-07-30T00:30:13.289148Z","submitted_at":"2026-05-13T17:08:41Z","title":"ENSEMBITS: an alphabet of protein conformational ensembles","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-15T04:51:32.821791Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2605.13789"},"observation_digest":"sha256:59553b20538ade96d43b0386917b9c3f0bce66a57212bfc0d7d094e2b6ffb415","observation_id":"732d25d3-0863-4f1e-a2f9-a5f8ba34b17e","resolution":{"observed_at":"2026-05-15T04:55:03.987291Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2605.20519","last_updated":"2026-05-21T17:49:53Z","snapshot_observed_at":"2026-07-06T23:31:03.878152Z","submitted_at":"2026-05-19T21:39:52Z","title":"Codec-Robust Attacks on Audio LLMs","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-21T06:43:52.735211Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2605.20519"},"observation_digest":"sha256:125a2de92ce0b996e8fbb87b7196d1fd5481a782046b7bdd9c9ff1147f0db1b4","observation_id":"854eea0f-0743-49e4-8e56-6a0053862c56","resolution":{"observed_at":"2026-05-21T06:44:00.663100Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2605.20519","last_updated":"2026-05-21T17:49:53Z","snapshot_observed_at":"2026-07-06T23:31:03.878152Z","submitted_at":"2026-05-19T21:39:52Z","title":"Codec-Robust Attacks on Audio LLMs","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-25T05:44:46.831360Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2605.20519"},"observation_digest":"sha256:13012ae003008adc9ff1efb777bd247ff812b2b5abba5aa7b9c645c3d6635194","observation_id":"592f59bb-4ce9-4807-b40e-97ec1e5107ac","resolution":{"observed_at":"2026-05-25T05:45:23.656482Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2605.20649","last_updated":"2026-05-20T03:09:45Z","snapshot_observed_at":"2026-08-04T03:27:22.737171Z","submitted_at":"2026-05-20T03:09:45Z","title":"AMAR: Lightweight Attention-Based Multi-User Activity Recognition from Wi-Fi CSI","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-21T03:11:44.222945Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2605.20649"},"observation_digest":"sha256:e4ed8f15c16c3dcffec211653b1c8cf93f3478df4ca6c22fcac91988b7dae9cd","observation_id":"1c163d7f-68ac-47d0-8275-0921e3cd588b","resolution":{"observed_at":"2026-05-21T03:13:56.115057Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2605.21081","last_updated":"2026-05-20T12:16:28Z","snapshot_observed_at":"2026-08-06T01:21:51.969876Z","submitted_at":"2026-05-20T12:16:28Z","title":"Musical Attention Transformer: Music Generation Using a Music-Specific Attention Model","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-21T01:44:20.282896Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2605.21081"},"observation_digest":"sha256:f7b09f505931def7c3357ddf9149b3094643e98b3192591d65b4ae5fb8f0a0b2","observation_id":"57f43a6d-5d8e-4b41-9c21-f261757f3f0c","resolution":{"observed_at":"2026-05-21T01:44:22.778590Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2606.02631","last_updated":"2026-05-30T14:59:57Z","snapshot_observed_at":"2026-08-08T06:42:18.072725Z","submitted_at":"2026-05-30T14:59:57Z","title":"Wavelet as Tokenizer: Preliminary Results on a Shared Wavelet Token Schema for Natural Signals","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-28T18:05:11.392504Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2606.02631"},"observation_digest":"sha256:4ea417b40a2e5682f24d85a3facff8bed06d4d874db65110e9a65baafaa45cac","observation_id":"a7b0c50a-f784-47ab-a015-258f2a3c6361","resolution":{"observed_at":"2026-07-01T20:46:13.380021Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2606.02739","last_updated":"2026-06-01T18:05:18Z","snapshot_observed_at":"2026-07-06T23:43:07.940839Z","submitted_at":"2026-06-01T18:05:18Z","title":"EntangleCodec: A Unified Discrete Audio Tokenizer via Semantic-Acoustic Entanglement","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-06-28T12:34:06.024192Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2606.02739"},"observation_digest":"sha256:f3527c4618dbd4d8462dc383832949c6bd072f8f9542f1c1e7cea26eb99100b5","observation_id":"51248e43-eef7-4c61-a155-310d1a5acc03","resolution":{"observed_at":"2026-07-02T01:06:24.820994Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2606.06357","last_updated":"2026-06-04T16:25:07Z","snapshot_observed_at":"2026-07-06T23:46:10.612537Z","submitted_at":"2026-06-04T16:25:07Z","title":"F3-Tokenizer: Taming Audio Autoencoder Latents for Understanding and Generation","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-27T23:36:27.369551Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2606.06357"},"observation_digest":"sha256:5ffc2a8eac6cc9862649119577f0c95d9650fdab08492bc25134d9026132c80c","observation_id":"169498f5-45de-4c48-b16c-6ba7b734649d","resolution":{"observed_at":"2026-07-02T15:47:06.019471Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2606.21893","last_updated":"2026-06-20T05:58:50Z","snapshot_observed_at":"2026-08-06T08:56:11.950592Z","submitted_at":"2026-06-20T05:58:50Z","title":"AugCodec: A Low-Bitrate Disentangled Neural Speech Codec via Data Augmentation","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-06-26T11:37:39.960212Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2606.21893"},"observation_digest":"sha256:4f15149760c1bdbd913f2792a80fbc2cc084c2b69d740ac7f06ad70d89f87cf3","observation_id":"c54b9722-9099-4269-a7bd-ec2b71ab558d","resolution":{"observed_at":"2026-07-04T08:29:41.903636Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2606.28779","last_updated":"2026-06-27T07:14:04Z","snapshot_observed_at":"2026-07-07T00:02:47.924620Z","submitted_at":"2026-06-27T07:14:04Z","title":"Telephony Voice Agent for Banking Services","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-30T08:57:32.212566Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2606.28779"},"observation_digest":"sha256:e0143ba4e5a4422e1d8d106efa28182f72c27f0d9402f82e9631225300d076bf","observation_id":"04b19d85-4832-4674-a219-a630ceffafca","resolution":{"observed_at":"2026-06-30T09:04:32.871308Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2607.02266","last_updated":"2026-07-02T14:51:42Z","snapshot_observed_at":"2026-08-08T15:31:37.930459Z","submitted_at":"2026-07-02T14:51:42Z","title":"HERMES: A Multi-Granularity Labeling Substrate for Pre-training Data Mixtures","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-07-03T16:44:41.720388Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2607.02266"},"observation_digest":"sha256:6895463868098b34daefa12531db2bea462f00316b2e2d51b868c62a2853e0e9","observation_id":"fb7580f2-7310-4645-9a55-bba81d679027","resolution":{"observed_at":"2026-07-03T16:48:39.426776Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-11T21:55:00.514210Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.04068","last_updated":"2026-07-05T01:00:58Z","snapshot_observed_at":"2026-07-11T21:55:00.070098Z","submitted_at":"2026-07-05T01:00:58Z","title":"UniSGR: Unified Framework for Semantic ID Generation and Ranking","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-07-11T21:55:00.514210Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2607.04068"},"observation_digest":"sha256:449b13408a9c6fcfa7d5a29ddb11f28de0ed990f29474b16f631ba6afbcea453","observation_id":"9cf854e5-57ba-441c-b304-1deda6cfeccb","resolution":{"observed_at":"2026-07-11T21:55:00.514210Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-08-02T06:23:32.601187Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.14148","last_updated":"2026-07-14T13:05:53Z","snapshot_observed_at":"2026-08-07T05:42:24.503168Z","submitted_at":"2026-07-14T13:05:53Z","title":"ITGPT: A Transformer Based Architecture for the Generation of Dance Dance Revolution and In the Groove Charts","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-02T06:23:32.601187Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2607.14148"},"observation_digest":"sha256:a43caeda0553438f1b847748c4f94d82959cb3717c39c1dae6e38a1ab50ff905","observation_id":"7c1e7c73-db9d-4ad5-8761-1ec6f0521e5f","resolution":{"observed_at":"2026-08-02T06:23:32.601187Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2107.03312/citation-record","integrity":"/paper/2107.03312/integrity","json":"/paper/2107.03312/citation-record.json","paper":"/paper/2107.03312"},"outbound":[],"paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","latest_version":1,"primary_category":"cs.SD","snapshot_observed_at":"2026-08-06T14:11:39.295017Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 22 inbound Pith citation observations for arXiv:2107.03312."}