{"as_of":"2026-08-22T03:54:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:eb86011f88b27beab8e717c170a243fe9d1e046a5a43a8e7d1b01e708bd1ebf6","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":6,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":6,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-21T06:32:19.484+00:00","state":"measured"},{"denominator":6,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":6,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T13:13:31.174390Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2410.07168","last_updated":"2025-03-02T09:16:05Z","snapshot_observed_at":"2026-08-16T13:10:59.428090Z","submitted_at":"2024-10-09T17:59:04Z","title":"Sylber: Syllabic Embedding Representation of Speech from Raw Audio","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.07168","snapshot_observed_at":"2026-08-07T13:13:31.174390Z","title":"Sylber: Syllabic embedding representation of speech from raw audio,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.22455","last_updated":"2025-05-28T15:12:53Z","snapshot_observed_at":"2026-08-15T22:31:30.601359Z","submitted_at":"2025-05-28T15:12:53Z","title":"Articulatory modeling of the S-shaped F2 trajectories observed in \\\"Ohman's spectrographic analysis of VCV syllables","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T13:13:31.174390Z"},"links":{"cited_paper":"/paper/2410.07168","citing_paper":"/paper/2505.22455"},"observation_digest":"sha256:4d7a586a0fd51eadf115bb8b6820a961803a3f05e001b73385ce5b7b52d218ed","observation_id":"52e671f3-f547-4dad-a99e-2d9ea79f1b6c","resolution":{"observed_at":"2026-08-07T13:13:31.174390Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.07168","last_updated":"2025-03-02T09:16:05Z","snapshot_observed_at":"2026-08-16T13:10:59.428090Z","submitted_at":"2024-10-09T17:59:04Z","title":"Sylber: Syllabic Embedding Representation of Speech from Raw Audio","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.07168","snapshot_observed_at":"2026-08-07T12:27:25.946050Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24496","last_updated":"2025-05-30T11:47:29Z","snapshot_observed_at":"2026-08-15T08:11:51.259191Z","submitted_at":"2025-05-30T11:47:29Z","title":"Speech Token Prediction via Compressed-to-fine Language Modeling for Speech Generation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T12:27:25.946050Z"},"links":{"cited_paper":"/paper/2410.07168","citing_paper":"/paper/2505.24496"},"observation_digest":"sha256:354b369ee582a621b9b517b81b69d168a58b0103ad44c1fdc352019ffcca5b36","observation_id":"ac8039c6-d6e6-4c4d-a396-592356a8a347","resolution":{"observed_at":"2026-08-07T12:27:25.946050Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.07168","last_updated":"2025-03-02T09:16:05Z","snapshot_observed_at":"2026-08-16T13:10:59.428090Z","submitted_at":"2024-10-09T17:59:04Z","title":"Sylber: Syllabic Embedding Representation of Speech from Raw Audio","version":2},"cited_work":{"arxiv_id":"2410.07168","doi":"10.48550/arxiv.2410.07168","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.07168","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Sylber: Syllabic embedding representation of speech from raw audio","venue":"arXiv (Cornell University)","work_id":"5d50b12f-cafd-4dbf-b851-888ff2f21b6b","year":2024},"citing_paper":{"arxiv_id":"2604.16445","last_updated":"2026-05-12T09:43:45Z","snapshot_observed_at":"2026-08-16T20:02:56.927930Z","submitted_at":"2026-04-07T13:24:13Z","title":"SAND: The Challenge on Speech Analysis for Neurodegenerative Disease Assessment","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-10T18:40:48.363579Z"},"links":{"cited_paper":"/paper/2410.07168","citing_paper":"/paper/2604.16445"},"observation_digest":"sha256:a51836e30c96ab6dd791254262ad0005653b9a5f72895efb1dcc0c94c7daddd4","observation_id":"0617091d-28fe-484e-a7a6-b382c23c19ec","resolution":{"observed_at":"2026-05-11T00:10:51.101256Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.07168","last_updated":"2025-03-02T09:16:05Z","snapshot_observed_at":"2026-08-16T13:10:59.428090Z","submitted_at":"2024-10-09T17:59:04Z","title":"Sylber: Syllabic Embedding Representation of Speech from Raw Audio","version":2},"cited_work":{"arxiv_id":"2410.07168","doi":"10.48550/arxiv.2410.07168","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.07168","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Sylber: Syllabic embedding representation of speech from raw audio","venue":"arXiv (Cornell University)","work_id":"5d50b12f-cafd-4dbf-b851-888ff2f21b6b","year":2024},"citing_paper":{"arxiv_id":"2604.16445","last_updated":"2026-05-12T09:43:45Z","snapshot_observed_at":"2026-08-16T20:02:56.927930Z","submitted_at":"2026-04-07T13:24:13Z","title":"SAND: The Challenge on Speech Analysis for Neurodegenerative Disease Assessment","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-13T07:55:39.622962Z"},"links":{"cited_paper":"/paper/2410.07168","citing_paper":"/paper/2604.16445"},"observation_digest":"sha256:e309ce48ef3c6fd4bca5dfabea56d554489c052c92eb78906de8b4b64a8d0132","observation_id":"c6bcdabb-acdb-46da-8e56-aec662200c84","resolution":{"observed_at":"2026-05-13T07:57:31.421612Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.07168","last_updated":"2025-03-02T09:16:05Z","snapshot_observed_at":"2026-08-16T13:10:59.428090Z","submitted_at":"2024-10-09T17:59:04Z","title":"Sylber: Syllabic Embedding Representation of Speech from Raw Audio","version":2},"cited_work":{"arxiv_id":"2410.07168","doi":"10.48550/arxiv.2410.07168","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.07168","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Sylber: Syllabic embedding representation of speech from raw audio","venue":"arXiv (Cornell University)","work_id":"5d50b12f-cafd-4dbf-b851-888ff2f21b6b","year":2024},"citing_paper":{"arxiv_id":"2605.01381","last_updated":"2026-05-02T11:08:10Z","snapshot_observed_at":"2026-08-16T03:40:49.345088Z","submitted_at":"2026-05-02T11:08:10Z","title":"A framework for analyzing concept representations in neural models","version":1},"reference_index":181,"source":"arxiv_source","source_observed_at":"2026-05-09T14:49:22.776209Z"},"links":{"cited_paper":"/paper/2410.07168","citing_paper":"/paper/2605.01381"},"observation_digest":"sha256:1bba590e0df32500001428cdc41cccb0ea543836cd0b4c38fc370cf30a7ae8f4","observation_id":"fc205fb7-2739-40d3-8500-9d0fbbc34c9a","resolution":{"observed_at":"2026-05-09T22:18:58.939689Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.07168","last_updated":"2025-03-02T09:16:05Z","snapshot_observed_at":"2026-08-16T13:10:59.428090Z","submitted_at":"2024-10-09T17:59:04Z","title":"Sylber: Syllabic Embedding Representation of Speech from Raw Audio","version":2},"cited_work":{"arxiv_id":"2410.07168","doi":"10.48550/arxiv.2410.07168","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.07168","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Sylber: Syllabic embedding representation of speech from raw audio","venue":"arXiv (Cornell University)","work_id":"5d50b12f-cafd-4dbf-b851-888ff2f21b6b","year":2024},"citing_paper":{"arxiv_id":"2606.31247","last_updated":"2026-06-30T07:24:10Z","snapshot_observed_at":"2026-08-13T09:35:09.713263Z","submitted_at":"2026-06-30T07:24:10Z","title":"FlexiSLM: A Dynamic and Controllable Frame Rate Spoken Language Model","version":1},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-07-01T03:50:26.873406Z"},"links":{"cited_paper":"/paper/2410.07168","citing_paper":"/paper/2606.31247"},"observation_digest":"sha256:4afc2b42560173f520d767f2dd49d7184c3111cd87b6fcd74dadcbf6afe74510","observation_id":"0572e94a-562b-4c55-b2b8-fd9d0a7d2f01","resolution":{"observed_at":"2026-07-01T11:55:42.171160Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2410.07168/citation-record","integrity":"/paper/2410.07168/integrity","json":"/paper/2410.07168/citation-record.json","paper":"/paper/2410.07168"},"outbound":[],"paper":{"arxiv_id":"2410.07168","last_updated":"2025-03-02T09:16:05Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-16T13:10:59.428090Z","submitted_at":"2024-10-09T17:59:04Z","title":"Sylber: Syllabic Embedding Representation of Speech from Raw Audio"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"thesis":"As of 22 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 6 inbound Pith citation observations for arXiv:2410.07168."}