{"as_of":"2026-08-23T03:20:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:e67c4cfc92d99cf35efba41246b1926120c763dc809b187f972f8fd0bf3585f4","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":23,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":23,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-22T06:32:14.747728+00:00","state":"measured"},{"denominator":23,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":23,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T00:03:40.478721Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-05T13:21:06.391412Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":"2308.11466","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-07-05T13:21:06.391412Z","title":"Sonar: Sentence-level multimodal and language-agnostic representations","venue":"cs.CL","work_id":"74071e8e-19c2-44e6-9972-e052b130c9d5","year":2023},"citing_paper":{"arxiv_id":"2401.08281","last_updated":"2025-10-23T09:36:08Z","snapshot_observed_at":"2026-07-31T05:45:37.385210Z","submitted_at":"2024-01-16T11:12:36Z","title":"The Faiss library","version":4},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-12T01:47:19.947054Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2401.08281"},"observation_digest":"sha256:72a94840722b102b37a9c53fba5cdd1be10e042e8fa95524290aff5df8c0db46","observation_id":"c69ffb56-6de6-4e38-82fb-30bff4c0f88a","resolution":{"observed_at":"2026-05-12T01:47:20.089514Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-08-12T22:00:59.825077Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.08135","last_updated":"2025-05-16T15:45:14Z","snapshot_observed_at":"2026-08-12T23:23:37.689372Z","submitted_at":"2024-11-12T19:26:43Z","title":"On the Role of Speech Data in Reducing Toxicity Detection Bias","version":2},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-12T22:00:59.825077Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2411.08135"},"observation_digest":"sha256:1920a06e2978330a730a2d6baced76686d402138c615f064138a326d2311b32b","observation_id":"69b61b3e-36af-435d-9003-97026787bd93","resolution":{"observed_at":"2026-08-12T22:00:59.825077Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-08-12T11:54:52.797320Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.17666","last_updated":"2025-02-20T18:04:45Z","snapshot_observed_at":"2026-08-19T03:26:02.745409Z","submitted_at":"2024-11-26T18:29:11Z","title":"How do Multimodal Foundation Models Encode Text and Speech? An Analysis of Cross-Lingual and Cross-Modal Representations","version":2},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-12T11:54:52.797320Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2411.17666"},"observation_digest":"sha256:5074449310c716cc68c5267c6f207f60b2c7608c7c7496f12bdfefd81666a8b8","observation_id":"da985be5-f93f-4e5d-b102-61e506f2a235","resolution":{"observed_at":"2026-08-12T11:54:52.797320Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-08-12T10:46:23.603547Z","title":"SONAR: Sentence- level multimodal and language-agnostic representations,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.18953","last_updated":"2024-11-28T06:50:13Z","snapshot_observed_at":"2026-08-21T11:14:56.479942Z","submitted_at":"2024-11-28T06:50:13Z","title":"AudioSetCaps: An Enriched Audio-Caption Dataset using Automated Generation Pipeline with Large Audio and Language Models","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-12T10:46:23.603547Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2411.18953"},"observation_digest":"sha256:b1a311c7794785d177d436eafbbaae967040ac147c5e1c43abf91ba848ca8a70","observation_id":"8ed14750-33cc-4f36-a9a6-d3f5d8ebf5de","resolution":{"observed_at":"2026-08-12T10:46:23.603547Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-08-11T18:05:26.518878Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.08268","last_updated":"2025-07-09T17:25:55Z","snapshot_observed_at":"2026-08-20T09:37:36.200203Z","submitted_at":"2024-12-11T10:35:45Z","title":"LCFO: Long Context and Long Form Output Dataset and Benchmarking","version":3},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-11T18:05:26.518878Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2412.08268"},"observation_digest":"sha256:34728c8ee8a3cb9bbc78893bd11b1692f73717fa324c99fa1397892f6dabf72d","observation_id":"c4242122-49b6-4a29-8d87-c3c8e25b946d","resolution":{"observed_at":"2026-08-11T18:05:26.518878Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-08-11T18:02:27.410280Z","title":"Sonar: Sentence-level multimodal and language-agnostic representations, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.08279","last_updated":"2024-12-11T10:52:29Z","snapshot_observed_at":"2026-08-20T00:14:26.771946Z","submitted_at":"2024-12-11T10:52:29Z","title":"Y-NQ: English-Yor\\`ub\\'a Evaluation dataset for Open-Book Reading Comprehension and Text Generation","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-11T18:02:27.410280Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2412.08279"},"observation_digest":"sha256:dd9f3c6ef5c3e38bfa807d69367e4c6c8657c219eec94050d7d4f22de2292bd4","observation_id":"4eff8563-9cc9-47f3-a1c1-85b24bc8dae1","resolution":{"observed_at":"2026-08-11T18:02:27.410280Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-08-11T17:37:01.080871Z","title":"Duquenne, H","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.08821","last_updated":"2024-12-15T21:20:12Z","snapshot_observed_at":"2026-08-16T08:10:37.041600Z","submitted_at":"2024-12-11T23:36:20Z","title":"Large Concept Models: Language Modeling in a Sentence Representation Space","version":2},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-11T17:37:01.080871Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2412.08821"},"observation_digest":"sha256:6cd8a6e2e2846ecf3d9c5ef388575f90f410b27eee735354ee06a9037dc037d7","observation_id":"415e0821-ce4c-49d8-a342-2341d4f5431a","resolution":{"observed_at":"2026-08-11T17:37:01.080871Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-08-08T22:52:31.524775Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.04314","last_updated":"2025-06-13T19:08:13Z","snapshot_observed_at":"2026-08-15T09:57:41.503011Z","submitted_at":"2025-02-06T18:56:37Z","title":"BOUQuET: dataset, Benchmark and Open initiative for Universal Quality Evaluation in Translation","version":2},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-08T22:52:31.524775Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2502.04314"},"observation_digest":"sha256:740c4881353e84907f1b41d11adde6f898508262d935ec53a3cbd5657f56091d","observation_id":"5f9b5bf3-eb2b-4209-a405-23f198b0a5ef","resolution":{"observed_at":"2026-08-08T22:52:31.524775Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-08-15T20:16:40.641890Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.13628","last_updated":"2025-05-19T18:06:45Z","snapshot_observed_at":"2026-08-17T16:36:37.659546Z","submitted_at":"2025-05-19T18:06:45Z","title":"Cross-Lingual Representation Alignment Through Contrastive Image-Caption Tuning","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-15T20:16:40.641890Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2505.13628"},"observation_digest":"sha256:d867fe3f0a985f0f24d2d4889d49ced75dfb31e91159e580547fc28830f75130","observation_id":"33608e34-5679-43f0-bbb3-34a134100ab2","resolution":{"observed_at":"2026-08-15T20:16:40.641890Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-08-07T12:35:26.208489Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24561","last_updated":"2025-05-30T13:16:08Z","snapshot_observed_at":"2026-08-12T15:49:15.982660Z","submitted_at":"2025-05-30T13:16:08Z","title":"Improving Language and Modality Transfer in Translation by Character-level Modeling","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:26.208489Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2505.24561"},"observation_digest":"sha256:9c54e257cb45ec5cf32669f5035ab9cbf944f1943d4e950eee678493bde0218e","observation_id":"12b1c7e8-e56a-4b95-8caf-381fb22e8c50","resolution":{"observed_at":"2026-08-07T12:35:26.208489Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-08-06T12:43:42.077357Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21568","last_updated":"2025-07-31T08:13:22Z","snapshot_observed_at":"2026-08-16T06:56:24.298883Z","submitted_at":"2025-07-29T07:59:20Z","title":"Multi-Hypothesis Distillation of Multilingual Neural Translation Models for Low-Resource Languages","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T12:43:42.077357Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2507.21568"},"observation_digest":"sha256:b3f1f59f99eabc5c5dbfe82cd186e183b87d3e14e63c4ea165e198acafc13a23","observation_id":"3eec4d76-8157-48af-af1a-5a1e1981bb07","resolution":{"observed_at":"2026-08-06T12:43:42.077357Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":"2308.11466","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-07-05T13:21:06.391412Z","title":"Sonar: Sentence-level multimodal and language-agnostic representations","venue":"cs.CL","work_id":"74071e8e-19c2-44e6-9972-e052b130c9d5","year":2023},"citing_paper":{"arxiv_id":"2604.18109","last_updated":"2026-06-17T13:20:24Z","snapshot_observed_at":"2026-08-14T04:36:27.589977Z","submitted_at":"2026-04-20T11:27:14Z","title":"FLiP: Towards understanding and interpreting multimodal multilingual sentence embeddings","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-10T05:23:25.136602Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2604.18109"},"observation_digest":"sha256:82d0001cb67e6c2824c8c4e8e448f11f1f6103ab19d3191cfb9ae3c6a09ee1c8","observation_id":"ecfa0112-d451-4ae1-a56a-4200595acce5","resolution":{"observed_at":"2026-05-10T05:25:55.103342Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":"2308.11466","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-07-05T13:21:06.391412Z","title":"Sonar: Sentence-level multimodal and language-agnostic representations","venue":"cs.CL","work_id":"74071e8e-19c2-44e6-9972-e052b130c9d5","year":2023},"citing_paper":{"arxiv_id":"2604.18109","last_updated":"2026-06-17T13:20:24Z","snapshot_observed_at":"2026-08-14T04:36:27.589977Z","submitted_at":"2026-04-20T11:27:14Z","title":"FLiP: Towards understanding and interpreting multimodal multilingual sentence embeddings","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-07-05T13:12:10.256070Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2604.18109"},"observation_digest":"sha256:98c1778a2c6301cb921044c4b1c1abf013611cce25ba0f29a892f4af10831295","observation_id":"0abda97d-e46b-40cd-9774-2d4be8248c3a","resolution":{"observed_at":"2026-07-05T13:21:06.394295Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":"2308.11466","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-07-05T13:21:06.391412Z","title":"Sonar: Sentence-level multimodal and language-agnostic representations","venue":"cs.CL","work_id":"74071e8e-19c2-44e6-9972-e052b130c9d5","year":2023},"citing_paper":{"arxiv_id":"2604.21786","last_updated":"2026-04-23T15:44:14Z","snapshot_observed_at":"2026-08-13T11:29:52.904567Z","submitted_at":"2026-04-23T15:44:14Z","title":"From Codebooks to VLMs: Evaluating Automated Visual Discourse Analysis for Climate Change on Social Media","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-05-09T22:55:33.247678Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2604.21786"},"observation_digest":"sha256:e3d1eee58f9acf60c60c29668ce7f50e34c9be34bb14d92facb1e016c5b18a73","observation_id":"5306eda5-50dd-40bc-9945-9ff64f952c9d","resolution":{"observed_at":"2026-05-11T14:16:06.886812Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":"2308.11466","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-07-05T13:21:06.391412Z","title":"Sonar: Sentence-level multimodal and language-agnostic representations","venue":"cs.CL","work_id":"74071e8e-19c2-44e6-9972-e052b130c9d5","year":2023},"citing_paper":{"arxiv_id":"2605.05103","last_updated":"2026-08-12T16:10:50Z","snapshot_observed_at":"2026-08-16T13:35:49.790107Z","submitted_at":"2026-05-06T16:38:49Z","title":"Text Corpora as Concept Fields: Black-Box Hallucination and Novelty Measurement","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-08T16:59:10.015928Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2605.05103"},"observation_digest":"sha256:30443d83d9f1596dd4204b5a5203a39d0fe66184c12ba8b4dc598f87ee855920","observation_id":"f2a9c850-824a-4947-a199-0702fba228d5","resolution":{"observed_at":"2026-05-11T17:51:09.974684Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":"2308.11466","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-07-05T13:21:06.391412Z","title":"Sonar: Sentence-level multimodal and language-agnostic representations","venue":"cs.CL","work_id":"74071e8e-19c2-44e6-9972-e052b130c9d5","year":2023},"citing_paper":{"arxiv_id":"2605.05103","last_updated":"2026-08-12T16:10:50Z","snapshot_observed_at":"2026-08-16T13:35:49.790107Z","submitted_at":"2026-05-06T16:38:49Z","title":"Text Corpora as Concept Fields: Black-Box Hallucination and Novelty Measurement","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-12T03:43:44.821785Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2605.05103"},"observation_digest":"sha256:3cbec9cde512e9e0334f93f53b672c825a83accfb6eaa958cf7e90f0f37e85d4","observation_id":"1ad4b7f7-ef8e-4734-b4f7-ce02ca4c6901","resolution":{"observed_at":"2026-05-12T07:06:35.962481Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":"2308.11466","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-07-05T13:21:06.391412Z","title":"Sonar: Sentence-level multimodal and language-agnostic representations","venue":"cs.CL","work_id":"74071e8e-19c2-44e6-9972-e052b130c9d5","year":2023},"citing_paper":{"arxiv_id":"2605.09476","last_updated":"2026-05-10T11:07:24Z","snapshot_observed_at":"2026-07-06T23:21:35.327066Z","submitted_at":"2026-05-10T11:07:24Z","title":"Align and Shine: Building High-Quality Sentence-Aligned Corpora for Multilingual Text Simplification","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-12T03:01:56.609663Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2605.09476"},"observation_digest":"sha256:d5a442fb9374df5b34c4106808d2a087465cd22f838c61be444fa46654ef5e02","observation_id":"9505da44-7dbd-4ddf-ac5e-5090447f6cd6","resolution":{"observed_at":"2026-05-12T07:26:27.728277Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":"2308.11466","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-07-05T13:21:06.391412Z","title":"Sonar: Sentence-level multimodal and language-agnostic representations","venue":"cs.CL","work_id":"74071e8e-19c2-44e6-9972-e052b130c9d5","year":2023},"citing_paper":{"arxiv_id":"2605.19130","last_updated":"2026-05-18T21:30:54Z","snapshot_observed_at":"2026-08-17T20:39:21.641071Z","submitted_at":"2026-05-18T21:30:54Z","title":"EgoBabyVLM: Benchmarking Cross-Modal Learning from Naturalistic Egocentric Video Data","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-05-20T12:12:45.251924Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2605.19130"},"observation_digest":"sha256:d0badd5cf30cbc1393e8b3f9bfe57f88f16ea96589dfba948337fac2ca01af27","observation_id":"10bc5201-126f-4f65-8d14-c4ffeba9468b","resolution":{"observed_at":"2026-05-20T12:13:16.016456Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":"2308.11466","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-07-05T13:21:06.391412Z","title":"Sonar: Sentence-level multimodal and language-agnostic representations","venue":"cs.CL","work_id":"74071e8e-19c2-44e6-9972-e052b130c9d5","year":2023},"citing_paper":{"arxiv_id":"2606.08748","last_updated":"2026-06-07T17:38:15Z","snapshot_observed_at":"2026-08-13T15:51:24.866213Z","submitted_at":"2026-06-07T17:38:15Z","title":"HydraQE: OSU's Submission for the IWSLT 2026 Speech Translation Metrics Shared Task","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-06-27T18:39:30.287695Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2606.08748"},"observation_digest":"sha256:cc28233550da18731727b03992ed1febe8b4d2b911058c24b7bc18c42e882673","observation_id":"b09ac0d8-9dea-4c3e-b99c-9332e83b1aa5","resolution":{"observed_at":"2026-07-02T22:47:26.108305Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":"2308.11466","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-07-05T13:21:06.391412Z","title":"Sonar: Sentence-level multimodal and language-agnostic representations","venue":"cs.CL","work_id":"74071e8e-19c2-44e6-9972-e052b130c9d5","year":2023},"citing_paper":{"arxiv_id":"2606.17967","last_updated":"2026-06-24T09:36:41Z","snapshot_observed_at":"2026-08-07T23:18:02.961979Z","submitted_at":"2026-06-16T14:18:20Z","title":"Learning task-specific subspaces via interventional post-training of speech foundation models","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-06-27T00:55:12.863379Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2606.17967"},"observation_digest":"sha256:a41532f6abbf600eec9baac9d4fe81080ab89353f161a34fb8405a05cf0fd128","observation_id":"87b242ab-7069-4f9b-8b8d-1eb0c253f667","resolution":{"observed_at":"2026-07-03T21:08:58.106748Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":"2308.11466","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-07-05T13:21:06.391412Z","title":"Sonar: Sentence-level multimodal and language-agnostic representations","venue":"cs.CL","work_id":"74071e8e-19c2-44e6-9972-e052b130c9d5","year":2023},"citing_paper":{"arxiv_id":"2606.27378","last_updated":"2026-05-07T04:04:25Z","snapshot_observed_at":"2026-08-13T09:55:08.244211Z","submitted_at":"2026-05-07T04:04:25Z","title":"Formalizing Latent Thoughts: Four Axioms of Thought Representation in LLMs","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-06-30T23:50:48.584391Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2606.27378"},"observation_digest":"sha256:9beaf21623ade34ce18b8d1d1334c5b923045ad5ef41fcdfd4777311d50c037e","observation_id":"7f3f7b11-06f2-40c5-81cf-425bbc638f5c","resolution":{"observed_at":"2026-07-01T13:15:45.678850Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":"2308.11466","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-07-05T13:21:06.391412Z","title":"Sonar: Sentence-level multimodal and language-agnostic representations","venue":"cs.CL","work_id":"74071e8e-19c2-44e6-9972-e052b130c9d5","year":2023},"citing_paper":{"arxiv_id":"2606.31411","last_updated":"2026-06-30T09:36:55Z","snapshot_observed_at":"2026-08-19T06:10:40.837525Z","submitted_at":"2026-06-30T09:36:55Z","title":"Linguistic Bias Mitigation for Spoofing Detection via Gradient Reversal and A Variational Information Bottleneck","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-07-01T05:56:24.699192Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2606.31411"},"observation_digest":"sha256:b65395607340f7be898182dbf0400bb38b35ca36df40d8e9aded04b4dbe682ae","observation_id":"b51b4af1-a270-430c-b481-64b0cc261b42","resolution":{"observed_at":"2026-07-01T10:05:40.497056Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.11466","snapshot_observed_at":"2026-08-16T00:03:40.478721Z","title":"arXiv preprint arXiv:2308.11466 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.12756","last_updated":"2026-08-13T03:08:40Z","snapshot_observed_at":"2026-08-21T15:21:44.188122Z","submitted_at":"2026-08-13T03:08:40Z","title":"ReconSpan: Reconstruction-Guided Adaptive Latent Tokenization","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-16T00:03:40.478721Z"},"links":{"cited_paper":"/paper/2308.11466","citing_paper":"/paper/2608.12756"},"observation_digest":"sha256:835369628a9ef7e1a44b1f90bd602d27903ce4734203419b9316584c6c786719","observation_id":"942d2f37-97da-451b-b91c-9e5c4dbb5cc6","resolution":{"observed_at":"2026-08-16T00:03:40.478721Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2308.11466/citation-record","integrity":"/paper/2308.11466/integrity","json":"/paper/2308.11466/citation-record.json","paper":"/paper/2308.11466"},"outbound":[],"paper":{"arxiv_id":"2308.11466","last_updated":"2023-08-23T10:46:16Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-21T07:59:09.675264Z","submitted_at":"2023-08-22T14:25:15Z","title":"SONAR: Sentence-Level Multimodal and Language-Agnostic Representations"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"thesis":"As of 23 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 23 inbound Pith citation observations for arXiv:2308.11466."}