{"as_of":"2026-08-08T06:59:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:0f51578c95b85342fc6c55b4c721468b0142c46cc2e5180021c1ceb63a9016ef","coverage":[{"denominator":35,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":35,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T11:37:06.635382Z","state":"measured"},{"denominator":35,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":35,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.01808/citation-record","integrity":"/paper/2506.01808/integrity","json":"/paper/2506.01808/citation-record.json","paper":"/paper/2506.01808"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:07.071330Z","title":"McCrae, Salima Mdhaffar, Yasmin Moslem, Kenton Murray, Satoshi Nakamura, Matteo Negri, Jan Niehues, Atul Kr","venue":null,"work_id":"17071f52-35dc-467d-b212-6c7298eeaaf0","year":2025},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.499009Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:3ececa0a6dd5d18dc4ab6348a0749f62bda37baeef390bfe3e861293fe32a048","observation_id":"7c3fec06-d499-4612-b211-78e096641959","resolution":{"observed_at":"2026-08-07T11:37:07.074974Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-07T11:37:06.503289Z","title":"Gpt-4 technical report.arXiv preprint arXiv:2303.08774, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.503289Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:e2309fbdbfee7d8414c9245074aabf864124aef37e3eeb8c677238b6a8ab225d","observation_id":"69e96998-aabf-4311-952d-110a99aacc46","resolution":{"observed_at":"2026-08-07T11:37:06.503289Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:07.060591Z","title":null,"venue":null,"work_id":"6852bd0f-fa03-4991-aa47-d484fafbdc01","year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.506761Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:7169c2a4f488a77c2011e580bfc8df6dfb0974bd1727bac3e017c38658ea2ad3","observation_id":"713f3730-07bf-4cb9-b4ff-ef6cf8a553d3","resolution":{"observed_at":"2026-08-07T11:37:07.063918Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.510426Z","title":"From tower to spire: Adding the speech modality to a text-only llm.arXiv preprint arXiv:2503.10620, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.510426Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:c70c31b90dd39508805594944d7ff8a42377cc57c858578222c7811c53e37b0e","observation_id":"a48b94cd-a300-44c4-a7d9-51d03d7d982d","resolution":{"observed_at":"2026-08-07T11:37:06.510426Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11596","last_updated":"2023-10-25T03:52:07Z","snapshot_observed_at":"2026-07-06T16:09:09.018523Z","submitted_at":"2023-08-22T17:44:18Z","title":"SeamlessM4T: Massively Multilingual & Multimodal Machine Translation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.11596","snapshot_observed_at":"2026-08-07T11:37:06.514410Z","title":"Seamlessm4t: Massively multilingual & multimodal machine translation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.514410Z"},"links":{"cited_paper":"/paper/2308.11596","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:ad8f421123b4004d3c48650d774d7656ebbe4b1a0d9e03e450de2f738b0f65d8","observation_id":"29dbeb19-39f9-4ba6-ace1-20d89998f4fe","resolution":{"observed_at":"2026-08-07T11:37:06.514410Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.00037","last_updated":"2024-10-02T09:11:45Z","snapshot_observed_at":"2026-07-30T10:21:14.474746Z","submitted_at":"2024-09-17T17:55:39Z","title":"Moshi: a speech-text foundation model for real-time dialogue","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.00037","snapshot_observed_at":"2026-08-07T11:37:06.518264Z","title":"Moshi: a speech-text foun- dation model for real-time dialogue","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.518264Z"},"links":{"cited_paper":"/paper/2410.00037","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:2b51b069bb70c7d4fc724071c1334c91042d4f60596f69d85e3379b6eb7bce1c","observation_id":"9e235be1-43f9-4d9c-a77c-18d8c2621ad0","resolution":{"observed_at":"2026-08-07T11:37:06.518264Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.03378","last_updated":"2023-03-06T18:58:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-03-06T18:58:06Z","title":"PaLM-E: An Embodied Multimodal Language Model","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.03378","snapshot_observed_at":"2026-08-07T11:37:06.522198Z","title":"Palm- e: An embodied multimodal language model.arXiv preprint arXiv:2303.03378, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.522198Z"},"links":{"cited_paper":"/paper/2303.03378","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:59baadd54a0b36a323e63e71d86ec129a666d67cf8ad8c5ff552c1be1d33acf4","observation_id":"6c8ffb9c-fd9f-4064-bc0c-1f614b804218","resolution":{"observed_at":"2026-08-07T11:37:06.522198Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-07-06T18:55:11.576666Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-07T11:37:06.526560Z","title":"The llama 3 herd of models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.526560Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:0645f1b6df6548ce4b944281f4c9c0d70abb405307666ae13165ef5735017de6","observation_id":"78473648-fb55-408c-becd-bcd5199fc2f3","resolution":{"observed_at":"2026-08-07T11:37:06.526560Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:07.049298Z","title":"LoRA: Low-rank adaptation of large language 7 NAVER LABS Europe Submission to the Instruction-following Track models","venue":null,"work_id":"9ad09866-a0dd-48d5-9529-596e889cd8a5","year":2022},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.535763Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:3a425c1be84536dc48d0fad2115fbbc83fdbc41592de653eea5cd41d3130f8d4","observation_id":"0527e5fb-6d4e-4c48-a7bc-12898a03eb3a","resolution":{"observed_at":"2026-08-07T11:37:07.053020Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.00656","last_updated":"2024-09-21T15:27:30Z","snapshot_observed_at":"2026-08-08T03:00:53.438875Z","submitted_at":"2024-03-31T12:01:32Z","title":"WavLLM: Towards Robust and Adaptive Speech Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.00656","snapshot_observed_at":"2026-08-07T11:37:06.539254Z","title":"Wavllm: Towards ro- bust and adaptive speech large language model.arXiv preprint arXiv:2404.00656, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.539254Z"},"links":{"cited_paper":"/paper/2404.00656","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:1758d78b627b337ce858a7929a1e9eadb0f2a5a62d64fe6020326f884a315686","observation_id":"e44347e4-7050-40ad-adc0-08546cffa511","resolution":{"observed_at":"2026-08-07T11:37:06.539254Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.543185Z","title":"Audiogpt: Understand- ing and generating speech, music, sound, and talking head","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.543185Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:beab05abd55e69c31aeb98ec7a96aba622f6d2ea5787f5352a32b28d3d3be290","observation_id":"8eda5c9f-a019-4a83-abf5-9999fef582d8","resolution":{"observed_at":"2026-08-07T11:37:06.543185Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.546654Z","title":"Iranzo-Sánchez, J","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.546654Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:698e9686cff80bb119e0b157378a0bf60aab8e1168e4524fdebf165590162315","observation_id":"268ac7ca-3cac-4d7d-aa18-08e6c5e59813","resolution":{"observed_at":"2026-08-07T11:37:06.546654Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04088","last_updated":"2024-01-08T18:47:34Z","snapshot_observed_at":"2026-08-08T06:16:25.839566Z","submitted_at":"2024-01-08T18:47:34Z","title":"Mixtral of Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.04088","snapshot_observed_at":"2026-08-07T11:37:06.550805Z","title":"Mixtral of experts.arXiv preprint arXiv:2401.04088, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.550805Z"},"links":{"cited_paper":"/paper/2401.04088","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:f5825acbaa4a7761a31049c48f296585e8a1586b93082327e811f8f3d185faf7","observation_id":"a2d57ed4-61f6-49a2-b936-66c799a19564","resolution":{"observed_at":"2026-08-07T11:37:06.550805Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.02246","last_updated":"2024-05-03T17:00:00Z","snapshot_observed_at":"2026-07-06T18:09:29.876755Z","submitted_at":"2024-05-03T17:00:00Z","title":"What matters when building vision-language models?","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.02246","snapshot_observed_at":"2026-08-07T11:37:06.554978Z","title":"What matters when building vision- language models?, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.554978Z"},"links":{"cited_paper":"/paper/2405.02246","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:2c78285520d15ae1786a29f673deb4181a4831559d4cfb4407d020877a8b35ab","observation_id":"7b786262-dace-4649-815d-49deec82777d","resolution":{"observed_at":"2026-08-07T11:37:06.554978Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:07.025794Z","title":"Spoken squad: A study of mitigating the impact ofspeechrecognitionerrorsonlisteningcomprehension","venue":null,"work_id":"64e5c969-bffb-4b44-b62b-f024a907f210","year":2018},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.558652Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:70b9685196d2280022a67014f052e1fdb41da06546bcfcee1463018852d9f2cd","observation_id":"90ff5998-a95c-4657-b95c-5e87a9c4e4a7","resolution":{"observed_at":"2026-08-07T11:37:07.029748Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:07.014292Z","title":"ROUGE: A package for automatic eval- uation of summaries","venue":null,"work_id":"20b7cf8d-348c-4f06-9556-b8240c344b73","year":2004},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.562104Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:e7524eebc73a52d685dd4db1cfea1f04680bc5c977841280ecf8ece1d8245c34","observation_id":"67ea7aa6-3736-46d4-9236-bc298fdd2d9d","resolution":{"observed_at":"2026-08-07T11:37:07.018393Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.08485","last_updated":"2023-12-11T17:46:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-04-17T17:59:25Z","title":"Visual Instruction Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.08485","snapshot_observed_at":"2026-08-07T11:37:06.566706Z","title":"Visual Instruction Tuning (LLaVA), 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.566706Z"},"links":{"cited_paper":"/paper/2304.08485","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:922acbc26bf1fdff653f27979f3c48821d86ab00315b3d43592c2acb6ef59318","observation_id":"529dc7ec-0d63-4f5c-b0c1-57b19886e295","resolution":{"observed_at":"2026-08-07T11:37:06.566706Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:07.002983Z","title":"Guerreiro, Ricardo Rei, Duarte M","venue":null,"work_id":"dee13529-43a1-4b82-b511-0bb75bb7613b","year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.571065Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:39581ed52935b8d6e1f2d023c8630be4a736abe818227590f7983b71f8eed053","observation_id":"48b3f722-9daa-4496-a77a-3203c0b1c53c","resolution":{"observed_at":"2026-08-07T11:37:07.007019Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.16235","last_updated":"2024-09-24T16:51:36Z","snapshot_observed_at":"2026-07-06T19:21:20.826877Z","submitted_at":"2024-09-24T16:51:36Z","title":"EuroLLM: Multilingual Language Models for Europe","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.16235","snapshot_observed_at":"2026-08-07T11:37:06.575250Z","title":"Eurollm: Multilingual language mod- els for europe.arXiv preprint arXiv:2409.16235, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.575250Z"},"links":{"cited_paper":"/paper/2409.16235","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:3d0fb23775a7643411a51e628eef6eeb004060344ef4e0f196f7e76f610deab3","observation_id":"2920b360-4f54-485c-a034-4d39e88c725b","resolution":{"observed_at":"2026-08-07T11:37:06.575250Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.579223Z","title":"Spirit-lm: Interleaved spoken and written language model.Transactions of the Association for Computational Linguistics, 13:30–52, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.579223Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:51c98004b1d13fe06fca3dd244b42879b22c66ad15a89c21c35661c9a25bf9af","observation_id":"5fdc741b-b046-466f-a913-11c68e3f2221","resolution":{"observed_at":"2026-08-07T11:37:06.579223Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.22577","last_updated":"2025-05-20T10:29:41Z","snapshot_observed_at":"2026-08-07T16:29:48.727258Z","submitted_at":"2025-03-28T16:26:52Z","title":"Breaking Language Barriers in Visual Language Models via Multilingual Textual Regularization","version":2},"cited_work":{"arxiv_id":"2503.22577","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.22577","snapshot_observed_at":"2026-08-07T11:37:06.706367Z","title":"Breaking Language Barriers in Visual Language Models via Multilingual Textual Regularization","venue":"cs.CV","work_id":"c19ef716-546f-4228-9f3e-6c8ed4748b25","year":2025},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.582912Z"},"links":{"cited_paper":"/paper/2503.22577","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:0f5c7b0654f8e61072a6459846a4a4ad33e01a4dc1e778d4db8a122633ae7aab","observation_id":"8517dfc8-5bcd-484f-aed7-718695d290c6","resolution":{"observed_at":"2026-08-07T11:37:06.712079Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.985809Z","title":"A call for clarity in reporting BLEU scores","venue":null,"work_id":"22eafb73-4382-4b01-a374-521068d3beb7","year":null},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.586670Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:e8bdcb5a9ea941df4df6f3e58df2ff62ec1b07fca8a4b915b9838fcdac32e2ca","observation_id":"8f305645-c550-446f-abe4-20f706ecd273","resolution":{"observed_at":"2026-08-07T11:37:06.989451Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.594248Z","title":"Scaling speech technology to 1,000+ languages.Jour- nal of Machine Learning Research, 25(97):1–52, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.594248Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:cb558c8609862e82bfd8cd5df59e9053ed76b9d2196edfa6bbafa6dc08882c3b","observation_id":"8e793ba6-d5c3-4bef-b36d-73488094783e","resolution":{"observed_at":"2026-08-07T11:37:06.594248Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.955202Z","title":"BERGEN: A benchmarking library for retrieval-augmented generation","venue":null,"work_id":"decb62ae-3ad9-46ea-9b4b-2419e4ee1de0","year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.598180Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:7537f5988bfe21a4932ec1916b40d20a94ba73f27cf2635762bbf6aba92f6a9b","observation_id":"28db40f3-5550-4fa2-bf91-c9a8a521bed4","resolution":{"observed_at":"2026-08-07T11:37:06.959433Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.943700Z","title":null,"venue":null,"work_id":"29e7a67f-9cc6-4b04-97b7-8390235da712","year":2022},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.601596Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:8a4cf43c556ba15aa1404c528e6ff581f41abd0f31f8fb98ab1a0e25c812528c","observation_id":"fca08bbe-71e6-472f-954f-5c7bae3edea8","resolution":{"observed_at":"2026-08-07T11:37:06.947205Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.12925","last_updated":"2023-06-22T14:37:54Z","snapshot_observed_at":"2026-08-07T01:18:02.068918Z","submitted_at":"2023-06-22T14:37:54Z","title":"AudioPaLM: A Large Language Model That Can Speak and Listen","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.12925","snapshot_observed_at":"2026-08-07T11:37:06.604936Z","title":"Audiopalm: A large language model that can speak and listen.arXiv preprint arXiv:2306.12925, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.604936Z"},"links":{"cited_paper":"/paper/2306.12925","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:97e7a5638d34e4448b6645fa5fcb5048747f8bfa5200286da3b47af9521b0570","observation_id":"f79628f3-5985-4d3b-a747-e3bf10535588","resolution":{"observed_at":"2026-08-07T11:37:06.604936Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.932681Z","title":"Evaluating multilingual speech translation under realistic condi- tions with resegmentation and terminology","venue":null,"work_id":"2979ad94-6df9-48e1-b264-846c046d57fb","year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.608488Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:27d74e2b7545995d9d2aa5ebcc475d257084d9f144c26eacde08448447bdc538","observation_id":"0b0423d7-3a1f-4ae8-870d-1fdb098dbd99","resolution":{"observed_at":"2026-08-07T11:37:06.936469Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13289","last_updated":"2024-04-08T06:12:52Z","snapshot_observed_at":"2026-08-02T21:51:36.809095Z","submitted_at":"2023-10-20T05:41:57Z","title":"SALMONN: Towards Generic Hearing Abilities for Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13289","snapshot_observed_at":"2026-08-07T11:37:06.612173Z","title":"Salmonn: Towardsgenerichearingabilitiesforlargelan- guage models.arXiv preprint arXiv:2310.13289, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.612173Z"},"links":{"cited_paper":"/paper/2310.13289","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:e44e136b27d455b10e844f65a0d5e00df116e5d3149ebee1c564324ae210db4d","observation_id":"ed556425-ebb2-49de-852d-dd04fbaa7627","resolution":{"observed_at":"2026-08-07T11:37:06.612173Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.11805","last_updated":"2025-05-09T21:04:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-19T02:39:27Z","title":"Gemini: A Family of Highly Capable Multimodal Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.11805","snapshot_observed_at":"2026-08-07T11:37:06.616463Z","title":"Gemini: a family of highly capable multimodal models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.616463Z"},"links":{"cited_paper":"/paper/2312.11805","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:89a3a953d35599289733841a6b8940cdf4c5e8fd9234bf1364e5e3e315e8e848","observation_id":"e494bf00-d5f7-4d07-9158-8e15878564ef","resolution":{"observed_at":"2026-08-07T11:37:06.616463Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.921289Z","title":null,"venue":null,"work_id":"435d6ca1-a044-4c57-85cb-8568dcda180d","year":2025},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.620780Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:37ccccd409e4a48617ca65686b2ea87780f4c03e4eec6d9a855fbe3f891ee1ff","observation_id":"b934a0aa-6c4f-4856-b6ff-6ae52e3985c8","resolution":{"observed_at":"2026-08-07T11:37:06.925827Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.910153Z","title":"torchtune: Py- torch’s finetuning library, 2024","venue":null,"work_id":"b3400eb8-11f3-4929-8f89-59a9179e1ed2","year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.625535Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:ca3b456441ec0de96edf4a71758ea2ea58eba6f969455796acaa35127555005e","observation_id":"59dcf11b-3005-4317-8278-092b9be9ec2e","resolution":{"observed_at":"2026-08-07T11:37:06.914231Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.899057Z","title":"Llama 2: Open foundation and fine-tuned chat models, 2023","venue":null,"work_id":"cda4a8dc-7603-43e4-b243-9aed4e89d58a","year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.629098Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:d1aa14ea1b3d4446baa4e9663f03a0c333c823bef7dc35e461f180848ea4cab9","observation_id":"7205e1a5-7ee9-43a0-ace1-77133abc0661","resolution":{"observed_at":"2026-08-07T11:37:06.903300Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.632229Z","title":"Covost 2: A massively multilingual speech-to-text translation corpus, 2020","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.632229Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:62f23d3d719857fb1ba478a68f5ea8a04e37783c4df8936f58411a2e0fe594c5","observation_id":"40ce9703-3feb-44d9-82d6-6a15240c47ca","resolution":{"observed_at":"2026-08-07T11:37:06.632229Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15115","last_updated":"2025-01-03T02:18:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-19T17:56:09Z","title":"Qwen2.5 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.15115","snapshot_observed_at":"2026-08-07T11:37:06.635382Z","title":"1”, instead of “0","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.635382Z"},"links":{"cited_paper":"/paper/2412.15115","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:81c197f99d71b29a6ecb5772b4c67037484f92e3cc5b0c9ca5b1989b77ea9879","observation_id":"ad24d2c1-314d-4162-b1b1-2e27ab122559","resolution":{"observed_at":"2026-08-07T11:37:06.635382Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.974138Z","title":null,"venue":null,"work_id":"7cf6ac76-0edd-42ee-81a7-790925a30048","year":null},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":2018,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.590475Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:0827015e1d0c61fbfd51d77493d218e77c8b82a97102f3a09713b9c2259ee478","observation_id":"da1a30ef-2e35-40d7-bb0a-2e32bdb6cd22","resolution":{"observed_at":"2026-08-07T11:37:06.978146Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T11:30:46.088140Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track"},"reference_resolution":{"displayed":35,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":23,"verified_exact":1,"verified_fuzzy":10},"total_outbound_references":35},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 35 of 35 outbound references and 0 inbound Pith citation observations for arXiv:2506.01808."}