{"as_of":"2026-08-19T06:38:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:77c21d580a3bcdc9a6b94b590f5e046aa673a364ecffd38834f71182a2e65d54","coverage":[{"denominator":37,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":37,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T05:00:34.142649Z","state":"measured"},{"denominator":38,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":38,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-19T06:32:44.657259+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-18T13:02:08.000814Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"cited_work":{"arxiv_id":"2504.21681","doi":"10.48550/arxiv.2504.21681","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.21681","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":null,"venue":"arXiv (Cornell University)","work_id":"eb51fef3-f171-4bcc-835c-ee4247071d72","year":2025},"citing_paper":{"arxiv_id":"2509.22123","last_updated":"2026-05-13T12:28:55Z","snapshot_observed_at":"2026-08-08T13:18:54.000166Z","submitted_at":"2025-09-26T09:46:13Z","title":"Multilingual Vision-Language Models, A Survey","version":2},"reference_index":94,"source":"pdf_text","source_observed_at":"2026-05-18T13:02:08.000814Z"},"links":{"cited_paper":"/paper/2504.21681","citing_paper":"/paper/2509.22123"},"observation_digest":"sha256:8dd1e33d8ad853afeeaa0938e52f0fd1c4ddf2fea88e4c2c24b403bdf4595cde","observation_id":"1f2d0755-0e5e-4dbf-84c4-923de3705c22","resolution":{"observed_at":"2026-05-18T13:02:37.069380Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2504.21681/citation-record","integrity":"/paper/2504.21681/integrity","json":"/paper/2504.21681/citation-record.json","paper":"/paper/2504.21681"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2201.11732","last_updated":"2022-07-17T13:01:43Z","snapshot_observed_at":"2026-08-18T21:08:21.787380Z","submitted_at":"2022-01-27T18:53:22Z","title":"IGLUE: A Benchmark for Transfer Learning across Modalities, Tasks, and Languages","version":2},"cited_work":{"arxiv_id":"2201.11732","doi":null,"metadata_source":"pith","pith_arxiv_id":"2201.11732","snapshot_observed_at":"2026-08-16T05:00:34.406694Z","title":"IGLUE: A Benchmark for Transfer Learning across Modalities, Tasks, and Languages","venue":"cs.CL","work_id":"ccc7a096-0c48-443d-be89-8dda77c8a90f","year":2022},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:33.999144Z"},"links":{"cited_paper":"/paper/2201.11732","citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:c2a41407e8866a795a249df731eea9b7608c6d4d0964db5c6e873cdad71f9928","observation_id":"8441ec24-8cea-4c63-9168-43a32fa4ebcb","resolution":{"observed_at":"2026-08-16T05:00:34.410989Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.547970Z","title":"In: International Conference on Learning Representations (2020), https://openreview.net/forum?id=r1xCMyBtPS","venue":null,"work_id":"84ad7464-2290-4bf2-98f9-60ac58d85bba","year":2020},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.004484Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:b158222e17dd510be4bec312af3d5276c0b3403af7298c911831fec438661140","observation_id":"4f5fcd50-bc3a-439a-a74e-be04d0245e83","resolution":{"observed_at":"2026-08-16T05:00:34.552826Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.536882Z","title":"In: Proceedings of the Thirteenth Language Resources and Evaluation Conference","venue":null,"work_id":"0b17f04d-0ff6-49b4-8a5c-144a8c6b9bff","year":2022},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.009062Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:1e804b1bce393ff80f8240dd5d4645b0538adc01474903ae854e691e53e068d7","observation_id":"edeedecb-7052-4df3-8a14-60d0b4f52e42","resolution":{"observed_at":"2026-08-16T05:00:34.540790Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.524310Z","title":"In: Computer Vision - ECCV 2020 - 16th European Conference, Glasgow, UK, August 23-28, 2020, Proceedings, Part XXX","venue":null,"work_id":"7a618875-885f-44d5-9d73-254bff606c4b","year":2020},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.013073Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:ee205fc095dc57e50a05fec9f1fed74082431245f989549fb9f7ee3c0c7e49e2","observation_id":"79c13419-db3e-4387-a9d3-94b13ba101eb","resolution":{"observed_at":"2026-08-16T05:00:34.528399Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.021341Z","title":"In: Proceedings of the 58th Annual Meeting of the Association for Computational Linguistics","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.021341Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:26c4f0ee4f78ec43ade2d6a86557567f912b7c2624477abeb3ef84096533447d","observation_id":"cd942499-f8a7-410e-8db1-153ca302960b","resolution":{"observed_at":"2026-08-16T05:00:34.021341Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.511650Z","title":"In: Proceedings of the 2019 Con- ferenceoftheNorthAmericanChapteroftheAssociationforComputationalLinguis- tics: Human Language Technologies, Volume 1 (Long and Short Papers)","venue":null,"work_id":"3596f218-2dfb-48f9-ae45-60b784420e01","year":2019},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.025968Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:6728139575b300804829da17bd5707fbe8bc240f22ca14c1ef32fe0eabc8ebb4","observation_id":"3a369c23-fdce-41ba-9ae8-f8b95d4b4548","resolution":{"observed_at":"2026-08-16T05:00:34.515692Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2101.08231","last_updated":"2021-08-12T03:07:58Z","snapshot_observed_at":"2026-08-18T01:16:06.791694Z","submitted_at":"2021-01-20T17:54:47Z","title":"Word Alignment by Fine-tuning Embeddings on Parallel Corpora","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2101.08231","snapshot_observed_at":"2026-08-16T05:00:34.033740Z","title":"CoRR abs/2101.08231 (2021), https://arxiv.org/abs/2101.08231","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.033740Z"},"links":{"cited_paper":"/paper/2101.08231","citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:c224572e3ab6d4376a25579213d9be66d315f7ded86136a1bdb70f94f81eb395","observation_id":"81257ecf-ec3c-46ab-b817-b8b5b7513c7f","resolution":{"observed_at":"2026-08-16T05:00:34.033740Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.037780Z","title":"In: Proceedings of the 60th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.037780Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:c2750d8e9c3e6fb9d055c5091040d299d2ec74e60ab53534e3c9f37e5a04b466","observation_id":"f695c9da-ed62-48b8-86d4-90133bfa34a8","resolution":{"observed_at":"2026-08-16T05:00:34.037780Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.041596Z","title":"In: Proceedings of the 3rd Workshop on Advances in Language and Vision Research (ALVR)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.041596Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:5187c4f543e958e943d827c6a18236a99b5555e6835c66b49256a7717e3bab55","observation_id":"ad426ea2-e574-4cea-a784-eca4d71becd7","resolution":{"observed_at":"2026-08-16T05:00:34.041596Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.045432Z","title":"In: Proceedings of the 61st Annual Meeting of the Association for Computational Linguistics (Vol- ume 2: Short Papers)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.045432Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:a8357f1bbbf0a8a069ee955a054b8c1030e731b0df7ca3858010f0ee2d750e64","observation_id":"244a1261-c664-4f0a-aca9-50e339e052c8","resolution":{"observed_at":"2026-08-16T05:00:34.045432Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2208.02131","last_updated":"2023-03-14T23:51:53Z","snapshot_observed_at":"2026-08-18T03:39:08.432838Z","submitted_at":"2022-08-03T15:11:01Z","title":"Masked Vision and Language Modeling for Multi-modal Representation Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2208.02131","snapshot_observed_at":"2026-08-16T05:00:34.049294Z","title":"CoRR abs/2208.02131 (2022)","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.049294Z"},"links":{"cited_paper":"/paper/2208.02131","citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:1437a68e9601ceeede51c261968da8e7c58f259240772a8b2f31f502c9d90a2a","observation_id":"63114a85-0dcd-4e80-be63-30d9cd94b3b0","resolution":{"observed_at":"2026-08-16T05:00:34.049294Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.498719Z","title":"In: Advances in Neural Information Processing Systems","venue":null,"work_id":"fb69524a-b4ae-435b-a5aa-70ac0d8ac5d2","year":2021},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.053566Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:a8edc4ed63741f17b752587b1d94d0923a1831042c0f978e7ed1d8c646befa94","observation_id":"9d0e086c-48c7-4c58-af3d-8b9124f50001","resolution":{"observed_at":"2026-08-16T05:00:34.503154Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.485971Z","title":"In: Proceedings of the 2021 Conference on Empirical Methods in Natural Language Processing","venue":null,"work_id":"be18c08b-2db8-44b3-90e6-fd6c6e33a9e2","year":2021},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.057716Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:2f31815bde17d863f9857e64506500c70ea6d207cdd861e590c99d8a91e81969","observation_id":"9b29c8b1-1156-4907-8389-bfb1b4d670bf","resolution":{"observed_at":"2026-08-16T05:00:34.490647Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2109.13238","last_updated":"2021-10-21T18:13:28Z","snapshot_observed_at":"2026-08-18T01:14:33.450421Z","submitted_at":"2021-09-28T16:51:38Z","title":"Visually Grounded Reasoning across Languages and Cultures","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2109.13238","snapshot_observed_at":"2026-08-16T05:00:34.061453Z","title":"CoRRabs/2109.13238 (2021), https://arxiv.org/abs/2109.13238","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.061453Z"},"links":{"cited_paper":"/paper/2109.13238","citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:63fb7f0f6a13b779b31d3dcc3ca5328a6cf5e9d01fad0b8c4cb1b2bec3604654","observation_id":"72544182-6f58-4d89-933f-f54c76462e63","resolution":{"observed_at":"2026-08-16T05:00:34.061453Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1907.11692","last_updated":"2019-07-26T17:48:29Z","snapshot_observed_at":"2026-08-16T14:33:50.657682Z","submitted_at":"2019-07-26T17:48:29Z","title":"RoBERTa: A Robustly Optimized BERT Pretraining Approach","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1907.11692","snapshot_observed_at":"2026-08-16T05:00:34.065641Z","title":"CoRR abs/1907.11692 (2019), http://arxiv.org/abs/1907.11692","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.065641Z"},"links":{"cited_paper":"/paper/1907.11692","citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:b74814f5fa16095065903db212e3b5f4337d279622e04ac4ab77a70dd98d72c2","observation_id":"acbc61e3-7833-4eb3-bc24-f17ea275bfff","resolution":{"observed_at":"2026-08-16T05:00:34.065641Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2006.02635","last_updated":"2021-04-01T03:43:53Z","snapshot_observed_at":"2026-08-04T07:21:11.970404Z","submitted_at":"2020-06-04T03:54:29Z","title":"M3P: Learning Universal Representations via Multitask Multilingual Multimodal Pre-training","version":4},"cited_work":{"arxiv_id":"2006.02635","doi":null,"metadata_source":"pith","pith_arxiv_id":"2006.02635","snapshot_observed_at":"2026-08-16T05:00:34.357289Z","title":"M3P: Learning Universal Representations via Multitask Multilingual Multimodal Pre-training","venue":"cs.CL","work_id":"06c468a8-f7eb-42c5-a349-535080029d8f","year":2020},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.069540Z"},"links":{"cited_paper":"/paper/2006.02635","citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:73e2bee41ec262f508dc7455d76ffebf7896dc509d37e4680bff72882bd36efb","observation_id":"6a3c5e3f-f796-4bc9-8749-2a29d451cd1e","resolution":{"observed_at":"2026-08-16T05:00:34.362062Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.474132Z","title":"The Prague Bulletin of Mathematical Linguistics106(1), 125 (2016)","venue":null,"work_id":"b5045964-9177-4e50-bbd2-cb2b02b9cfaf","year":2016},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.073489Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:3754bf0e0a6d178342ecbb00baf0e156c8681923f44548dd0fa693a8283c68f1","observation_id":"67d8550f-8e00-413e-8745-94eb5af6e402","resolution":{"observed_at":"2026-08-16T05:00:34.478366Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.077173Z","title":"In: Proceedings of the 2022 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.077173Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:1a113ce594ce4c538f774352ba056da0b1854203e876ffffc14f83f4317d4ff6","observation_id":"71a46e20-2e7d-4ee3-b3f4-45dcff521613","resolution":{"observed_at":"2026-08-16T05:00:34.077173Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.080669Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.080669Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:cf5cf2beede00d4d5f8d614644a013a063267028c44b502bf1fd356a2467b263","observation_id":"0c8d60a0-efc7-4f75-9b51-724cb9c6aeeb","resolution":{"observed_at":"2026-08-16T05:00:34.080669Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.462380Z","title":"In: Proceedings of the 38th International Conference on Machine Learning","venue":null,"work_id":"0928400c-ebf5-4049-9b06-d17499a43718","year":2021},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.084854Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:7beb364aae17ef7f3d03a80d052e377d85ca02068257628e0b06e3121092f18a","observation_id":"913d5d1d-9e5d-4708-8518-0043ca5240d5","resolution":{"observed_at":"2026-08-16T05:00:34.466214Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.088709Z","title":"In: International conference on machine learning","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.088709Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:2253716b3754a550ad381b605ec551ef29281267a937ade30c872577f231c5ed","observation_id":"a97f8d4f-df17-45be-b41a-7433b502597a","resolution":{"observed_at":"2026-08-16T05:00:34.088709Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.092143Z","title":"In: Proceedings of the 2020 Conference on Empirical Investigating Cross-Lingual Transfer for VL Encoders 9 Methods in Natural Language Processing (EMNLP)","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.092143Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:4e03a7c7328a7b15f58c24e751e1a3f9ee6408d3b77d048dd805e443544b6584","observation_id":"1d3aaa83-8275-47fa-b991-9ba8c0c27483","resolution":{"observed_at":"2026-08-16T05:00:34.092143Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2024.findings-emnlp.250","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.214752Z","title":"In: Findings of the Association for Computational Linguis- tics: EMNLP 2024","venue":null,"work_id":"19e4973d-918e-45df-ad57-4f09e5afb1e1","year":2024},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.095711Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:9200ac070616346d5a81499de189c03d666a5b1d0fece863a726450c2cc29adf","observation_id":"ce5e7640-92c0-4ba0-b0a5-07683e384dbf","resolution":{"observed_at":"2026-08-16T05:00:34.220920Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1811.00491","last_updated":"2019-07-21T05:26:36Z","snapshot_observed_at":"2026-08-14T18:05:37.795597Z","submitted_at":"2018-11-01T16:47:44Z","title":"A Corpus for Reasoning About Natural Language Grounded in Photographs","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1811.00491","snapshot_observed_at":"2026-08-16T05:00:34.099641Z","title":"CoRR abs/1811.00491 (2018), http://arxiv.org/abs/1811.00491","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.099641Z"},"links":{"cited_paper":"/paper/1811.00491","citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:737fd37b15084f7aeb1633daf81bce76dda2d1bfbe0e15a902294bb1187eec21","observation_id":"18090eb9-c4bc-4998-b6d6-ceb06465e0c6","resolution":{"observed_at":"2026-08-16T05:00:34.099641Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.103417Z","title":"In: Proceedings of the 2019 Conference on Empirical Methods in Natural Language Processing and the 9th International Joint Conference on Natural Language Processing (EMNLP-IJCNLP)","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.103417Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:56fb9b029d78c96999220a985b1189e384d4a5f32e5b5319bdf01dcd58156dc2","observation_id":"c9d742db-bfbc-4b6e-939b-6357d1ca6b5e","resolution":{"observed_at":"2026-08-16T05:00:34.103417Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.444396Z","title":"In: Proceedings of the Eighth International Conference on Language Resources and Evaluation (LREC’12)","venue":null,"work_id":"0e90bbd9-77b9-4674-a343-1a793e315926","year":2012},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.107302Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:a537e8ac8b3cf2ef3eb0e0e807d3f635c804e9d1280bb0e34a34928f4ac319c8","observation_id":"43037746-35d9-41a0-a7b4-2511423ab79b","resolution":{"observed_at":"2026-08-16T05:00:34.448449Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.110762Z","title":"In: Proceed- ings of the 2022 Conference on Empirical Methods in Natural Language Pro- cessing","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.110762Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:8a868ebf1502951ceaf452e8fe341374738e30770dd45c82f17edf81be9b06a9","observation_id":"785356bb-5ccd-4b85-896f-058a8af1210b","resolution":{"observed_at":"2026-08-16T05:00:34.110762Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.115376Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.115376Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:4a99ab46294b5b18af3d0f9563ef1dc38e5360e1d0fd1bd57567eeea6fc2e8c3","observation_id":"24ca2ce3-230a-4c9f-a8ce-8651a50880f5","resolution":{"observed_at":"2026-08-16T05:00:34.115376Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.431757Z","title":"In: Thirty-Seventh AAAI Conference on Artificial Intelligence, AAAI","venue":null,"work_id":"8753017c-726b-42eb-8bbd-9203e62471eb","year":null},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.119113Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:934302c63a56d38afb10d899e22aed8dbf352f6b9b1aa46d9844b9ec45a233c1","observation_id":"75a30008-1026-4aee-9786-7d67430b1e27","resolution":{"observed_at":"2026-08-16T05:00:34.435956Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.12402","last_updated":"2023-07-30T13:20:13Z","snapshot_observed_at":"2026-08-16T16:14:08.080032Z","submitted_at":"2022-11-22T16:48:01Z","title":"X$^2$-VLM: All-In-One Pre-trained Model For Vision-Language Tasks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.12402","snapshot_observed_at":"2026-08-16T05:00:34.127001Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.127001Z"},"links":{"cited_paper":"/paper/2211.12402","citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:223b0f648b2ed8d20322e029c1fe8c4570e4781d7700b3fe8947cdb7b8a47ed2","observation_id":"005fd9fe-3143-454f-aafd-b152152a3729","resolution":{"observed_at":"2026-08-16T05:00:34.127001Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.130847Z","title":"In: Proceedings of the 61st Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.130847Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:3ea5f19277f5b12e91488017cf2b05a2744ff8a06a0b42e5024fcd32b5d9ef29","observation_id":"91905c27-c92f-49cf-87ec-4878bd59817a","resolution":{"observed_at":"2026-08-16T05:00:34.130847Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.135097Z","title":"In: Proceedings of the 58th Annual Meeting of the Association for Computational Linguistics","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.135097Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:1565366d4ff835087d1efa9763e7217e465f2feae3b43c51163e83672088d2f5","observation_id":"c1b76ca0-d7cb-44db-9aae-c79c112734cf","resolution":{"observed_at":"2026-08-16T05:00:34.135097Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1904.09675","last_updated":"2020-02-24T18:59:28Z","snapshot_observed_at":"2026-07-29T15:42:51.774083Z","submitted_at":"2019-04-21T23:08:53Z","title":"BERTScore: Evaluating Text Generation with BERT","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1904.09675","snapshot_observed_at":"2026-08-16T05:00:34.138823Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.138823Z"},"links":{"cited_paper":"/paper/1904.09675","citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:b516c9e8117bdddac43f970d4c85d48e83e2cf15d742b0c395a6b38a7914a3c4","observation_id":"5c7ce4b6-bd54-4f3d-b25f-14f9f94b6021","resolution":{"observed_at":"2026-08-16T05:00:34.138823Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2104.00332","last_updated":"2021-04-01T08:30:53Z","snapshot_observed_at":"2026-08-16T18:34:33.750788Z","submitted_at":"2021-04-01T08:30:53Z","title":"UC2: Universal Cross-lingual Cross-modal Vision-and-Language Pre-training","version":1},"cited_work":{"arxiv_id":"2104.00332","doi":null,"metadata_source":"pith","pith_arxiv_id":"2104.00332","snapshot_observed_at":"2026-08-16T05:00:34.309724Z","title":"UC2: Universal Cross-lingual Cross-modal Vision-and-Language Pre-training","venue":"cs.CV","work_id":"0e433ed4-659b-4a0c-b7d9-ce8843b54ae6","year":2021},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.142649Z"},"links":{"cited_paper":"/paper/2104.00332","citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:4eae653f001c2019e28900a8020eaa1336a40ab7011e8bc7e89bf8ae581d5eed","observation_id":"c475e7a2-ba6b-45b2-87e7-eead080a3431","resolution":{"observed_at":"2026-08-16T05:00:34.314591Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.017007Z","title":"https://doi.org/10.1007/978-3-030-58577-8\\_7","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":120,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.017007Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:9743ed7a43cf52e88cc349c05bf6a44ef01064c34f4147f96eb7a935a670d114","observation_id":"9c78d0a9-123a-44e1-aa71-fed418e4f624","resolution":{"observed_at":"2026-08-16T05:00:34.017007Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.419413Z","title":"10637–10647","venue":null,"work_id":"1d37de35-37ed-4098-a705-aca6e2b75346","year":2023},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.123316Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:2062c20ffa7988012cc0ccccd222f25a2c7021d6a6e45fe8a6f1273986db889b","observation_id":"bbd8c83c-b60a-46c4-bc25-57370c6ddcd9","resolution":{"observed_at":"2026-08-16T05:00:34.423497Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:00:34.029599Z","title":"https://doi.org/10.18653/v1/N19- 1423","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders","version":2},"reference_index":4186,"source":"pdf_text","source_observed_at":"2026-08-16T05:00:34.029599Z"},"links":{"citing_paper":"/paper/2504.21681"},"observation_digest":"sha256:55a770878af8da8b2a8ebc4802fc93787875591ef83eab213655d2c813f03a31","observation_id":"fe0205ea-b09f-4686-8210-e4d41020fc07","resolution":{"observed_at":"2026-08-16T05:00:34.029599Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2504.21681","last_updated":"2025-08-15T08:17:15Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-18T21:07:54.724698Z","submitted_at":"2025-04-30T14:19:15Z","title":"Investigating the Effect of Parallel Data in the Cross-Lingual Transfer for Vision-Language Encoders"},"reference_resolution":{"displayed":37,"state_counts":{"malformed_identifier":1,"metadata_mismatch":3,"parse_uncertain":0,"unresolved":21,"verified_exact":1,"verified_fuzzy":11},"total_outbound_references":37},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"thesis":"As of 19 August 2026, this Paper Citation Record lists 37 of 37 outbound references and 1 inbound Pith citation observation for arXiv:2504.21681."}