{"as_of":"2026-08-08T23:52:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:a92b8ea188b61a49bcfe4997ba72bd06dbf14d4c653635a45a02ebadba18fcef","coverage":[{"denominator":105,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":100,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T15:19:26.982116Z","state":"measured"},{"denominator":100,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":100,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2505.15628/citation-record","integrity":"/paper/2505.15628/integrity","json":"/paper/2505.15628/citation-record.json","paper":"/paper/2505.15628"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:19.013478Z","title":"http://www.gphoto.com/","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.013478Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:ddf085cf24479a02659df5121bf459b89b114763c4ff610ea57b7c7dbc5360bf","observation_id":"14dba2fa-1787-4da8-aa4e-89319a0530c0","resolution":{"observed_at":"2026-08-07T15:19:19.013478Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:19.082195Z","title":"On sensor bias in experimental methods for comparing interest-point, saliency, and recognition algorithms","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.082195Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:1df8ef3e75fef39a99da61409e25644f25b62391b1e4001a725a445a600f2710","observation_id":"47111471-79b0-4b72-8a76-017b58a117ed","resolution":{"observed_at":"2026-08-07T15:19:19.082195Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1805.12177","last_updated":"2019-12-31T13:40:12Z","snapshot_observed_at":"2026-07-06T06:42:08.107734Z","submitted_at":"2018-05-30T18:56:33Z","title":"Why do deep convolutional networks generalize so poorly to small image transformations?","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1805.12177","snapshot_observed_at":"2026-08-07T15:19:19.139371Z","title":"Why do deep convolutional networks generalize so poorly to small image transformations? arXiv:1805.12177, 2018","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.139371Z"},"links":{"cited_paper":"/paper/1805.12177","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:030a999165a5ff5595bd98c49096713bc6a8de931cc5cdf78bdaaec51fae1d48","observation_id":"7cc116aa-85c7-47b1-a9d7-e2e84425f1c8","resolution":{"observed_at":"2026-08-07T15:19:19.139371Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:19.210665Z","title":"Unexplored faces of robustness and out-of-distribution: Covariate shifts in environment and sensor domains","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.210665Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:80da28a90b22171e4965bfcae3e34e0536a86dbd54fcd07b3453bae2f29cfb8a","observation_id":"1db757c1-dfa4-40ea-9d34-5e688d723502","resolution":{"observed_at":"2026-08-07T15:19:19.210665Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:19.286738Z","title":"ObjectNet: A large-scale bias-controlled dataset for pushing the limits of object recognition models","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.286738Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:ccd0eb4e923e3b212f0cd68e79366e592d84b7d9bbb3626addcd05bb624eb7a0","observation_id":"25d7dde1-16fe-45a5-8e4b-37a4113d9712","resolution":{"observed_at":"2026-08-07T15:19:19.286738Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2210.14142","last_updated":"2022-11-17T16:20:21Z","snapshot_observed_at":"2026-07-06T14:10:16.603292Z","submitted_at":"2022-10-25T16:42:03Z","title":"From colouring-in to pointillism: revisiting semantic segmentation supervision","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.14142","snapshot_observed_at":"2026-08-07T15:19:19.341802Z","title":"From colouring-in to pointillism: revisiting semantic segmentation supervision","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.341802Z"},"links":{"cited_paper":"/paper/2210.14142","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:de97ffa2a9f83d4f5eb8b9c481504b2ff466e68e843fa00eb3c1d4f0f340a5b9","observation_id":"c94fac47-20b7-4802-819e-d768b9494457","resolution":{"observed_at":"2026-08-07T15:19:19.341802Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2006.07159","last_updated":"2020-06-12T13:17:25Z","snapshot_observed_at":"2026-08-07T08:01:08.867333Z","submitted_at":"2020-06-12T13:17:25Z","title":"Are we done with ImageNet?","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2006.07159","snapshot_observed_at":"2026-08-07T15:19:19.411111Z","title":"Are we done with ImageNet? arXiv:2006.07159, 2020","venue":null,"work_id":null,"year":2006},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.411111Z"},"links":{"cited_paper":"/paper/2006.07159","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:584a3825de6243e93a24b321e4e2412188f03bb46bdc8d63349ba993f7bb12c2","observation_id":"dcebfebb-d997-4b6a-8655-b4a7189c2013","resolution":{"observed_at":"2026-08-07T15:19:19.411111Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.07726","last_updated":"2024-10-10T17:28:23Z","snapshot_observed_at":"2026-08-08T07:16:45.596308Z","submitted_at":"2024-07-10T14:57:46Z","title":"PaliGemma: A versatile 3B VLM for transfer","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.07726","snapshot_observed_at":"2026-08-07T15:19:19.468698Z","title":"PaliGemma: A versatile 3B VLM for transfer","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.468698Z"},"links":{"cited_paper":"/paper/2407.07726","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:26ef01f7dbfcef429acd6034203e2b7025ae099bdc82a10fa01350e9c4539d1e","observation_id":"09c8659f-e152-4432-8cd8-4d52f52eba35","resolution":{"observed_at":"2026-08-07T15:19:19.468698Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:19.533564Z","title":"A digital image processing pipeline for modelling of realistic noise in synthetic images","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.533564Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:0be850ac4cb2b99c28fe74a843920007fa80d268252f73e01f9dc2cf9f474e4b","observation_id":"cfd36915-25c0-475c-8d4f-5c18e72bd469","resolution":{"observed_at":"2026-08-07T15:19:19.533564Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2004.10934","last_updated":"2020-04-23T02:10:02Z","snapshot_observed_at":"2026-07-06T09:14:32.318388Z","submitted_at":"2020-04-23T02:10:02Z","title":"YOLOv4: Optimal Speed and Accuracy of Object Detection","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2004.10934","snapshot_observed_at":"2026-08-07T15:19:19.604219Z","title":"YOLOv4: Optimal speed and accuracy of object detection","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.604219Z"},"links":{"cited_paper":"/paper/2004.10934","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:9b5151173186ad63fbebfc977bf7d212f0586ffaa8032a24c5340cf4220f637d","observation_id":"87c011a9-d149-4871-8740-a6bcb9ab40de","resolution":{"observed_at":"2026-08-07T15:19:19.604219Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:19.658188Z","title":"COYO-700M: Image-text pair dataset","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.658188Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:bac5a5102eee3ec114979a61c9e072e6b8d33f5214c7dd594fdcda403a4bf38d","observation_id":"daed7ba7-026b-4f01-a224-ce45a4c97c7a","resolution":{"observed_at":"2026-08-07T15:19:19.658188Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:19.725636Z","title":"End-to-end object detection with Transformers","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.725636Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:99d8ee997f40485012feb76e0064f06c0765b15cb6d38afba658bc5bdd121bf2","observation_id":"855a54ea-ad68-4366-abbd-579d6b118115","resolution":{"observed_at":"2026-08-07T15:19:19.725636Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:19.780664Z","title":"Conceptual 12M: Pushing web-scale image-text pre-training to recognize long-tail visual concepts","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.780664Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:67f44b48c62e54093a8f175a1ecd36c489147d2d36fbc51bc89159ae60fbbf84","observation_id":"7ee0e2cf-e7f0-44d8-bb00-fb55ffcf05ed","resolution":{"observed_at":"2026-08-07T15:19:19.780664Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2203.10036","last_updated":"2022-06-03T19:31:06Z","snapshot_observed_at":"2026-07-06T12:49:36.698431Z","submitted_at":"2022-03-18T16:09:53Z","title":"On the Generalization Mystery in Deep Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2203.10036","snapshot_observed_at":"2026-08-07T15:19:19.845283Z","title":"On the generalization mystery in deep learning","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.845283Z"},"links":{"cited_paper":"/paper/2203.10036","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:5039feeb3306c1ca63b1e7a8e78db6b3c355cc8e1ae4578cc070dd4399d2b54a","observation_id":"a1bf887f-f3ce-4661-90f6-4cea0df39587","resolution":{"observed_at":"2026-08-07T15:19:19.845283Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1905.06803","last_updated":"2019-10-03T16:10:46Z","snapshot_observed_at":"2026-08-07T20:39:30.321350Z","submitted_at":"2019-05-16T14:48:29Z","title":"How is Gaze Influenced by Image Transformations? Dataset and Model","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1905.06803","snapshot_observed_at":"2026-08-07T15:19:19.897209Z","title":"GazeGAN: A generative adversarial saliency model based on invariance analysis of human gaze during scene free viewing","venue":null,"work_id":null,"year":1905},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.897209Z"},"links":{"cited_paper":"/paper/1905.06803","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:3118ae0f69dc93a144f0ccbd70a74f1fd3f125b2890b9823735e9b615413384d","observation_id":"b476f343-eb58-4cc2-9c09-fa97034c4149","resolution":{"observed_at":"2026-08-07T15:19:19.897209Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:19.980008Z","title":"Benchmarking robustness of adaptation methods on pre-trained vision-language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:19.980008Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:ee1eb2819ca0a601a6b9461d226173a3def431e0eba55bed51ac4809c37660cc","observation_id":"bcebb1b8-27f7-4142-b2a2-20dea98afa22","resolution":{"observed_at":"2026-08-07T15:19:19.980008Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.039635Z","title":"PaLI: A jointly-scaled multilingual language-image model","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.039635Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:af00eef972ffb36b6f39a7b9bdf262390c9aa1b5a160e6767fbcfc13fdcbc45e","observation_id":"598a0aad-36a8-481c-a28f-9a6836da87f9","resolution":{"observed_at":"2026-08-07T15:19:20.039635Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.109559Z","title":"Reproducible scaling laws for contrastive language-image learning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.109559Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:42abe8dcef7b05aca5bde0c82590e55cd032e15d53d7e5fee8d2b91cfe68cd9f","observation_id":"ef6b2b1c-47c2-4147-bd8e-9b2dd89f72a3","resolution":{"observed_at":"2026-08-07T15:19:20.109559Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.157721Z","title":"ImageNet: A large-scale hierarchical image database","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.157721Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:3c4301e1d5335afc6dfbf3a197a6bdd6f91dfd3f0287a63817ab8e65ad02914b","observation_id":"c975d861-cfcb-4e43-8dee-c6d837ebcb6b","resolution":{"observed_at":"2026-08-07T15:19:20.157721Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.217827Z","title":"Understanding how image quality affects deep neural networks","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.217827Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:b437dce1f53834baf42ef4f0634bf038a8742741f3198da589b045881de93ba6","observation_id":"209d1777-b1a8-4071-b21d-35fe7cc44426","resolution":{"observed_at":"2026-08-07T15:19:20.217827Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.279623Z","title":"A study and comparison of human and deep learning recognition performance under visual distortions","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.279623Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:e0d7cb4de62882cf2ef7de2ea1aa6f89882e925b7329a1c8fa5c48751af3a26b","observation_id":"19da56fc-afdd-4ff4-b5fd-47edb87c4181","resolution":{"observed_at":"2026-08-07T15:19:20.279623Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.343402Z","title":"An image is worth 16x16 words: Transformers for image recognition at scale","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.343402Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:1b245d6f4412809b9f46a80249f6cae77fae76872e1f113a78aa5ea7e91d3c4f","observation_id":"71279c86-a262-4c11-94d9-597b79254a37","resolution":{"observed_at":"2026-08-07T15:19:20.343402Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.388459Z","title":"In search of robust measures of generalization","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.388459Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:254f1976e2d83b098afba724531d2f81d8f58385feb72349555967a54bdd50c4","observation_id":"5786a502-afa5-4e69-99c6-86e84196b361","resolution":{"observed_at":"2026-08-07T15:19:20.388459Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.428465Z","title":"The Pascal Visual Object Classes (VOC) Challenge","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.428465Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:be0ca4050dfae22dbc37e00841427fae16aa916b7f10241b69d301e9dd97fc07","observation_id":"17789bfb-4d71-441f-9880-3a1ccccbb1e1","resolution":{"observed_at":"2026-08-07T15:19:20.428465Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.466197Z","title":"Data filtering networks","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.466197Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:a35c7c4a1170f832eed8a9cae5e8e86b2335d4cad7c38fd04129586089d10e85","observation_id":"08b5f6c1-a538-4b65-9216-d55b191a7bb9","resolution":{"observed_at":"2026-08-07T15:19:20.466197Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1804.02767","last_updated":"2018-04-08T22:27:57Z","snapshot_observed_at":"2026-08-06T11:09:16.409556Z","submitted_at":"2018-04-08T22:27:57Z","title":"YOLOv3: An Incremental Improvement","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1804.02767","snapshot_observed_at":"2026-08-07T15:19:20.527690Z","title":"Yolov3: An incremental improvement","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.527690Z"},"links":{"cited_paper":"/paper/1804.02767","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:5b4671db45437881a6e1d7618df23c517cd0cee39c19c7b1413bbb85347997e5","observation_id":"493a9d8f-9524-41c3-b561-993dd1b28f0c","resolution":{"observed_at":"2026-08-07T15:19:20.527690Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.610942Z","title":"Shortcut learning in deep neural networks","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.610942Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:2d0479048ae61264c1819a6e1d4ead051eafffae5966209c2dae23de5f38d0d7","observation_id":"0c651a35-6db8-4968-a5f3-0445b3855f12","resolution":{"observed_at":"2026-08-07T15:19:20.610942Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.673436Z","title":"Generalisation in humans and deep neural networks","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.673436Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:dfefc94e2d2358d6a2f8ce02d4135bff2642d881c246d4f5f8deb7d84ccbd31c","observation_id":"65a6b1bd-b2f2-432d-8e40-ad8ceff170ec","resolution":{"observed_at":"2026-08-07T15:19:20.673436Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.735227Z","title":"Fast R-CNN","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.735227Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:eb6f6bb5382ae2249a2a4ce91427c7a5ad171c2afe34647670967986bb42fd4b","observation_id":"4c0a5a5e-f6f7-4d7c-b64a-133d1282be8d","resolution":{"observed_at":"2026-08-07T15:19:20.735227Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.796264Z","title":"Truth or backpropaganda? an empirical investigation of deep learning theory","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.796264Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:d43bc11bdadf8b7e9cddd12bc44e14106d73b4c62ae21ca6890db6aebb51c469","observation_id":"1b8741b6-fe87-425d-ab03-ce96fdb62e04","resolution":{"observed_at":"2026-08-07T15:19:20.796264Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.862947Z","title":"Strengths and weaknesses of deep learning models for face recognition against image degradations","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.862947Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:134007a100735a564c0c9be5f75ad12e1e97620369abd63f839ddbe3dfe72b2b","observation_id":"640954a7-9c64-4b22-9a0d-4d30ef52d832","resolution":{"observed_at":"2026-08-07T15:19:20.862947Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.937774Z","title":"Wukong: A 100 million large-scale Chinese cross-modal pre-training benchmark","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.937774Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:fb604e23274ba4db077f042e952d623f2494bc549da79049a3eb9ae6e21fc372","observation_id":"53ff878f-d79b-4332-bd3e-702e71ad4b61","resolution":{"observed_at":"2026-08-07T15:19:20.937774Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:20.996131Z","title":"ExifTool","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:20.996131Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:7fe8c49f1da4d3d0ce0f191b1f00ed8c2dbb979226312332ee2d08b5b04e6e3f","observation_id":"55ed9123-ca89-4702-9196-7ec69871bc29","resolution":{"observed_at":"2026-08-07T15:19:20.996131Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2012.10931","last_updated":"2021-03-11T09:16:45Z","snapshot_observed_at":"2026-07-06T10:26:16.221193Z","submitted_at":"2020-12-20T14:16:41Z","title":"Recent advances in deep learning theory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2012.10931","snapshot_observed_at":"2026-08-07T15:19:21.047681Z","title":"Recent advances in deep learning theory","venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:21.047681Z"},"links":{"cited_paper":"/paper/2012.10931","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:b10202bfe3c3a8d6e5aa31e5d450b2ac47203892266ba4311c6fa3aa0306b0c2","observation_id":"6cf7d35d-0f35-4928-b9b8-ce52546ffdb5","resolution":{"observed_at":"2026-08-07T15:19:21.047681Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:21.106520Z","title":"Masked autoencoders are scalable vision learners","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:21.106520Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:60d9744e40b087ff119272066fb62f027e7854378b7741a756105edb1c962f91","observation_id":"288c359c-eaa6-4bb9-8d6b-4c21a280a8d6","resolution":{"observed_at":"2026-08-07T15:19:21.106520Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:21.183151Z","title":"Deep residual learning for image recognition","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:21.183151Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:5a417564363045dc7315b59a05a2fdfdf0f923e67b2e1f210b77aa14db79a5ab","observation_id":"29206b86-5906-434f-8db6-ab582d887b88","resolution":{"observed_at":"2026-08-07T15:19:21.183151Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:34.607572Z","title":"Benchmarking neural network robustness to common corruptions and perturbations","venue":null,"work_id":"a1e76926-6ebc-4b25-afef-0096750056bb","year":2019},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:21.251721Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:e3cc13db6f1c7cd2b3470ae0971283c4b8f62f9297b44808394ebcd394614766","observation_id":"706f2648-e6da-414d-a1b9-9faeae834e69","resolution":{"observed_at":"2026-08-07T15:19:34.637540Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:21.360393Z","title":"Natural adversarial examples","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:21.360393Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:581d9f293548aaa4718d44610899e41250e76fdb28417126a649bb0ba658de25","observation_id":"f89b8dbb-840b-4e63-ad09-449a39ed00cd","resolution":{"observed_at":"2026-08-07T15:19:21.360393Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:34.479431Z","title":"Scene recognition with CNNs: Objects, scales and dataset bias","venue":null,"work_id":"7c706262-0441-4c8e-a2bb-e17384f25688","year":2016},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:21.484034Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:957023f507a74804f6434c9b1365055db335d133c30040cb149730e2d0a72bfb","observation_id":"ebc3972f-0de0-4363-8a0c-1131ee6aa84a","resolution":{"observed_at":"2026-08-07T15:19:34.528094Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:34.381446Z","title":"A survey on hallucination in large language models: Principles, taxonomy, challenges, and open questions.ACM Transactions on Information Systems, 43(2):1–55, 2025","venue":null,"work_id":"5be26f1a-cf25-4e1a-99e4-2e7807e004d5","year":2025},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:21.674685Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:63daf6f88a78ff3d2c39e9f1daccb77bb4f91a53e708d2b0cf84a162d020d6c0","observation_id":"31b179a8-e86f-44f9-a0b9-05dd8b7137d6","resolution":{"observed_at":"2026-08-07T15:19:34.425257Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:21.803108Z","title":"Visual Genome: Connecting language and vision using crowdsourced dense image annotations","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:21.803108Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:035fe9500c024cc01193a2d1559730f6674a4d7b162fddc5bfaee6feb2b3b819","observation_id":"5ba2f02f-4899-4cbc-be6c-9653de5ae094","resolution":{"observed_at":"2026-08-07T15:19:21.803108Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:34.263940Z","title":"Robustness and repeatability of saliency models subjected to visual degradations","venue":null,"work_id":"6889000c-d7d7-4d97-a273-fbb7881ffd5b","year":2011},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:21.906664Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:15a835de3439dccdaff97a7067375cad0809e2adcdaab1ff2b09968912fbbde5","observation_id":"aed55045-7d95-49bc-9929-5c445321feee","resolution":{"observed_at":"2026-08-07T15:19:34.321821Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:34.158853Z","title":"Can multiple-choice questions really be useful in detecting the abilities of llms? In LREC-COLING, 2024","venue":null,"work_id":"e1d55674-ec45-4142-b03b-3f866d4dfb57","year":2024},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:22.039374Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:703c51602271afece7429121eb1c7b97f0a017c959e0f67d4785fb5ca14e0834","observation_id":"1c5abedc-0721-415f-a46d-3ab6432486a1","resolution":{"observed_at":"2026-08-07T15:19:34.207967Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:22.173824Z","title":"Exploring plain vision transformer backbones for object detection","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:22.173824Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:d4dd7815d43ea1999f4800607b351f8b63af7847bb039d4496809eae5cd0e776","observation_id":"3a2ab2d9-63fb-45c2-b118-8ba3db672f89","resolution":{"observed_at":"2026-08-07T15:19:22.173824Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:33.998574Z","title":"VILA: On pre-training for visual language models","venue":null,"work_id":"38b68219-0faf-40da-a086-361779cb7076","year":2024},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:22.300921Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:93e0eda9fe22d3018b7ff1b288785637ebde6e29b6f019eaa396565038852c8f","observation_id":"dc3fc9cc-7409-4d14-aaeb-78aaec0348ff","resolution":{"observed_at":"2026-08-07T15:19:34.063584Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:33.759745Z","title":"Microsoft COCO: Common objects in context","venue":null,"work_id":"dc72ea3a-8d52-44a6-8431-4983dbd418a1","year":2014},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:22.407479Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:0a05506d03043cede24d794ae2122b4a5855a84dff189990a0565d85283e1b1f","observation_id":"619d1d5a-41c9-4116-b254-ebcef56bf7b6","resolution":{"observed_at":"2026-08-07T15:19:33.865626Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:22.497695Z","title":"Improved baselines with visual instruction tuning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:22.497695Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:a1cc01a1a9f57e9444a1e347b5f289538652fa0b0c8b9030c278ee550ae4b147","observation_id":"26c1112d-03f5-425f-8319-5b2dd650fb27","resolution":{"observed_at":"2026-08-07T15:19:22.497695Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:33.523343Z","title":"Evaluation of LBP and deep texture descriptors with a new robustness benchmark","venue":null,"work_id":"4177a964-acc3-462c-b686-9327c951170d","year":2016},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:22.588683Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:f7010153e80f7863d3fa9a286fb3b63bd0d85f071b9aa6de6cfddd95b6578a36","observation_id":"639e21f4-c038-4061-919b-9a83df9285bd","resolution":{"observed_at":"2026-08-07T15:19:33.648459Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:33.202318Z","title":"Grounding DINO: Marrying DINO with grounded pre-training for open-set object detection","venue":null,"work_id":"8d8e331c-5232-4709-8d55-159a24ed9ccc","year":2024},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:22.677135Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:f4e6b1d68a11dc39aca1ed7471287d02a0586be02ceb9b31e0193313098604a0","observation_id":"324b4052-66ef-4ccc-99a2-ff3506e8d997","resolution":{"observed_at":"2026-08-07T15:19:33.337150Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:32.892398Z","title":"SSD: Single shot multibox detector","venue":null,"work_id":"04cab912-c648-45ea-8b08-6aede838d93a","year":2016},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:22.786761Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:e8cff4cfd71aaa794a9ceeee10e160692ae99d2c2df34519b1c1387726267e51","observation_id":"fef0914d-8f5d-4e1c-a0c5-427538865366","resolution":{"observed_at":"2026-08-07T15:19:33.046187Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:32.679123Z","title":"Swin Transformer: Hierarchical vision transformer using shifted windows","venue":null,"work_id":"462400c4-dd35-4e85-ab0f-7cd9fef04027","year":2021},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:22.877568Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:8021afaa1ae85e0f17762799b6a519a2a0eb9101e0b86bc6db01d8a760a92dab","observation_id":"945619b5-00c2-49ca-9b58-a611e4e7e8c1","resolution":{"observed_at":"2026-08-07T15:19:32.791117Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:32.479200Z","title":"A decade’s battle on dataset bias: Are we there yet? In ICLR, 2025","venue":null,"work_id":"f97cd2e6-a275-4e64-a613-6610d9aaebb4","year":2025},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:22.945464Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:65451e0682683c5b8eb84290329e74b8c219ff1f1c2fd92a042dc886a04ac747","observation_id":"51899197-c343-4d66-a52b-9e86f51699dd","resolution":{"observed_at":"2026-08-07T15:19:32.525259Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:23.039693Z","title":"A convnet for the 2020s","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:23.039693Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:a0c90529fa01ce5ca5de5e525037a8d0d485f91100f7f824a695025ec04615b5","observation_id":"15128ed5-502c-48b1-b05b-96d118678c84","resolution":{"observed_at":"2026-08-07T15:19:23.039693Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.05525","last_updated":"2024-03-11T16:47:41Z","snapshot_observed_at":"2026-08-05T16:51:32.094151Z","submitted_at":"2024-03-08T18:46:00Z","title":"DeepSeek-VL: Towards Real-World Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.05525","snapshot_observed_at":"2026-08-07T15:19:23.133913Z","title":"DeepSeek-VL: Towards real-world vision-language understand- ing","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:23.133913Z"},"links":{"cited_paper":"/paper/2403.05525","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:e9b57886a3bd4b09dae5cd821377e685d7668291517292d5123c8542f1e54d42","observation_id":"98d14744-9442-4ca0-85a3-2cd336c1ca93","resolution":{"observed_at":"2026-08-07T15:19:23.133913Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1907.07484","last_updated":"2020-03-31T08:42:46Z","snapshot_observed_at":"2026-08-03T06:25:53.775445Z","submitted_at":"2019-07-17T12:51:10Z","title":"Benchmarking Robustness in Object Detection: Autonomous Driving when Winter is Coming","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1907.07484","snapshot_observed_at":"2026-08-07T15:19:23.249352Z","title":"Benchmarking robustness in object detection: Autonomous driving when winter is coming","venue":null,"work_id":null,"year":1907},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:23.249352Z"},"links":{"cited_paper":"/paper/1907.07484","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:c6076f9b8263285ba550ce6a18a0cc8e0abb93c9f29b34c2f9355a0ca4110476","observation_id":"c59d4fc5-25a7-41e5-be83-97b936c9a535","resolution":{"observed_at":"2026-08-07T15:19:23.249352Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:23.403274Z","title":"Simple open-vocabulary object detection","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:23.403274Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:288278bddd035897fa6442fcfedac81779c9a74a3531ce118abcc1fb8ba4c807","observation_id":"1db5fb59-572a-4b99-9767-06d9a780c805","resolution":{"observed_at":"2026-08-07T15:19:23.403274Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1906.02337","last_updated":"2019-06-05T22:23:43Z","snapshot_observed_at":"2026-07-06T07:58:18.558050Z","submitted_at":"2019-06-05T22:23:43Z","title":"MNIST-C: A Robustness Benchmark for Computer Vision","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1906.02337","snapshot_observed_at":"2026-08-07T15:19:23.494642Z","title":"MNIST-C: A robustness benchmark for computer vision","venue":null,"work_id":null,"year":1906},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:23.494642Z"},"links":{"cited_paper":"/paper/1906.02337","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:0e5326333f2fc79d8858daf354553f39674b90bae6c99623e425740d692bddc5","observation_id":"f91d21e4-f869-4068-b3bd-be07f29d57fe","resolution":{"observed_at":"2026-08-07T15:19:23.494642Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:32.207588Z","title":"Uniform convergence may be unable to explain generalization in deep learning","venue":null,"work_id":"520f3156-c45d-42cf-a0ad-21f39533ea05","year":2019},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:23.582498Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:23c274f6fa0c897b2e9caaa78dce153311b5344bc055fa49a077bbdc54da32af","observation_id":"8e1e7b56-9955-405c-9b9a-4405ac89ffa3","resolution":{"observed_at":"2026-08-07T15:19:32.338981Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:32.004362Z","title":"Evaluation of visual saliency analysis algorithms in noisy images","venue":null,"work_id":"576f2f2e-1132-433e-a2ac-97998b152819","year":2016},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:23.694774Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:8ccb7211e34dee2840235b4003aa8539403518fe03bb70de6980461f0f266a13","observation_id":"678895ed-6722-4606-8d46-19259637b6b4","resolution":{"observed_at":"2026-08-07T15:19:32.050442Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:31.931519Z","title":"Localization recall precision (LRP): A new performance metric for object detection","venue":null,"work_id":"50af079d-77ed-4404-97be-f8e64749e7b3","year":2018},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:23.764029Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:eaba5e6847d6a3aaac6317f6840243260cb43c544a3ab5b87907bad601c7c1d5","observation_id":"8836ead4-f92a-470f-9555-310c911324b6","resolution":{"observed_at":"2026-08-07T15:19:31.965490Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:31.841758Z","title":"One metric to measure them all: Localisation recall precision (LRP) for evaluating visual detection tasks","venue":null,"work_id":"adebffa0-2e92-4441-98fe-c32197796098","year":2021},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:23.867701Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:0d85b038b0e6d585cef58b8e7022192e13f604eb90451bffe3665ff0851adffc","observation_id":"be18f93a-1175-47b4-bf17-96c2b7fd1be6","resolution":{"observed_at":"2026-08-07T15:19:31.887237Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:31.729613Z","title":"Im2Text: Describing images using 1 million captioned photographs","venue":null,"work_id":"b46a283e-9519-4536-8f92-0ae34689e69f","year":2011},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:23.965521Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:e71bbc198cd4afb9bc165ebd8b1327c0f6d74eef2b699b2f8eb3de0a28d1cd31","observation_id":"064e4042-e7b4-4860-889f-2c49815d9fa8","resolution":{"observed_at":"2026-08-07T15:19:31.777566Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2212.06137","last_updated":"2022-12-12T18:59:58Z","snapshot_observed_at":"2026-07-06T14:29:40.215055Z","submitted_at":"2022-12-12T18:59:58Z","title":"NMS Strikes Back","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2212.06137","snapshot_observed_at":"2026-08-07T15:19:24.054663Z","title":"NMS strikes back","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:24.054663Z"},"links":{"cited_paper":"/paper/2212.06137","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:d32d14cd56cc724b795e9b157ed5493122b6c3b7c391cabe570d3943276f7c74","observation_id":"1de0c7c4-da40-4962-bf08-9f49e885da97","resolution":{"observed_at":"2026-08-07T15:19:24.054663Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:31.630835Z","title":"PsychoPy2: Experiments in behavior made easy","venue":null,"work_id":"350b5e6f-29b0-42cd-a7b8-34dabf315a5c","year":2019},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:24.164737Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:d57a7a2daf71ad0b19c7be3d28f49fc3d4d7beaccd6b9e1ea847b751a39f69ca","observation_id":"3aeee163-9358-42d2-870b-aa451f762c05","resolution":{"observed_at":"2026-08-07T15:19:31.677414Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:31.564960Z","title":"Dataset issues in object recognition","venue":null,"work_id":"8ae90f9a-49f8-4e83-b33d-a89fdd588d43","year":2006},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:24.255104Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:d12cc21db6539683659dc14422d9aeb3b1c458ceff5068068f33abf2a659f831","observation_id":"38f5fa3c-983b-4cce-be54-2bdeb12e91b3","resolution":{"observed_at":"2026-08-07T15:19:31.581371Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:31.505786Z","title":"Basics Photography 07: Exposure","venue":null,"work_id":"68ec408e-354b-435b-8e20-701509d2c7e9","year":2009},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:24.352300Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:e80bdbfce761541e687b6eb17261c2f4a7c4d43c79cbed71f6d494a146b40dd7","observation_id":"e4e383db-aaf2-42bd-bfcf-e4e0328589ca","resolution":{"observed_at":"2026-08-07T15:19:31.524617Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:31.423408Z","title":"Dataset growth","venue":null,"work_id":"c5698e5b-3a59-4b93-8510-c588997b41ca","year":2024},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:24.422701Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:a958c122b646612f921355b2fab8aee9ebb2054bd9380bcd566e63b257f46e17","observation_id":"52068316-e8b2-42be-87c1-d430c621eadb","resolution":{"observed_at":"2026-08-07T15:19:31.450882Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:24.493728Z","title":"Learning transferable visual models from natural language supervision","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:24.493728Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:e65094be50fa6a46af06dce6c2de1028413b9b230bc59afa0655a4e75d226e75","observation_id":"a6241244-6c3c-4a4c-81c8-44d348e5e00c","resolution":{"observed_at":"2026-08-07T15:19:24.493728Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:31.350586Z","title":"Camera exposure determination","venue":null,"work_id":"fbe5b429-d642-4db2-942c-d4cbbd0683d4","year":2000},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:24.570752Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:6f93d9e0b54853b7a872cb9811cdd53c955bbd56175501c88b0e95f9fda541da","observation_id":"5c21c847-ff57-49aa-8b73-93f0da92333c","resolution":{"observed_at":"2026-08-07T15:19:31.376115Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2104.10972","last_updated":"2021-08-05T15:04:28Z","snapshot_observed_at":"2026-08-06T21:17:51.553405Z","submitted_at":"2021-04-22T10:10:14Z","title":"ImageNet-21K Pretraining for the Masses","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2104.10972","snapshot_observed_at":"2026-08-07T15:19:24.655809Z","title":"ImageNet-21K pretraining for the masses","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:24.655809Z"},"links":{"cited_paper":"/paper/2104.10972","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:fba2b5a46fdb801c8d4631efb3f44cb3f7151624420b9693573e51ed02006597","observation_id":"176c2969-7266-406d-8936-776ed5122aba","resolution":{"observed_at":"2026-08-07T15:19:24.655809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1610.06756","last_updated":"2016-10-21T12:14:50Z","snapshot_observed_at":"2026-07-06T05:15:29.158422Z","submitted_at":"2016-10-21T12:14:50Z","title":"Fine-grained Recognition in the Noisy Wild: Sensitivity Analysis of Convolutional Neural Networks Approaches","version":1},"cited_work":{"arxiv_id":"1610.06756","doi":null,"metadata_source":"pith","pith_arxiv_id":"1610.06756","snapshot_observed_at":"2026-08-07T15:19:28.087097Z","title":"Fine-grained Recognition in the Noisy Wild: Sensitivity Analysis of Convolutional Neural Networks Approaches","venue":"cs.CV","work_id":"b4ba972d-d724-4042-bcab-8f8413fe35ad","year":2016},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:24.722976Z"},"links":{"cited_paper":"/paper/1610.06756","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:e4fe0cd66a64286f8a0e8f9e32c580e0ab05267c8de1009a0508459073deea36","observation_id":"575405ef-8b33-45d8-ab56-c7b6c52c3891","resolution":{"observed_at":"2026-08-07T15:19:28.124268Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1807.10108","last_updated":"2025-02-18T16:40:45Z","snapshot_observed_at":"2026-07-06T06:52:28.777662Z","submitted_at":"2018-07-26T13:20:57Z","title":"Effects of Degradations on Deep Neural Network Architectures","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1807.10108","snapshot_observed_at":"2026-08-07T15:19:24.808208Z","title":"Effects of degradations on deep neural network architectures","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:24.808208Z"},"links":{"cited_paper":"/paper/1807.10108","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:6c30cb3536f39a8f28fb497aaa63b435afca4d65676d3614e66b036bd1b2b09a","observation_id":"2b14ba7b-7756-4ebd-b137-94e0e4e61ac0","resolution":{"observed_at":"2026-08-07T15:19:24.808208Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:31.274812Z","title":"LAION-400M: Open dataset of CLIP-filtered 400 million image-text pairs","venue":null,"work_id":"1c97df61-8933-4235-b9df-d6f9c97785a3","year":2021},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:24.884493Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:69d67fdb3b64bed8c6025d2f1b726e7f6d15cac855f95729111e2d45206e2e9d","observation_id":"77f833de-0a1b-48ee-a872-ad1eec381b19","resolution":{"observed_at":"2026-08-07T15:19:31.332125Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:31.159482Z","title":"Objects365: A large-scale, high-quality dataset for object detection","venue":null,"work_id":"87f32234-4009-4426-a00e-a9098a5fa607","year":2019},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:24.952231Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:ad826ea4482f08a6948a18ae5a2eb950b95f72bbb7cd10f0de912b551fbb4639","observation_id":"f3fbede3-007b-4e27-9875-cd04042e6249","resolution":{"observed_at":"2026-08-07T15:19:31.205607Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:25.025894Z","title":"Conceptual Captions: A cleaned, hypernymed, image alt-text dataset for automatic image captioning","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:25.025894Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:c0ec65c62f9903e219c8157fe0be5dc6f2d7724593c5a318dfc4f434181ffaa9","observation_id":"2d0a0782-231a-4b46-997f-25ea2148dbf1","resolution":{"observed_at":"2026-08-07T15:19:25.025894Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:31.058079Z","title":"Assessing visually-continuous corruption robustness of neural networks relative to human performance","venue":null,"work_id":"9d477479-2fb5-46fe-847e-6ecc8111b8a9","year":2025},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:25.157576Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:239395e9b2cd51963ec12528b5d7e18d4ffac71d989be35aff4b7fb826a7eece","observation_id":"b741f4dd-794b-4a5e-941b-21061730606f","resolution":{"observed_at":"2026-08-07T15:19:31.094307Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:25.244910Z","title":"Very deep convolutional networks for large-scale image recognition","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:25.244910Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:a101b70e6619264a286f2c2ccef542b261cdf593ff73ffcbf9c02a78b298cfc9","observation_id":"056ae939-c8ba-4e55-bc43-add5667829f3","resolution":{"observed_at":"2026-08-07T15:19:25.244910Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:30.964282Z","title":"WIT: Wikipedia- based image text dataset for multimodal multilingual machine learning","venue":null,"work_id":"74903f08-4985-4f88-9d4d-db398b9daeec","year":2021},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:25.327393Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:3a18fcf672bb533f55af498bd65873805aa51fba8f0ba909c0e81e94c11b547d","observation_id":"253681a5-fd2d-4724-9f71-2e85fe25c7eb","resolution":{"observed_at":"2026-08-07T15:19:31.000899Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:30.826707Z","title":"Convnets and ImageNet beyond accuracy: Understanding mistakes and uncovering biases","venue":null,"work_id":"5e45b466-11c4-443c-a129-b0eedb2c890f","year":2018},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:25.393093Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:dcd85a022998eae2ad8f36e5ec97480f09b662e464bee4b09d6e38074575c834","observation_id":"cc738568-53af-43ba-a433-1add98757456","resolution":{"observed_at":"2026-08-07T15:19:30.877662Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:30.709164Z","title":"A survey on statistical theory of deep learning: Approximation, training dynamics, and generative models","venue":null,"work_id":"12150560-026d-4024-ab2d-4f0c967e213f","year":2024},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:25.482974Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:a5039aa35b3f1d421cc65fbd3d59f9e326d9bb3271ba394b6b8408c898faea59","observation_id":"881eaa8d-7e37-4c1c-8b4a-7a58ce644793","resolution":{"observed_at":"2026-08-07T15:19:30.773502Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:30.572024Z","title":"YFCC100M: The new data in multimedia research","venue":null,"work_id":"f1587566-6587-45a3-acb9-0bcf5eb3150d","year":2016},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:25.575770Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:5e054701a8fb3ea3ee856074ed1a8233611a1878eeb3ea9c61fdd644feb3c76e","observation_id":"da6e972b-1125-4cf3-8e08-3d1c62d1d3ca","resolution":{"observed_at":"2026-08-07T15:19:30.633096Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:30.419258Z","title":"A deeper look at dataset bias","venue":null,"work_id":"325cd835-8bf0-4a83-94ab-6682cf195de2","year":2017},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:25.662541Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:f13a4cafd3a41705a4f27b9c1b1e7d37fb55c01f0a3b18f89e201b51b8e09ff8","observation_id":"a247d5a6-53dd-4614-957c-824d161e4755","resolution":{"observed_at":"2026-08-07T15:19:30.458429Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:30.319603Z","title":"Unbiased look at dataset bias","venue":null,"work_id":"9507185f-0d0b-434e-9a75-c6728367ed80","year":2011},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:25.726074Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:28088e07a6b994505f244bcedabf079621fe29350acb51c15b8eace3c093b42f","observation_id":"9b54e344-8bab-4f6b-93cb-6ce3cb914cdc","resolution":{"observed_at":"2026-08-07T15:19:30.363494Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:30.225278Z","title":"A robustness analysis of deep q networks","venue":null,"work_id":"3e97d907-69e1-40a4-92a2-d81b9eaa3461","year":2016},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:25.787742Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:3ae697b079f9ed165156063d9464d2ee10de467450dddc68e7a948beaf0957a2","observation_id":"dea93efe-0c10-40c3-97dc-5ea4a64fb4af","resolution":{"observed_at":"2026-08-07T15:19:30.263952Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:30.105684Z","title":"From ImageNet to image classification: Contextualizing progress on benchmarks","venue":null,"work_id":"6742f3e8-ea7a-4748-be34-635bdc504b37","year":2020},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:25.850341Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:98e71ee0a59d199728078729d3843853a39c23997b9dd034f0ecfb599aea0d05","observation_id":"addde9cb-4718-4d73-b4a1-254769b72ec9","resolution":{"observed_at":"2026-08-07T15:19:30.173180Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:30.020590Z","title":"Why does data-driven beat theory-driven computer vision? In ICCVW, 2019","venue":null,"work_id":"a82568fb-1ca8-424b-9894-e00d36e68b24","year":2019},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:25.934515Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:db231a95001532e33a84d6b5a86a7e0b8cf47e8faa33634de6b740b54c976c28","observation_id":"1f2f6423-3c31-40dc-a7d6-fd87e0c8fa2e","resolution":{"observed_at":"2026-08-07T15:19:30.054639Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2105.09934","last_updated":"2022-04-30T14:20:14Z","snapshot_observed_at":"2026-07-06T11:11:22.798053Z","submitted_at":"2021-05-20T17:54:48Z","title":"Probing the Effect of Selection Bias on Generalization: A Thought Experiment","version":2},"cited_work":{"arxiv_id":"2105.09934","doi":null,"metadata_source":"pith","pith_arxiv_id":"2105.09934","snapshot_observed_at":"2026-08-07T15:19:27.878419Z","title":"Probing the Effect of Selection Bias on Generalization: A Thought Experiment","venue":"cs.CV","work_id":"604f5990-7f63-4d00-8352-5fb6048aeba5","year":2021},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:26.001686Z"},"links":{"cited_paper":"/paper/2105.09934","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:0219469d80a03ffd3968d8fea23479217c2b8e42d4173d55611c049b729fa9ab","observation_id":"4fac7bb5-78b1-488f-897f-8b0637face16","resolution":{"observed_at":"2026-08-07T15:19:27.944051Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:29.906999Z","title":"CSPNet: A new backbone that can enhance learning capability of CNN","venue":null,"work_id":"c51c83ee-eba8-4c33-9b14-a1f1c8fed13e","year":2020},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:26.149027Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:27529fad6b30e842179730c5bbd8c56f7ee734eec8c59fef29ae278f1230a36c","observation_id":"ed120f39-f863-4446-97e2-8248020b118d","resolution":{"observed_at":"2026-08-07T15:19:29.947802Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12191","last_updated":"2024-10-03T15:54:49Z","snapshot_observed_at":"2026-08-06T05:35:29.109022Z","submitted_at":"2024-09-18T17:59:32Z","title":"Qwen2-VL: Enhancing Vision-Language Model's Perception of the World at Any Resolution","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12191","snapshot_observed_at":"2026-08-07T15:19:26.245997Z","title":"QwenV2-VL: Enhancing vision-language model’s perception of the world at any resolution","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:26.245997Z"},"links":{"cited_paper":"/paper/2409.12191","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:cfb6aec411425272fbd00ca4b1fa4eca8f73693002fe8254dc18d15fa06c0d45","observation_id":"3e95fd5b-16fe-4d36-aa17-ab4035056fbd","resolution":{"observed_at":"2026-08-07T15:19:26.245997Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.10442","last_updated":"2025-04-07T09:09:39Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-11-15T18:59:27Z","title":"Enhancing the Reasoning Ability of Multimodal Large Language Models via Mixed Preference Optimization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.10442","snapshot_observed_at":"2026-08-07T15:19:26.317605Z","title":"Enhancing the reasoning ability of multimodal large language models via mixed preference optimization","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":90,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:26.317605Z"},"links":{"cited_paper":"/paper/2411.10442","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:badbf3e1a368cb5ff93ce9bb04aab705ee09df0cd86df475af535cbf30366a9e","observation_id":"a11ce7f6-fd21-4471-9cd0-9925aa7ef474","resolution":{"observed_at":"2026-08-07T15:19:26.317605Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:29.808276Z","title":"InternImage: Exploring large-scale vision foundation models with deformable convolutions","venue":null,"work_id":"2f908223-8e7d-414e-9e8a-5d99aa3ecfc6","year":2023},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:26.393445Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:b0fefd34bbc84314b0c5ad0f0cf7d34183f0084e2ff732f36a9fcb24d4737fe4","observation_id":"001ce651-1ee8-4038-865f-30f59cb49224","resolution":{"observed_at":"2026-08-07T15:19:29.858990Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:29.673646Z","title":"Physical adversarial attack meets computer vision: A decade survey","venue":null,"work_id":"88a2f3b0-6616-437b-adfc-94e4c96e451f","year":2024},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":92,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:26.456177Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:79bf7ccdc1783f76bf06a9613a49df6fe3134a22bf7d3d0af639b1af9b327eb7","observation_id":"a02a840d-b72a-407e-af89-5ec714aaa83e","resolution":{"observed_at":"2026-08-07T15:19:29.734336Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1705.05685","last_updated":"2017-05-16T12:47:47Z","snapshot_observed_at":"2026-07-06T05:42:51.694461Z","submitted_at":"2017-05-16T12:47:47Z","title":"Active Control of Camera Parameters for Object Detection Algorithms","version":1},"cited_work":{"arxiv_id":"1705.05685","doi":null,"metadata_source":"pith","pith_arxiv_id":"1705.05685","snapshot_observed_at":"2026-08-07T15:19:27.678341Z","title":"Active Control of Camera Parameters for Object Detection Algorithms","venue":"cs.CV","work_id":"a4fbda3d-2571-4403-8919-08bbd3eef264","year":2017},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":93,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:26.512057Z"},"links":{"cited_paper":"/paper/1705.05685","citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:df95aef5c3f962283bb6406cedfc6ddf6cc516ac583469fbc85f34c7de18e614","observation_id":"69e7edf9-65f0-4e72-b23f-31a4bffd27ec","resolution":{"observed_at":"2026-08-07T15:19:27.775717Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:29.532240Z","title":"Statistic analysis of millions of digital photos","venue":null,"work_id":"5a1e4c69-ab5c-442a-9752-285bdfee72a7","year":2008},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":94,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:26.568118Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:4d6d66b8a77371561f05c8a554453b3a71b0c2546d9db07b35a673a8f39e6dab","observation_id":"475974c2-f41e-409d-9dca-43fd7da44412","resolution":{"observed_at":"2026-08-07T15:19:29.575010Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:29.402011Z","title":"Statistic analysis of millions of digital photos 2017","venue":null,"work_id":"2e1db41c-38df-4bb1-afd8-5b791246ec85","year":2017},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":95,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:26.648030Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:8f06042ba8368f300e14296e35b01e778803f6e4cd0237dccbdac559037232b1","observation_id":"94322036-7102-44de-96b0-41ce336efb80","resolution":{"observed_at":"2026-08-07T15:19:29.453066Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:29.297337Z","title":"Does robustness on ImageNet transfer to downstream tasks? In CVPR, 2022","venue":null,"work_id":"95cab148-2a19-4f4c-8917-4e32973f3f04","year":2022},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":96,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:26.705156Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:719b9ab46a1f55b9cf77415a10bbb59098e4bf514a2a4688b4c7ef71f89825a6","observation_id":"d4629c19-863f-480b-9726-f29c5216ec4a","resolution":{"observed_at":"2026-08-07T15:19:29.345097Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:29.169435Z","title":"Yocto-Light-V3 product page","venue":null,"work_id":"9dc9e4b3-7887-45b8-8896-ad37e8f51661","year":2024},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":97,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:26.766084Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:b3ec3cfd1fe3ecf0d9d6b3742972f980808364ead388890007a73bdfb77a050d","observation_id":"61688f7f-2aa6-43d9-9a18-55e5301e22c4","resolution":{"observed_at":"2026-08-07T15:19:29.240502Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:29.065064Z","title":"Understanding bias in large-scale visual datasets","venue":null,"work_id":"1f83db04-5deb-473f-b6dd-970f9cbb643b","year":2024},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":98,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:26.875136Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:3f9329351a9f315655732fa02d5882e05738d68ae3f2fbe4cf86cf77ab7dbae2","observation_id":"c46afc1e-fd01-4207-9265-8cb60b5641ad","resolution":{"observed_at":"2026-08-07T15:19:29.105796Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:26.925706Z","title":"Sigmoid loss for language image pre-training","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":99,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:26.925706Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:795fa98a070ae969a8edab0e88ef5eaa17f7d3e0336661639677981c52fcb6f7","observation_id":"ceb31487-5a14-4cec-af1c-b673c2123afd","resolution":{"observed_at":"2026-08-07T15:19:26.925706Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:19:28.952360Z","title":"Understanding deep learning (still) requires rethinking generalization","venue":null,"work_id":"f6c3a023-16e7-4278-91a9-11c42519ce43","year":2021},"citing_paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks","version":1},"reference_index":100,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:26.982116Z"},"links":{"citing_paper":"/paper/2505.15628"},"observation_digest":"sha256:25f4781838af2b11e27797d5b3d6b256c29a6f0b1b612e958762cf797dd30753","observation_id":"5d1bd571-986b-4b9d-bfac-c48f7a19c2da","resolution":{"observed_at":"2026-08-07T15:19:28.993090Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2505.15628","last_updated":"2025-05-21T15:14:34Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-08T02:16:21.927195Z","submitted_at":"2025-05-21T15:14:34Z","title":"SNAP: A Benchmark for Testing the Effects of Capture Conditions on Fundamental Vision Tasks"},"reference_resolution":{"displayed":100,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":54,"verified_exact":3,"verified_fuzzy":43},"total_outbound_references":105},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 100 of 105 outbound references and 0 inbound Pith citation observations for arXiv:2505.15628."}