{"as_of":"2026-08-12T18:45:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:8edd9f635d49a3d14302078fd1fb77bbfd41bba69cfbb44059d2d1693cb35e26","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":13,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":13,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":13,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":13,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T12:43:47.568564Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-06-28T19:32:35.285668Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2211.05778","last_updated":"2023-04-17T11:51:12Z","snapshot_observed_at":"2026-08-11T17:05:48.093833Z","submitted_at":"2022-11-10T18:59:04Z","title":"InternImage: Exploring Large-Scale Vision Foundation Models with Deformable Convolutions","version":4},"cited_work":{"arxiv_id":"2211.05778","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2211.05778","snapshot_observed_at":"2026-06-28T19:32:35.285668Z","title":"Internim- age: Exploring large-scale vision foundation mod- els with deformable convolutions","venue":null,"work_id":"f6269124-1471-45b6-863b-d94f58db1a6b","year":2022},"citing_paper":{"arxiv_id":"2301.01201","last_updated":"2026-04-06T09:09:39Z","snapshot_observed_at":"2026-07-06T14:37:04.988254Z","submitted_at":"2022-12-20T07:32:12Z","title":"Uncertainty in Real-Time Semantic Segmentation on Embedded Systems","version":6},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-24T10:09:23.030973Z"},"links":{"cited_paper":"/paper/2211.05778","citing_paper":"/paper/2301.01201"},"observation_digest":"sha256:f95e98f4582a2bec155e1816520996eb6150902b8ba480b70febf9a82744c420","observation_id":"7ce8d913-bc22-4381-928d-0abbdae14e56","resolution":{"observed_at":"2026-05-24T10:14:19.010661Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.05778","last_updated":"2023-04-17T11:51:12Z","snapshot_observed_at":"2026-08-11T17:05:48.093833Z","submitted_at":"2022-11-10T18:59:04Z","title":"InternImage: Exploring Large-Scale Vision Foundation Models with Deformable Convolutions","version":4},"cited_work":{"arxiv_id":"2211.05778","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2211.05778","snapshot_observed_at":"2026-06-28T19:32:35.285668Z","title":"Internim- age: Exploring large-scale vision foundation mod- els with deformable convolutions","venue":null,"work_id":"f6269124-1471-45b6-863b-d94f58db1a6b","year":2022},"citing_paper":{"arxiv_id":"2303.16199","last_updated":"2024-09-18T23:54:36Z","snapshot_observed_at":"2026-08-06T06:36:02.994951Z","submitted_at":"2023-03-28T17:59:12Z","title":"LLaMA-Adapter: Efficient Fine-tuning of Language Models with Zero-init Attention","version":3},"reference_index":280,"source":"arxiv_source","source_observed_at":"2026-05-14T23:07:42.245641Z"},"links":{"cited_paper":"/paper/2211.05778","citing_paper":"/paper/2303.16199"},"observation_digest":"sha256:b9e28f896c25049f45867c26fc42817cb2b3f3eb2250deddedf597683027dec6","observation_id":"b0edfe37-5b78-4572-986e-9fe6a0bd3a1d","resolution":{"observed_at":"2026-05-14T23:07:42.852643Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.05778","last_updated":"2023-04-17T11:51:12Z","snapshot_observed_at":"2026-08-11T17:05:48.093833Z","submitted_at":"2022-11-10T18:59:04Z","title":"InternImage: Exploring Large-Scale Vision Foundation Models with Deformable Convolutions","version":4},"cited_work":{"arxiv_id":"2211.05778","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2211.05778","snapshot_observed_at":"2026-06-28T19:32:35.285668Z","title":"Internim- age: Exploring large-scale vision foundation mod- els with deformable convolutions","venue":null,"work_id":"f6269124-1471-45b6-863b-d94f58db1a6b","year":2022},"citing_paper":{"arxiv_id":"2305.07598","last_updated":"2026-05-05T04:25:55Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-05-12T16:42:54Z","title":"Hausdorff Distance Matching with Adaptive Query Denoising for Rotated Detection Transformer","version":6},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-24T08:32:21.680678Z"},"links":{"cited_paper":"/paper/2211.05778","citing_paper":"/paper/2305.07598"},"observation_digest":"sha256:37001e8576242fd4cb0a2145ef519f31a8e19ebf7ed199f24e7ee0f878bce25c","observation_id":"aa662122-8c47-4764-8188-c6bd474e69fb","resolution":{"observed_at":"2026-05-24T08:34:11.821516Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.05778","last_updated":"2023-04-17T11:51:12Z","snapshot_observed_at":"2026-08-11T17:05:48.093833Z","submitted_at":"2022-11-10T18:59:04Z","title":"InternImage: Exploring Large-Scale Vision Foundation Models with Deformable Convolutions","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.05778","snapshot_observed_at":"2026-08-12T12:43:47.568564Z","title":"Internimage: Exploring large-scale vi- sion foundation models with deformable convolutions.arXiv preprint arXiv:2211.05778, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2411.17027","last_updated":"2024-11-26T01:42:49Z","snapshot_observed_at":"2026-08-12T12:33:57.601279Z","submitted_at":"2024-11-26T01:42:49Z","title":"D$^2$-World: An Efficient World Model through Decoupled Dynamic Flow","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-12T12:43:47.568564Z"},"links":{"cited_paper":"/paper/2211.05778","citing_paper":"/paper/2411.17027"},"observation_digest":"sha256:00d5f25995996de6628ddb0ef71140a1930f007186d709d3c4ce060a4739317d","observation_id":"52a886f4-49f2-446a-81ba-9c9890f8b27e","resolution":{"observed_at":"2026-08-12T12:43:47.568564Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.05778","last_updated":"2023-04-17T11:51:12Z","snapshot_observed_at":"2026-08-11T17:05:48.093833Z","submitted_at":"2022-11-10T18:59:04Z","title":"InternImage: Exploring Large-Scale Vision Foundation Models with Deformable Convolutions","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.05778","snapshot_observed_at":"2026-08-09T18:36:48.331552Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.01675","last_updated":"2025-02-01T21:48:31Z","snapshot_observed_at":"2026-08-11T01:22:36.691340Z","submitted_at":"2025-02-01T21:48:31Z","title":"Semantic Communication based on Generative AI: A New Approach to Image Compression and Edge Optimization","version":1},"reference_index":116,"source":"arxiv_source","source_observed_at":"2026-08-09T18:36:48.331552Z"},"links":{"cited_paper":"/paper/2211.05778","citing_paper":"/paper/2502.01675"},"observation_digest":"sha256:6f2551a290e6db392952f1f492c26fecc7b97800e5864b88be7acc4d8e411b86","observation_id":"f42091e8-0b14-401a-a6a1-2a9a0da614e9","resolution":{"observed_at":"2026-08-09T18:36:48.331552Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.05778","last_updated":"2023-04-17T11:51:12Z","snapshot_observed_at":"2026-08-11T17:05:48.093833Z","submitted_at":"2022-11-10T18:59:04Z","title":"InternImage: Exploring Large-Scale Vision Foundation Models with Deformable Convolutions","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.05778","snapshot_observed_at":"2026-08-09T05:10:35.050125Z","title":"Internimage: Exploring large-scale vision foundation models with deformable convolutions","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.04377","last_updated":"2025-02-05T16:25:45Z","snapshot_observed_at":"2026-08-10T20:41:28.632169Z","submitted_at":"2025-02-05T16:25:45Z","title":"MapFusion: A Novel BEV Feature Fusion Network for Multi-modal Map Construction","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-09T05:10:35.050125Z"},"links":{"cited_paper":"/paper/2211.05778","citing_paper":"/paper/2502.04377"},"observation_digest":"sha256:1d6b7d87f8af899acaa3ba924686c995d2c570e2be90f7456d09ff2ce720b837","observation_id":"faf65b23-a8df-4ed4-8fd7-b4873109267e","resolution":{"observed_at":"2026-08-09T05:10:35.050125Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.05778","last_updated":"2023-04-17T11:51:12Z","snapshot_observed_at":"2026-08-11T17:05:48.093833Z","submitted_at":"2022-11-10T18:59:04Z","title":"InternImage: Exploring Large-Scale Vision Foundation Models with Deformable Convolutions","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.05778","snapshot_observed_at":"2026-08-06T18:25:20.693336Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.08384","last_updated":"2025-07-11T07:58:21Z","snapshot_observed_at":"2026-08-08T00:18:46.207107Z","submitted_at":"2025-07-11T07:58:21Z","title":"Smelly, dense, and spreaded: The Object Detection for Olfactory References (ODOR) dataset","version":1},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-06T18:25:20.693336Z"},"links":{"cited_paper":"/paper/2211.05778","citing_paper":"/paper/2507.08384"},"observation_digest":"sha256:c73486ca9d39416a5c12cf14877b91a89cf1adbe540acc2efe6e678925ff82bb","observation_id":"50c05e8d-c4f5-4e64-b432-2281af96cd29","resolution":{"observed_at":"2026-08-06T18:25:20.693336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.05778","last_updated":"2023-04-17T11:51:12Z","snapshot_observed_at":"2026-08-11T17:05:48.093833Z","submitted_at":"2022-11-10T18:59:04Z","title":"InternImage: Exploring Large-Scale Vision Foundation Models with Deformable Convolutions","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.05778","snapshot_observed_at":"2026-08-06T13:10:19.687685Z","title":"Internimage: Exploring large-scale vision foundation models with deformable convolutions,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.20963","last_updated":"2025-07-28T16:18:29Z","snapshot_observed_at":"2026-08-12T00:00:13.134750Z","submitted_at":"2025-07-28T16:18:29Z","title":"GTAD: Global Temporal Aggregation Denoising Learning for 3D Semantic Occupancy Prediction","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T13:10:19.687685Z"},"links":{"cited_paper":"/paper/2211.05778","citing_paper":"/paper/2507.20963"},"observation_digest":"sha256:8cc788b24b7f71e0f992a74b3d2c327048d9f3f4471daec8682643777780cc0c","observation_id":"16d200ed-55ca-44ae-9ac0-dd417168cd57","resolution":{"observed_at":"2026-08-06T13:10:19.687685Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.05778","last_updated":"2023-04-17T11:51:12Z","snapshot_observed_at":"2026-08-11T17:05:48.093833Z","submitted_at":"2022-11-10T18:59:04Z","title":"InternImage: Exploring Large-Scale Vision Foundation Models with Deformable Convolutions","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.05778","snapshot_observed_at":"2026-08-06T04:50:41.428702Z","title":"Internimage: Exploring large-scale vi- sion foundation models with deformable convolutions.arXiv preprint arXiv:2211.05778, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2508.02987","last_updated":"2025-08-05T01:31:10Z","snapshot_observed_at":"2026-08-08T02:57:23.719882Z","submitted_at":"2025-08-05T01:31:10Z","title":"Adversarial Attention Perturbations for Large Object Detection Transformers","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:41.428702Z"},"links":{"cited_paper":"/paper/2211.05778","citing_paper":"/paper/2508.02987"},"observation_digest":"sha256:d99977076637fd1ab08a56d96c8dcac07af5b094701ca5fb12017e99ff0aa6e0","observation_id":"55c4d454-7b08-46f4-8b37-487d1a93f090","resolution":{"observed_at":"2026-08-06T04:50:41.428702Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.05778","last_updated":"2023-04-17T11:51:12Z","snapshot_observed_at":"2026-08-11T17:05:48.093833Z","submitted_at":"2022-11-10T18:59:04Z","title":"InternImage: Exploring Large-Scale Vision Foundation Models with Deformable Convolutions","version":4},"cited_work":{"arxiv_id":"2211.05778","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2211.05778","snapshot_observed_at":"2026-06-28T19:32:35.285668Z","title":"Internim- age: Exploring large-scale vision foundation mod- els with deformable convolutions","venue":null,"work_id":"f6269124-1471-45b6-863b-d94f58db1a6b","year":2022},"citing_paper":{"arxiv_id":"2604.21119","last_updated":"2026-04-22T22:04:35Z","snapshot_observed_at":"2026-07-06T23:07:47.244208Z","submitted_at":"2026-04-22T22:04:35Z","title":"Materialistic RIR: Material Conditioned Realistic RIR Generation","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-05-10T00:01:22.336057Z"},"links":{"cited_paper":"/paper/2211.05778","citing_paper":"/paper/2604.21119"},"observation_digest":"sha256:86bfc426a3770b1401ba87a8b14c7e9853c33cc7a05ec2d9415edf9ee51588ae","observation_id":"fed3ebbb-2f6c-4008-80b6-01e6af1bc355","resolution":{"observed_at":"2026-05-11T13:51:02.855356Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.05778","last_updated":"2023-04-17T11:51:12Z","snapshot_observed_at":"2026-08-11T17:05:48.093833Z","submitted_at":"2022-11-10T18:59:04Z","title":"InternImage: Exploring Large-Scale Vision Foundation Models with Deformable Convolutions","version":4},"cited_work":{"arxiv_id":"2211.05778","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2211.05778","snapshot_observed_at":"2026-06-28T19:32:35.285668Z","title":"Internim- age: Exploring large-scale vision foundation mod- els with deformable convolutions","venue":null,"work_id":"f6269124-1471-45b6-863b-d94f58db1a6b","year":2022},"citing_paper":{"arxiv_id":"2606.00746","last_updated":"2026-05-30T14:29:43Z","snapshot_observed_at":"2026-07-06T23:41:24.911421Z","submitted_at":"2026-05-30T14:29:43Z","title":"Scaling Parallel Sequence Models to Foundation-Scale Vision Encoders","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-06-28T19:23:08.100056Z"},"links":{"cited_paper":"/paper/2211.05778","citing_paper":"/paper/2606.00746"},"observation_digest":"sha256:2845cf5b84dd6f2d07a053370be937141bd714115a2c1da03b9d2b34952b9414","observation_id":"374b012a-5cf0-47f4-8869-cc7684176c1d","resolution":{"observed_at":"2026-06-28T19:32:35.287136Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.05778","last_updated":"2023-04-17T11:51:12Z","snapshot_observed_at":"2026-08-11T17:05:48.093833Z","submitted_at":"2022-11-10T18:59:04Z","title":"InternImage: Exploring Large-Scale Vision Foundation Models with Deformable Convolutions","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.05778","snapshot_observed_at":"2026-08-01T04:56:31.547417Z","title":"doi:10.48550/arXiv.2211.05778 , note =","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.22408","last_updated":"2026-07-24T15:25:20Z","snapshot_observed_at":"2026-08-09T23:26:12.856482Z","submitted_at":"2026-07-24T15:25:20Z","title":"LunarFM: A Shared Multimodal Representation of the Moon's Surface","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-01T04:56:31.547417Z"},"links":{"cited_paper":"/paper/2211.05778","citing_paper":"/paper/2607.22408"},"observation_digest":"sha256:cf3035938088af84bd282487ee28ae503b6c62dc1cbdc968ecc5675d7b0811b2","observation_id":"3f0fa527-4fa6-4616-be63-874a524198ac","resolution":{"observed_at":"2026-08-01T04:56:31.547417Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.05778","last_updated":"2023-04-17T11:51:12Z","snapshot_observed_at":"2026-08-11T17:05:48.093833Z","submitted_at":"2022-11-10T18:59:04Z","title":"InternImage: Exploring Large-Scale Vision Foundation Models with Deformable Convolutions","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.05778","snapshot_observed_at":"2026-07-31T03:43:06.071039Z","title":"InternImage: Exploring large-scale vision foundation models with deformable convolutions,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.25011","last_updated":"2026-07-27T19:08:53Z","snapshot_observed_at":"2026-08-11T17:06:04.714344Z","submitted_at":"2026-07-27T19:08:53Z","title":"Optimization of Collaborative Semantic Communication Network Performance with Channel and Content Preference Feedback","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-07-31T03:43:06.071039Z"},"links":{"cited_paper":"/paper/2211.05778","citing_paper":"/paper/2607.25011"},"observation_digest":"sha256:ce9366c8e8523e2e8336940293d4dc48e9fbeed7f6dd6bc4590226c54d148965","observation_id":"e55bb1b7-4226-41eb-93d8-fc5b3db34d02","resolution":{"observed_at":"2026-07-31T03:43:06.071039Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2211.05778/citation-record","integrity":"/paper/2211.05778/integrity","json":"/paper/2211.05778/citation-record.json","paper":"/paper/2211.05778"},"outbound":[],"paper":{"arxiv_id":"2211.05778","last_updated":"2023-04-17T11:51:12Z","latest_version":4,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-11T17:05:48.093833Z","submitted_at":"2022-11-10T18:59:04Z","title":"InternImage: Exploring Large-Scale Vision Foundation Models with Deformable Convolutions"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 13 inbound Pith citation observations for arXiv:2211.05778."}