{"as_of":"2026-08-07T15:25:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:e15b9de711d557f7a390fb0ca841784494a4e29b7909e7f909b1936db6e6118d","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":19,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":19,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":19,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":19,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T01:06:18.216606Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":6,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":"2401.08541","doi":"10.48550/arxiv.2401.08541","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2401.08541 , year=","venue":"arXiv (Cornell University)","work_id":"c3ee0fe6-374c-40c8-889d-ca25c47cf1c3","year":2024},"citing_paper":{"arxiv_id":"2403.09611","last_updated":"2024-04-18T18:51:04Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-03-14T17:51:32Z","title":"MM1: Methods, Analysis & Insights from Multimodal LLM Pre-training","version":4},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-16T04:09:36.019146Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2403.09611"},"observation_digest":"sha256:9a6e84809934c1a8bed1c5d6a08a92e831d26ef3227193a1e3c29c246649a6b7","observation_id":"cc51c86e-8447-4528-87ac-473afe76094a","resolution":{"observed_at":"2026-05-16T04:09:36.398124Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":"2401.08541","doi":"10.48550/arxiv.2401.08541","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2401.08541 , year=","venue":"arXiv (Cornell University)","work_id":"c3ee0fe6-374c-40c8-889d-ca25c47cf1c3","year":2024},"citing_paper":{"arxiv_id":"2506.01844","last_updated":"2025-06-02T16:30:19Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-06-02T16:30:19Z","title":"SmolVLA: A Vision-Language-Action Model for Affordable and Efficient Robotics","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-11T21:22:36.902119Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2506.01844"},"observation_digest":"sha256:0cc1dd24e85db029f5c0d3763a06eae5215bfd7d6c640a3813453e9c0c59b15a","observation_id":"25897824-a30e-4eba-b7bb-6b61635c7873","resolution":{"observed_at":"2026-05-11T21:22:37.351022Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-07T01:06:18.216606Z","title":"El-Nouby, M","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.11967","last_updated":"2025-06-13T17:25:27Z","snapshot_observed_at":"2026-08-07T00:57:22.861281Z","submitted_at":"2025-06-13T17:25:27Z","title":"Visual Pre-Training on Unlabeled Images using Reinforcement Learning","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T01:06:18.216606Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2506.11967"},"observation_digest":"sha256:2d59118ed68720976e1155a3f5f8eb1112e80d4592916864d480911f22c5802d","observation_id":"cce1f160-3ff1-43ee-b35f-ccb84a9f3b8b","resolution":{"observed_at":"2026-08-07T01:06:18.216606Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-06T20:53:03.825159Z","title":"Scalable pre-training of large autoregressive image models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.01643","last_updated":"2025-07-02T12:17:23Z","snapshot_observed_at":"2026-08-06T20:44:05.101436Z","submitted_at":"2025-07-02T12:17:23Z","title":"SAILViT: Towards Robust and Generalizable Visual Backbones for MLLMs via Gradual Feature Refinement","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T20:53:03.825159Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2507.01643"},"observation_digest":"sha256:55dac2e81fda7379fd07e82fb3644c40d698c3cc803ef9c49b805b78bb3cc6bc","observation_id":"4d8ec47f-687d-4fea-ab49-a7e1a0a11765","resolution":{"observed_at":"2026-08-06T20:53:03.825159Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-06T20:38:53.354253Z","title":"Scalable pre- training of large autoregressive image models.arXiv preprint arXiv:2401.08541, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.02358","last_updated":"2025-07-11T09:06:39Z","snapshot_observed_at":"2026-08-06T21:34:24.738600Z","submitted_at":"2025-07-03T06:44:26Z","title":"Hita: Holistic Tokenizer for Autoregressive Image Generation","version":4},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T20:38:53.354253Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2507.02358"},"observation_digest":"sha256:049be0232710364eb169ca52de1ebca3d63e871e6f8945218cced61fbcd8a8ef","observation_id":"a87c79cf-73ab-4575-9598-d4760d41a977","resolution":{"observed_at":"2026-08-06T20:38:53.354253Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-05T12:22:32.780054Z","title":"Scalable pre- training of large autoregressive image models.arXiv preprint arXiv:2401.08541, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.01644","last_updated":"2025-09-01T17:38:21Z","snapshot_observed_at":"2026-08-05T12:22:30.859046Z","submitted_at":"2025-09-01T17:38:21Z","title":"OpenVision 2: A Family of Generative Pretrained Visual Encoders for Multimodal Learning","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-05T12:22:32.780054Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2509.01644"},"observation_digest":"sha256:ca3e2c186480478ab537bf5525727b689145b517ff5bf110b3e120586145238a","observation_id":"505638ea-f1f5-4a45-97f1-1463a38f024b","resolution":{"observed_at":"2026-08-05T12:22:32.780054Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":"2401.08541","doi":"10.48550/arxiv.2401.08541","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2401.08541 , year=","venue":"arXiv (Cornell University)","work_id":"c3ee0fe6-374c-40c8-889d-ca25c47cf1c3","year":2024},"citing_paper":{"arxiv_id":"2604.17961","last_updated":"2026-07-22T14:02:06Z","snapshot_observed_at":"2026-08-02T17:40:34.881723Z","submitted_at":"2026-04-20T08:41:30Z","title":"DifFoundMAD: Foundation Models meet Differential Morphing Attack Detection","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-10T05:31:36.180645Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2604.17961"},"observation_digest":"sha256:5e1106eb8ed421a575b3ab4807f1fd70a03fc016853022f74ed51b14f7cb43cf","observation_id":"028815b8-227f-4d79-ac1f-cbd50fab44f5","resolution":{"observed_at":"2026-05-10T05:36:02.098890Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-02T15:57:14.411668Z","title":"El-Nouby, M","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2604.17961","last_updated":"2026-07-22T14:02:06Z","snapshot_observed_at":"2026-08-02T17:40:34.881723Z","submitted_at":"2026-04-20T08:41:30Z","title":"DifFoundMAD: Foundation Models meet Differential Morphing Attack Detection","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-02T15:57:14.411668Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2604.17961"},"observation_digest":"sha256:a9d9bc514355b8c706255e83ece4fdca730d40b2ca7d67e85fffbbdec5ea80dc","observation_id":"520346e1-1991-4a4b-a34e-0ef4c86f31bb","resolution":{"observed_at":"2026-08-02T15:57:14.411668Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":"2401.08541","doi":"10.48550/arxiv.2401.08541","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2401.08541 , year=","venue":"arXiv (Cornell University)","work_id":"c3ee0fe6-374c-40c8-889d-ca25c47cf1c3","year":2024},"citing_paper":{"arxiv_id":"2605.08298","last_updated":"2026-05-08T11:09:30Z","snapshot_observed_at":"2026-08-02T05:07:12.715709Z","submitted_at":"2026-05-08T11:09:30Z","title":"What Cohort INRs Encode and Where to Freeze Them","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-12T01:32:38.878151Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2605.08298"},"observation_digest":"sha256:5ee2b05a531fcbda162ff6c3a6c18f189ad99690ccd64a77d369fd672208ebda","observation_id":"1d1d70fe-05f9-4f77-8685-52af7b180094","resolution":{"observed_at":"2026-05-12T07:56:26.363547Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":"2401.08541","doi":"10.48550/arxiv.2401.08541","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2401.08541 , year=","venue":"arXiv (Cornell University)","work_id":"c3ee0fe6-374c-40c8-889d-ca25c47cf1c3","year":2024},"citing_paper":{"arxiv_id":"2605.16384","last_updated":"2026-05-11T10:51:02Z","snapshot_observed_at":"2026-07-06T23:27:34.514541Z","submitted_at":"2026-05-11T10:51:02Z","title":"Mutual Enhancement Between Global Tokens and Patch Tokens: From Theory to Practice","version":1},"reference_index":91,"source":"arxiv_source","source_observed_at":"2026-05-20T22:41:44.510546Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2605.16384"},"observation_digest":"sha256:f6af787e713e7a7a38afc36477061a547c2dd161421d8118e12ecfeb63422920","observation_id":"bbe2c6da-c266-4685-8d4d-3287e7ca10df","resolution":{"observed_at":"2026-05-20T22:43:51.072841Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":"2401.08541","doi":"10.48550/arxiv.2401.08541","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2401.08541 , year=","venue":"arXiv (Cornell University)","work_id":"c3ee0fe6-374c-40c8-889d-ca25c47cf1c3","year":2024},"citing_paper":{"arxiv_id":"2605.17472","last_updated":"2026-05-20T04:42:46Z","snapshot_observed_at":"2026-07-06T23:28:26.397778Z","submitted_at":"2026-05-17T14:20:53Z","title":"Weighted Reverse Convolution for Feature Upsampling","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-20T14:24:28.728963Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2605.17472"},"observation_digest":"sha256:d7301473e59066be8092f021682b74a3739baf4753b786be98a9d9ab214ffb5b","observation_id":"3f8d4011-4e1d-4d6f-bcc8-43b69459d3a4","resolution":{"observed_at":"2026-05-20T14:28:21.683562Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":"2401.08541","doi":"10.48550/arxiv.2401.08541","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2401.08541 , year=","venue":"arXiv (Cornell University)","work_id":"c3ee0fe6-374c-40c8-889d-ca25c47cf1c3","year":2024},"citing_paper":{"arxiv_id":"2605.17472","last_updated":"2026-05-20T04:42:46Z","snapshot_observed_at":"2026-07-06T23:28:26.397778Z","submitted_at":"2026-05-17T14:20:53Z","title":"Weighted Reverse Convolution for Feature Upsampling","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-21T08:15:52.610554Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2605.17472"},"observation_digest":"sha256:10a53feae2dab988b5ded1b6e76c98c9cb77dcee186033dfc6a59ab31f69aae6","observation_id":"b111038b-75b4-4db8-82f5-9f4b1e6b2a22","resolution":{"observed_at":"2026-05-21T08:19:52.817998Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":"2401.08541","doi":"10.48550/arxiv.2401.08541","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2401.08541 , year=","venue":"arXiv (Cornell University)","work_id":"c3ee0fe6-374c-40c8-889d-ca25c47cf1c3","year":2024},"citing_paper":{"arxiv_id":"2605.23033","last_updated":"2026-05-21T20:58:42Z","snapshot_observed_at":"2026-07-06T23:33:19.375527Z","submitted_at":"2026-05-21T20:58:42Z","title":"Uncovering the Latent Potential of Deep Intermediate Representations","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-05-25T05:36:24.743558Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2605.23033"},"observation_digest":"sha256:d2a160774c04e0f41df857b6dc6559a1b5035c93eb0743c22ba9cd5580e896b7","observation_id":"58b74b70-2776-4b13-9694-c0a6fc467f3a","resolution":{"observed_at":"2026-05-25T05:36:39.257486Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":"2401.08541","doi":"10.48550/arxiv.2401.08541","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2401.08541 , year=","venue":"arXiv (Cornell University)","work_id":"c3ee0fe6-374c-40c8-889d-ca25c47cf1c3","year":2024},"citing_paper":{"arxiv_id":"2605.23992","last_updated":"2026-05-17T22:30:29Z","snapshot_observed_at":"2026-08-05T02:59:44.824431Z","submitted_at":"2026-05-17T22:30:29Z","title":"A World Model of Radiologist Reading for Medical Image Representation Learning","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-30T18:55:50.523140Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2605.23992"},"observation_digest":"sha256:436c7473a8408c516aaaafa188f3e21079d5dfcff3796a36a83bf97471c1a6fb","observation_id":"2545e4e5-890a-4a15-862c-1485f8a63296","resolution":{"observed_at":"2026-06-30T19:15:01.230775Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":"2401.08541","doi":"10.48550/arxiv.2401.08541","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2401.08541 , year=","venue":"arXiv (Cornell University)","work_id":"c3ee0fe6-374c-40c8-889d-ca25c47cf1c3","year":2024},"citing_paper":{"arxiv_id":"2606.21734","last_updated":"2026-06-19T20:43:49Z","snapshot_observed_at":"2026-08-05T18:05:51.515234Z","submitted_at":"2026-06-19T20:43:49Z","title":"HPP: Hierarchical Programmatic Probing for Long Video Understanding by Decoupling Perception and Reasoning","version":1},"reference_index":244,"source":"arxiv_source","source_observed_at":"2026-06-26T14:19:53.450263Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2606.21734"},"observation_digest":"sha256:4ab8772aa3a9fc438f817662644b90131fbd8f221e8e54c9d3c825352489d917","observation_id":"2d749fa7-129a-4666-9008-4868700605f6","resolution":{"observed_at":"2026-07-04T06:39:37.519117Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":"2401.08541","doi":"10.48550/arxiv.2401.08541","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2401.08541 , year=","venue":"arXiv (Cornell University)","work_id":"c3ee0fe6-374c-40c8-889d-ca25c47cf1c3","year":2024},"citing_paper":{"arxiv_id":"2606.25610","last_updated":"2026-06-24T09:17:17Z","snapshot_observed_at":"2026-08-02T23:05:09.045909Z","submitted_at":"2026-06-24T09:17:17Z","title":"The Galaxy's Guide to the Tokenizer: A Benchmark for Scientific Foundation Models","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-25T20:30:45.814418Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2606.25610"},"observation_digest":"sha256:b1c787d768ea235a29023ae655111ce4f74a6683d77de4b9dc1ece3371d8113e","observation_id":"a07b8580-0dfb-4a53-8ea3-122d735ad973","resolution":{"observed_at":"2026-06-25T20:38:19.798073Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":"2401.08541","doi":"10.48550/arxiv.2401.08541","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2401.08541 , year=","venue":"arXiv (Cornell University)","work_id":"c3ee0fe6-374c-40c8-889d-ca25c47cf1c3","year":2024},"citing_paper":{"arxiv_id":"2606.26794","last_updated":"2026-06-25T09:27:54Z","snapshot_observed_at":"2026-08-03T16:40:37.506200Z","submitted_at":"2026-06-25T09:27:54Z","title":"ReasonCLIP-58M: Visually Grounded Commonsense Reasoning Supervision for CLIP","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-06-26T05:03:15.044146Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2606.26794"},"observation_digest":"sha256:a1f50d8b79151ff37724950d3da93a67c0f21ac82872f5abcfe12276420688ff","observation_id":"d38ac3e9-eebc-4c52-b111-4f18ac8a2700","resolution":{"observed_at":"2026-07-04T13:39:50.719757Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T18:20:55.388664+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-07-12T03:24:13.814557Z","title":"https://doi.org/10.48550/arXiv.2401.08541, http://arxiv.org/abs/2401","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.03304","last_updated":"2026-07-03T13:16:30Z","snapshot_observed_at":"2026-08-06T16:38:09.160220Z","submitted_at":"2026-07-03T13:16:30Z","title":"Adaptive Loss Balancing for Multi-Task Bioacoustic Classification of Bird Species and Call Types","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-07-12T03:24:13.814557Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2607.03304"},"observation_digest":"sha256:09eb16cf862b28b9b30ed5f44abba46308c4b4a61cec29971461cd8f5f2027ab","observation_id":"6bf5396e-4475-40ac-b767-1f148fe582af","resolution":{"observed_at":"2026-07-12T03:24:13.814557Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.08541","snapshot_observed_at":"2026-08-04T04:31:48.707484Z","title":"Scalable Pre-training of Large Autoregressive Image Models.arXiv e-prints, page arXiv:2401.08541, January 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.02573","last_updated":"2026-08-03T17:49:40Z","snapshot_observed_at":"2026-08-06T23:39:11.550280Z","submitted_at":"2026-08-03T17:49:40Z","title":"Foundation Models for Astrophysics","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-04T04:31:48.707484Z"},"links":{"cited_paper":"/paper/2401.08541","citing_paper":"/paper/2608.02573"},"observation_digest":"sha256:953b584d6911c701d0ffcbf750ce43264819a2642b8ab9d20b6e61eedcfffc5b","observation_id":"9585a693-900a-4e5b-b603-895a16a09f84","resolution":{"observed_at":"2026-08-04T04:31:48.707484Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2401.08541/citation-record","integrity":"/paper/2401.08541/integrity","json":"/paper/2401.08541/citation-record.json","paper":"/paper/2401.08541"},"outbound":[],"paper":{"arxiv_id":"2401.08541","last_updated":"2024-01-16T18:03:37Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T17:16:17.240820Z","submitted_at":"2024-01-16T18:03:37Z","title":"Scalable Pre-training of Large Autoregressive Image Models"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 19 inbound Pith citation observations for arXiv:2401.08541."}