{"as_of":"2026-08-08T11:42:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:87e3010bb9a7bece8d76bfb9963e9c7ba79acec94532cdb9ae2d44bcd4f9d10c","coverage":[{"denominator":36,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":36,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T10:34:50.476017Z","state":"measured"},{"denominator":40,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":40,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":4,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":4,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-14T16:20:19.874172Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-16T10:17:44.669084Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"cited_work":{"arxiv_id":"2506.04997","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.04997","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Towards storage-efficient visual document retrieval: An empirical study on reducing patch-level embeddings.arXiv preprint arXiv:2506.04997, 2025a","venue":null,"work_id":"d4b6a762-b468-4ab0-b8be-002bf522ef16","year":2025},"citing_paper":{"arxiv_id":"2601.21262","last_updated":"2026-04-16T06:59:30Z","snapshot_observed_at":"2026-07-29T16:09:50.031781Z","submitted_at":"2026-01-29T04:47:27Z","title":"CausalEmbed: Auto-Regressive Multi-Vector Generation in Latent Space for Visual Document Embedding","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-16T10:14:15.589472Z"},"links":{"cited_paper":"/paper/2506.04997","citing_paper":"/paper/2601.21262"},"observation_digest":"sha256:dbaa9616c17672f42e43059fbba9a72e43d16c519ec5d632398bd8b6a2c82ca7","observation_id":"85970e61-20e3-4abb-9f6c-5d709d506f12","resolution":{"observed_at":"2026-05-16T10:17:44.671695Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"cited_work":{"arxiv_id":"2506.04997","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.04997","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Towards storage-efficient visual document retrieval: An empirical study on reducing patch-level embeddings.arXiv preprint arXiv:2506.04997, 2025a","venue":null,"work_id":"d4b6a762-b468-4ab0-b8be-002bf522ef16","year":2025},"citing_paper":{"arxiv_id":"2604.10167","last_updated":"2026-04-11T11:31:11Z","snapshot_observed_at":"2026-07-06T22:58:51.931478Z","submitted_at":"2026-04-11T11:31:11Z","title":"Visual Late Chunking: An Empirical Study of Contextual Chunking for Efficient Visual Document Retrieval","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-10T15:26:44.498777Z"},"links":{"cited_paper":"/paper/2506.04997","citing_paper":"/paper/2604.10167"},"observation_digest":"sha256:cf66aa4e2185cd14ab00f4d112e936827e262ba7a571a0a2bd3f3ab232832a87","observation_id":"5c566db8-93ae-4dd2-b054-73ab1c55c255","resolution":{"observed_at":"2026-05-11T10:31:04.261270Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.04997","snapshot_observed_at":"2026-07-11T16:32:55.757864Z","title":"arXiv preprint arXiv:2506.04997 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04605","last_updated":"2026-07-13T07:13:49Z","snapshot_observed_at":"2026-08-04T23:09:53.876810Z","submitted_at":"2026-07-06T02:19:11Z","title":"Do All Visual Tokens Matter Equally? Object-Evidence Preserving Token Merging for Vision-Language Retrieval","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-07-11T16:32:55.757864Z"},"links":{"cited_paper":"/paper/2506.04997","citing_paper":"/paper/2607.04605"},"observation_digest":"sha256:1e08fbc9b6eebf6e51d9238098666eb428aed7c34ac10d602bdcd0f2e530a790","observation_id":"3717ba08-d5b5-4052-b99b-43d26e915fa0","resolution":{"observed_at":"2026-07-11T16:32:55.757864Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.04997","snapshot_observed_at":"2026-07-14T16:20:19.874172Z","title":"arXiv preprint arXiv:2506.04997 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04605","last_updated":"2026-07-13T07:13:49Z","snapshot_observed_at":"2026-08-04T23:09:53.876810Z","submitted_at":"2026-07-06T02:19:11Z","title":"Do All Visual Tokens Matter Equally? Object-Evidence Preserving Token Merging for Vision-Language Retrieval","version":2},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-07-14T16:20:19.874172Z"},"links":{"cited_paper":"/paper/2506.04997","citing_paper":"/paper/2607.04605"},"observation_digest":"sha256:02632c326c9090d236591b1acced24125811dff3d802358d7665a03a9a709826","observation_id":"72881047-42d5-48b4-ae9d-5f032e5f9f77","resolution":{"observed_at":"2026-07-14T16:20:19.874172Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2506.04997/citation-record","integrity":"/paper/2506.04997/integrity","json":"/paper/2506.04997/citation-record.json","paper":"/paper/2506.04997"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2407.07726","last_updated":"2024-10-10T17:28:23Z","snapshot_observed_at":"2026-08-08T07:16:45.596308Z","submitted_at":"2024-07-10T14:57:46Z","title":"PaliGemma: A versatile 3B VLM for transfer","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.07726","snapshot_observed_at":"2026-08-07T10:34:47.967212Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:47.967212Z"},"links":{"cited_paper":"/paper/2407.07726","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:729aeb44953f5eb6efc87476c97ec4dcb6509407100185ab6d2fa1ffd297a493","observation_id":"c4335d3a-6921-4dc5-a871-086db87061c9","resolution":{"observed_at":"2026-08-07T10:34:47.967212Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:52.639548Z","title":null,"venue":null,"work_id":"b8815268-3893-432d-afaf-032edcfc2ac9","year":2023},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.022821Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:1d9980a82eb2ead664a831c24ba0386d09af18f975bf5a0fc3a60b0eb2568bec","observation_id":"514a11b6-ad70-408e-b90d-c9a9cc9db918","resolution":{"observed_at":"2026-08-07T10:34:52.679562Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:52.513183Z","title":"Rossi, Changyou Chen, and Tong Sun","venue":null,"work_id":"077f4943-598c-49b6-a75f-a5d807d0032b","year":2025},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.129493Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:b7eeb86c3e24224ffcf28a144edfc49e116e7c1748601effc3c7bc5d855b210f","observation_id":"973dc096-a943-40f6-ac5c-d5d9728186d4","resolution":{"observed_at":"2026-08-07T10:34:52.552822Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:52.403437Z","title":null,"venue":null,"work_id":"7b91589e-369f-4327-b239-77a293919f21","year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.214199Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:62a3ef275405a9ac03f549b0ee86f23b9aa0ec3e35a42ce2465d7b11400d1278","observation_id":"a987dba7-5de2-4296-aecc-049d4b327ee3","resolution":{"observed_at":"2026-08-07T10:34:52.469558Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.04952","last_updated":"2024-11-07T18:29:38Z","snapshot_observed_at":"2026-07-06T19:46:54.707852Z","submitted_at":"2024-11-07T18:29:38Z","title":"M3DocRAG: Multi-modal Retrieval is What You Need for Multi-page Multi-document Understanding","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.04952","snapshot_observed_at":"2026-08-07T10:34:48.298867Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.298867Z"},"links":{"cited_paper":"/paper/2411.04952","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:d0fb15d64ae4e594c007d3852f9d9d19281cee1e5fd288e5912a33b47bbb30d1","observation_id":"b5544092-a869-4a6b-9b09-20613d430e27","resolution":{"observed_at":"2026-08-07T10:34:48.298867Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.14683","last_updated":"2024-09-23T03:12:43Z","snapshot_observed_at":"2026-08-07T06:44:36.523865Z","submitted_at":"2024-09-23T03:12:43Z","title":"Reducing the Footprint of Multi-Vector Retrieval with Minimal Performance Impact via Token Pooling","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.14683","snapshot_observed_at":"2026-08-07T10:34:48.341670Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.341670Z"},"links":{"cited_paper":"/paper/2409.14683","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:9d696c5d7107c7318974910c997966d29dbe1177dfa7dfb9ade53208aa2874bf","observation_id":"3ff8675f-6b21-4411-b323-6c5dd1997472","resolution":{"observed_at":"2026-08-07T10:34:48.341670Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:52.252530Z","title":null,"venue":null,"work_id":"ed2f7d57-7af3-4bf7-916c-5aaedeb405f6","year":2025},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.398865Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:5a295c8d434b6f429dd295435c405af7733fe024c0a477c87fd30c771f25a544","observation_id":"4a81aebb-10e6-4c09-b7c2-81ddb657f1d5","resolution":{"observed_at":"2026-08-07T10:34:52.332510Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:48.469835Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.469835Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:c2008f064b29b73024d9b4482e2055207f1fb4fbcd1c782b42f7a593424a2098","observation_id":"4d6c3dcb-c2cb-46b4-9d59-6eb677129d3f","resolution":{"observed_at":"2026-08-07T10:34:48.469835Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:48.514132Z","title":null,"venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.514132Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:6df028c90f3cf8fb48ef1ac01319b4a03d3a4f74e7f9d8696aebd753f0abb75b","observation_id":"942c9fa1-6462-49b6-8772-85734335801f","resolution":{"observed_at":"2026-08-07T10:34:48.514132Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:48.591704Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.591704Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:f447d2cc7182777c956315f07a0ec935ed8a41c09987e40c592d0dc3ad6ea3a2","observation_id":"e6ba5de4-fad5-48ef-a921-e17af4e79d6a","resolution":{"observed_at":"2026-08-07T10:34:48.591704Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:48.660446Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.660446Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:b9e645a34fd03a608484f2d672b87e869b176ef42df784f7c212f120a42d242d","observation_id":"140cea8c-3e1b-40ec-9232-7fde73b63535","resolution":{"observed_at":"2026-08-07T10:34:48.660446Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:52.093963Z","title":null,"venue":null,"work_id":"894aefb7-b8d7-4f8d-b6d6-e91726966eed","year":2022},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.716111Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:2eb42daa6847597c607abf049f2638c8030d6d48f34ca91d673849a463389df2","observation_id":"b5b9740e-7b10-4378-9ce6-af4d37c1780b","resolution":{"observed_at":"2026-08-07T10:34:52.178094Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:48.819379Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.819379Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:7dee894367bd95bc6b33c489eeb5597479a0f12ab16417af43d90ebbabbafabf","observation_id":"fbf8466b-7f6b-45c8-b7c4-46c31086afc5","resolution":{"observed_at":"2026-08-07T10:34:48.819379Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.02392","last_updated":"2024-08-28T08:49:57Z","snapshot_observed_at":"2026-07-06T18:40:22.951929Z","submitted_at":"2024-07-02T16:10:55Z","title":"TokenPacker: Efficient Visual Projector for Multimodal LLM","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.02392","snapshot_observed_at":"2026-08-07T10:34:48.866029Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.866029Z"},"links":{"cited_paper":"/paper/2407.02392","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:46218405a5fd73ac2ae8da464708acf69c9e0fc985dd827ccb8dfc0e800ee5ee","observation_id":"f1413e60-66b2-4474-8cc5-4e90d4930ee2","resolution":{"observed_at":"2026-08-07T10:34:48.866029Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2202.07800","last_updated":"2022-04-13T23:46:08Z","snapshot_observed_at":"2026-08-03T03:15:29.208149Z","submitted_at":"2022-02-16T00:19:42Z","title":"Not All Patches are What You Need: Expediting Vision Transformers via Token Reorganizations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2202.07800","snapshot_observed_at":"2026-08-07T10:34:48.952540Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.952540Z"},"links":{"cited_paper":"/paper/2202.07800","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:56c6956f6ddb0221fc1fc71a5f8b7aa6a3ec0ad0d3045924af7d538bc643a361","observation_id":"fd4a96a3-e744-4317-88b6-1a3852736fd3","resolution":{"observed_at":"2026-08-07T10:34:48.952540Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:49.030889Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.030889Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:62ff28ead31856f55d82fc61e64072a12006ab0a8a66d61bc625b66811174f43","observation_id":"ba6c668c-259f-4a66-a4ca-b573e76dfbd1","resolution":{"observed_at":"2026-08-07T10:34:49.030889Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:51.903833Z","title":null,"venue":null,"work_id":"0f797af6-5b46-401d-9004-9a482f1e350d","year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.108336Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:17c0c75d417b5080d65ec41641ff41c80627fbafb770c52c558c0aff3d057599","observation_id":"a8c716eb-8f34-4fbe-b998-d63c1d26ea2b","resolution":{"observed_at":"2026-08-07T10:34:52.005884Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2203.10244","last_updated":"2022-03-19T05:00:30Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-03-19T05:00:30Z","title":"ChartQA: A Benchmark for Question Answering about Charts with Visual and Logical Reasoning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2203.10244","snapshot_observed_at":"2026-08-07T10:34:49.164741Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.164741Z"},"links":{"cited_paper":"/paper/2203.10244","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:83aa54296e99f1d15a10213e2253edae9409d0f7bb495067735e8d7c2c5f4e9a","observation_id":"1b0633fa-1f7a-47fe-b24d-34f1789bbd26","resolution":{"observed_at":"2026-08-07T10:34:49.164741Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:51.745639Z","title":null,"venue":null,"work_id":"017f2ecb-8834-4ead-8bb2-927d0aac5ca3","year":2021},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.231921Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:7949ef4739a184db845f83d8f79239f9e262e3559ec51cc84177c46d9674b528","observation_id":"b2d25baf-5494-41e2-b3a6-6f5ee645b6f7","resolution":{"observed_at":"2026-08-07T10:34:51.837276Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:51.528444Z","title":"Manmatha, and C","venue":null,"work_id":"4d6fcbd1-c2e5-42fc-b574-e9d2f677bd61","year":2020},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.301161Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:dfb17aaf3e48875c68b74957140fbfb4606778798c1f0ad1e6fb18af4f6793fe","observation_id":"58d024bb-4cf7-4fd8-9726-ccae9e18de51","resolution":{"observed_at":"2026-08-07T10:34:51.631527Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:51.302294Z","title":null,"venue":null,"work_id":"624a8c9c-7c07-4719-a616-067ed81aaa2a","year":2012},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.396840Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:75d4ee040c8514e1b883ca78310c8f10af5665a3e14618208086bd8f8c593d76","observation_id":"c4c44aed-6f8f-40c7-995c-540a58f88b35","resolution":{"observed_at":"2026-08-07T10:34:51.432829Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:49.478388Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.478388Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:9d38dda769385a7d6066b39513fa21874631088afcd5b7d66ae08db3ab538f10","observation_id":"6f29794e-0b4a-4dd2-8f1f-36d6e1d43ecc","resolution":{"observed_at":"2026-08-07T10:34:49.478388Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:49.536671Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.536671Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:1dc5521ab9bb43319a5599bef9034e0b01e4f591468b6c2104d290bc88201780","observation_id":"90de8716-48c5-4eb3-8f7b-035fa594ad38","resolution":{"observed_at":"2026-08-07T10:34:49.536671Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:51.118961Z","title":null,"venue":null,"work_id":"febdd1ba-ebff-43b9-8756-2319e8245dc7","year":2023},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.610618Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:eff44a4d8a50d0bd99c35cd89b057955522b4e719307694f81910aed95c2c748","observation_id":"36cfd537-d3b7-4f2a-9163-e58487aba55a","resolution":{"observed_at":"2026-08-07T10:34:51.235210Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12191","last_updated":"2024-10-03T15:54:49Z","snapshot_observed_at":"2026-08-06T05:35:29.109022Z","submitted_at":"2024-09-18T17:59:32Z","title":"Qwen2-VL: Enhancing Vision-Language Model's Perception of the World at Any Resolution","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12191","snapshot_observed_at":"2026-08-07T10:34:49.665368Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.665368Z"},"links":{"cited_paper":"/paper/2409.12191","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:a601e933bb48d52c35b02fdc8481bcd5025708709380fdd9a7e757784d94819c","observation_id":"27f5e534-644c-4f2a-875b-05c548853e35","resolution":{"observed_at":"2026-08-07T10:34:49.665368Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"status/1826238","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:50.829693Z","title":null,"venue":null,"work_id":"57c99a00-1716-4566-9458-0daca8b6fa88","year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.726150Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:d9cd16326dd2c4c7fed2cec008b3b47930536070c7cdd5d6e869fa7c69b04a00","observation_id":"a0814261-a18d-4dc3-9d68-ead305cccf77","resolution":{"observed_at":"2026-08-07T10:34:50.883764Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02686","last_updated":"2025-06-12T16:39:09Z","snapshot_observed_at":"2026-08-07T15:56:21.414821Z","submitted_at":"2025-05-05T14:33:49Z","title":"Sailing by the Stars: A Survey on Reward Models and Learning Strategies for Learning from Rewards","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.02686","snapshot_observed_at":"2026-08-07T10:34:49.797401Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.797401Z"},"links":{"cited_paper":"/paper/2505.02686","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:f3d28861a0d17a335129daf14c5b5e1e6f7997efa39c973cc018aa6dbee85a6a","observation_id":"ba0db544-a3dc-49a4-99ab-304f93204ed0","resolution":{"observed_at":"2026-08-07T10:34:49.797401Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:49.838297Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.838297Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:bfb4d7779c0699d7058ac0c1913b6d7d3b6f4666b6b228a1005366326efa92a3","observation_id":"4eb93546-8eaa-423c-a5cd-b4e8df3ee813","resolution":{"observed_at":"2026-08-07T10:34:49.838297Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.13670","last_updated":"2025-05-29T03:19:42Z","snapshot_observed_at":"2026-07-06T20:09:05.798010Z","submitted_at":"2024-12-18T09:53:12Z","title":"AntiLeakBench: Preventing Data Contamination by Automatically Constructing Benchmarks with Updated Real-World Knowledge","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.13670","snapshot_observed_at":"2026-08-07T10:34:49.916745Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.916745Z"},"links":{"cited_paper":"/paper/2412.13670","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:7ac354c48dd5385cd2a30e5bee88b63525eb18b56db177b4a1a44b7141958d32","observation_id":"914ecf84-8c46-4d30-b1a2-819367705d34","resolution":{"observed_at":"2026-08-07T10:34:49.916745Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.17247","last_updated":"2025-02-27T11:16:33Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-10-22T17:59:53Z","title":"PyramidDrop: Accelerating Your Large Vision-Language Models via Pyramid Visual Redundancy Reduction","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.17247","snapshot_observed_at":"2026-08-07T10:34:49.994840Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.994840Z"},"links":{"cited_paper":"/paper/2410.17247","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:84a539fae2515c04fc236c5eabbbee6c412e677a16b3648bbc01047fe18ac97e","observation_id":"629ba998-3c0a-4abc-9d38-c6eb14cba1c7","resolution":{"observed_at":"2026-08-07T10:34:49.994840Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:50.085393Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:50.085393Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:017126d48a54217f556d18892fab7f9059e070b0512a72efcd0ba3923333a25e","observation_id":"905b0ba9-1dd5-4e70-b9d2-c604e5930e05","resolution":{"observed_at":"2026-08-07T10:34:50.085393Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.10594","last_updated":"2025-03-02T01:19:51Z","snapshot_observed_at":"2026-08-02T13:24:00.339129Z","submitted_at":"2024-10-14T15:04:18Z","title":"VisRAG: Vision-based Retrieval-augmented Generation on Multi-modality Documents","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.10594","snapshot_observed_at":"2026-08-07T10:34:50.174984Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:50.174984Z"},"links":{"cited_paper":"/paper/2410.10594","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:8da53dc472105010dda428286f011621e224a733ada5a44b0f8262fc9c9a5076","observation_id":"c44d13d5-7738-4f9a-86ef-9a5e82d9b65f","resolution":{"observed_at":"2026-08-07T10:34:50.174984Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.04417","last_updated":"2025-06-03T04:12:10Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-10-06T09:18:04Z","title":"SparseVLM: Visual Token Sparsification for Efficient Vision-Language Model Inference","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.04417","snapshot_observed_at":"2026-08-07T10:34:50.294599Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:50.294599Z"},"links":{"cited_paper":"/paper/2410.04417","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:6c8d43abb554e59c0bbeffb8d9bf4150517db1e7cff2009c475d2a217ac4f9c6","observation_id":"89a7c29f-c639-4e22-ae57-216d82d7ed2d","resolution":{"observed_at":"2026-08-07T10:34:50.294599Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:50.374302Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:50.374302Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:ea6b30336800ac735f598ae03c4278e523c2fff31c060cc2c030c7dabb1043a3","observation_id":"ce8dcb9c-3b8f-4db1-a1cc-8ffda0837af3","resolution":{"observed_at":"2026-08-07T10:34:50.374302Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:50.415916Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:50.415916Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:4455539b86bfa800e5e966c564fa889514006f78111d60c9d4e9968e01d70d4a","observation_id":"9cf2edb2-42f3-42e0-bfad-e99caa105cfc","resolution":{"observed_at":"2026-08-07T10:34:50.415916Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:50.476017Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:50.476017Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:2333f7c26b9bc4815b5fcc806685ed85d3fe2a797b3f7a313f6946a55de768b2","observation_id":"5f679184-4f66-4d4c-9347-319b15eb8377","resolution":{"observed_at":"2026-08-07T10:34:50.476017Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","latest_version":1,"primary_category":"cs.IR","snapshot_observed_at":"2026-08-07T10:26:27.903134Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings"},"reference_resolution":{"displayed":36,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":33,"verified_exact":1,"verified_fuzzy":2},"total_outbound_references":36},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 36 of 36 outbound references and 4 inbound Pith citation observations for arXiv:2506.04997."}