{"as_of":"2026-08-10T16:17:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:e0649ef12b0ae439c4883a366252e5270774521ff2bf8fa12f9dc0a90105ad69","coverage":[{"denominator":43,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":43,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-07T19:32:20.095114Z","state":"measured"},{"denominator":43,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":43,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.05290/citation-record","integrity":"/paper/2607.05290/integrity","json":"/paper/2607.05290/citation-record.json","paper":"/paper/2607.05290"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.776717Z","title":"The Claude model family.Anthropic Technical Report, 2024","venue":null,"work_id":"27a0890f-1b63-4c9d-916d-d389bb18c2cd","year":2024},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:84906606c14a8b34e880854d3fd4870d24d4187126d8007a94fcff638c329b03","observation_id":"4ec586dc-a98e-407e-8b02-0a882ccabf28","resolution":{"observed_at":"2026-07-07T19:34:06.779307Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.729620Z","title":"VQA: Visual question answering","venue":null,"work_id":"e915d80b-eb75-4a53-a391-3e025d61a8bd","year":2015},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:ddde530b7ff4aa17571eff78af9cc36e393c5f8408cad46f079958b105de2093","observation_id":"b0018f56-a31a-4978-bc4e-5336c9377568","resolution":{"observed_at":"2026-07-07T19:34:06.732264Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.13391","last_updated":"2024-12-17T15:12:02Z","snapshot_observed_at":"2026-07-06T17:19:52.125879Z","submitted_at":"2024-01-24T11:41:30Z","title":"Reranking individuals: The effect of fair classification within-groups","version":3},"cited_work":{"arxiv_id":"2401.13391","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.13391","snapshot_observed_at":"2026-07-07T19:34:06.396887Z","title":"Reranking individuals: The effect of fair classification within-groups","venue":"cs.LG","work_id":"7792d451-8eb7-4678-b331-3b9b4c4e03de","year":2024},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"cited_paper":"/paper/2401.13391","citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:ea0738c8edf24a22eda3bf2f7c82f1e16067cdcfadd927af0d9622211f1ebc77","observation_id":"74c70bbd-67b6-4a25-8811-8b8f660bb9f1","resolution":{"observed_at":"2026-07-07T19:34:06.399306Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.746100Z","title":"InternVL: Scaling up vision foundation models and aligning for generic visual-linguistic tasks","venue":null,"work_id":"124afdf4-5f24-4e09-a4d6-a7120817ddef","year":2024},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:f304f4d5fdba49c6440c97758c70b12c4a5873930aa4dc1a016b0296f02ca809","observation_id":"2844780f-372a-4a63-8a13-5ce284880f8d","resolution":{"observed_at":"2026-07-07T19:34:06.748766Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.713583Z","title":"YOLO-World: Real-time open-vocabulary object detection.IEEE/CVF Conference on Computer Vision and Pattern Recognition, 2024","venue":null,"work_id":"9578df52-6b0a-4c52-80be-db563543f26a","year":2024},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:410eda861a3a221afe5560576cee0b2b28374aa608d30aef18aaabeea02c6b05","observation_id":"d6d0c24e-b42a-429b-a272-b9d3487a24dd","resolution":{"observed_at":"2026-07-07T19:34:06.716095Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.14390","last_updated":"2024-05-29T00:27:52Z","snapshot_observed_at":"2026-07-06T17:06:54.562888Z","submitted_at":"2023-12-22T02:34:56Z","title":"Concatenating Binomial Codes with the Planar Code","version":2},"cited_work":{"arxiv_id":"2312.14390","doi":null,"metadata_source":"pith","pith_arxiv_id":"2312.14390","snapshot_observed_at":"2026-07-07T19:34:06.378325Z","title":"Concatenating Binomial Codes with the Planar Code","venue":"quant-ph","work_id":"75ccf66f-5334-439e-86a9-94e06a5aabfa","year":2023},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"cited_paper":"/paper/2312.14390","citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:e50e50a81b9ea5df3205da5b1d585a9d5bd26289f1a06afa6f8e48e1e87ead70","observation_id":"0f126077-7b68-48e2-b686-125d20680f78","resolution":{"observed_at":"2026-07-07T19:34:06.380852Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.807613Z","title":"TransVG: Visual grounding with transformers","venue":null,"work_id":"10cb4f55-e8c4-4f0f-84f9-ec5766e82237","year":2021},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:9c671fa2c80f85078a2071747dc4265f643aed67093a64c803e4ee4f179b6c5a","observation_id":"98bbaa77-b24a-4e55-a618-50cebcf4dcec","resolution":{"observed_at":"2026-07-07T19:34:06.810593Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.816160Z","title":"LayoutGPT: Compositional visual planning and generation with large language models","venue":null,"work_id":"5ffabf42-3dbf-4403-8dee-74c8b928d39f","year":2023},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:ef6d3da45a83aae9f3df05d8cd2f606c071cd71ac4108b3a88812688439b9ae3","observation_id":"00d0de1c-f87d-4852-80bc-71594c2d534f","resolution":{"observed_at":"2026-07-07T19:34:06.818944Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.10997","last_updated":"2024-03-27T09:16:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-18T07:47:33Z","title":"Retrieval-Augmented Generation for Large Language Models: A Survey","version":5},"cited_work":{"arxiv_id":"2312.10997","doi":"10.1186/1476-072x-8-72","metadata_source":"pith","pith_arxiv_id":"2312.10997","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Retrieval-Augmented Generation for Large Language Models: A Survey","venue":"cs.CL","work_id":"b80d2790-6cd9-4c87-b3c4-de404f99a80e","year":2023},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"cited_paper":"/paper/2312.10997","citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:ff71d54e81542c6f225664df525e10621941795eb4c97547845f4fa1c0cbd315","observation_id":"0d620a3b-6552-4ca1-9cbb-90269bce75c7","resolution":{"observed_at":"2026-07-07T19:34:06.371018Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1612.00837","last_updated":"2017-05-15T17:58:49Z","snapshot_observed_at":"2026-07-06T05:21:10.284182Z","submitted_at":"2016-12-02T20:57:07Z","title":"Making the V in VQA Matter: Elevating the Role of Image Understanding in Visual Question Answering","version":3},"cited_work":{"arxiv_id":"1612.00837","doi":"10.48550/arxiv.1612.00837","metadata_source":"pith","pith_arxiv_id":"1612.00837","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Making the V in VQA Matter: Elevating the Role of Image Understanding in Visual Question Answering","venue":"cs.CV","work_id":"5224ce0f-65a2-404a-8b9d-27604f17acb9","year":2016},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"cited_paper":"/paper/1612.00837","citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:ce8e8cecc09855ea1d1b35ae71f5ec3ac5c8b0761e1b86b1dc7b660e4f1c9ae9","observation_id":"a32ba699-9e1a-4bd8-9cbc-a51f9eeb0929","resolution":{"observed_at":"2026-07-07T19:34:06.385691Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.767410Z","title":"Interactive poster visualization with PhD Online.IEEE Computer Graphics and Applications, 2004","venue":null,"work_id":"986c3eef-2d9c-44c7-9a29-de4d2e23e1d3","year":2004},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:31390ce712f7fe3c620792ecefd3b362c9c8af35a090e98ce92e70e8dd7e3943","observation_id":"914006fb-a983-4a99-9c14-e51baaa1631e","resolution":{"observed_at":"2026-07-07T19:34:06.770093Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10988","last_updated":"2024-02-16T10:56:45Z","snapshot_observed_at":"2026-07-06T17:31:21.158599Z","submitted_at":"2024-02-16T10:56:45Z","title":"Cryptography: Classical versus Post-Quantum","version":1},"cited_work":{"arxiv_id":"2402.10988","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.10988","snapshot_observed_at":"2026-07-07T19:34:06.363240Z","title":"Cryptography: Classical versus Post-Quantum","venue":"quant-ph","work_id":"ee9d68e5-dd1d-4f90-85b4-4be4e6443596","year":2024},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"cited_paper":"/paper/2402.10988","citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:4107f7689f80cb1563c53e76aac845624464d8011e5d92b9750a8eeef91aa486","observation_id":"a26503f0-6a60-4938-9956-752cc374522d","resolution":{"observed_at":"2026-07-07T19:34:06.365987Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2605.18636","last_updated":"2026-05-18T16:43:32Z","snapshot_observed_at":"2026-08-02T15:45:34.247215Z","submitted_at":"2026-05-18T16:43:32Z","title":"SPIKE: An Adaptive Dual Controller Framework for Cost-Efficient Long-Horizon Game Agents","version":1},"cited_work":{"arxiv_id":"2605.18636","doi":null,"metadata_source":"pith","pith_arxiv_id":"2605.18636","snapshot_observed_at":"2026-07-07T19:34:06.357620Z","title":"SPIKE: An Adaptive Dual Controller Framework for Cost-Efficient Long-Horizon Game Agents","venue":"cs.CV","work_id":"fcc98409-771b-467e-acf1-6fe07dfd5e01","year":2026},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"cited_paper":"/paper/2605.18636","citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:3d1ef079a4073f932fc6031c33a6d8b98e3ab607629cc2d0919891284b78e8da","observation_id":"d11de996-ca9b-49b0-85ce-2a462c0a5d13","resolution":{"observed_at":"2026-07-07T19:34:06.361002Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.724756Z","title":"MDETR: Modulated detection for end-to-end multi-modal understanding","venue":null,"work_id":"a4876d2b-b8e6-4d90-a41a-d610817bb346","year":2021},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:7cacfbdf41da78ca7b3d30bdc0aa9c55ad5a3c4d467dbaadb2638c9d9f30b5bd","observation_id":"657bf2bf-32a4-4e30-b94a-da0b94e6c0ba","resolution":{"observed_at":"2026-07-07T19:34:06.727764Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.737183Z","title":"Berg, Wan-Yen Lo, Piotr Dollár, and Ross Girshick","venue":null,"work_id":"a430edec-a4a8-4073-b02d-f6d95ad7b762","year":2023},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:43cb2e6cd7c9a033c2c7ee0b924d6733171d9e383b433389b80be8cdc1c61d66","observation_id":"d351c8ea-411e-46a5-b25b-c23a42be83f6","resolution":{"observed_at":"2026-07-07T19:34:06.740141Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.750157Z","title":"Retrieval-augmented generation for knowledge-intensive NLP tasks","venue":null,"work_id":"096fa0ee-49c1-4c43-a972-2ef87bdfa876","year":2020},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:4eb61cbc214788e43b493e1cc45098f1f90c99f4ed4459bc773819659fd67710","observation_id":"7c916d13-74a2-4c07-97cc-f775f64bb773","resolution":{"observed_at":"2026-07-07T19:34:06.753043Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.780606Z","title":"GLIGEN: Open-set grounded text-to-image generation","venue":null,"work_id":"6cd9602b-2624-44b7-8e5f-2385e09ef199","year":2023},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:67c385a55ca04b5c5712ff88fcce58bfd331f425a554d497a15fa6f2c3bb5215","observation_id":"7deb228a-bc66-448d-b927-6b8f07e9a303","resolution":{"observed_at":"2026-07-07T19:34:06.783486Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.763352Z","title":"Visual instruction tuning","venue":null,"work_id":"4af234cb-142f-4b22-8c5e-5d0b88357cc8","year":2023},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:ee56bef3adf2e37d3b2eb0d1623be47b5685cf04d484a85e401d9d6ddd94ff3f","observation_id":"635065f5-2d11-4490-828c-f8b63c9d8d19","resolution":{"observed_at":"2026-07-07T19:34:06.765998Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.758338Z","title":"Grounding DINO: Marrying DINO with grounded pre-training for open-set object detection","venue":null,"work_id":"96921ddc-94d9-4cf0-b85c-2d8d19144891","year":2024},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:67c5690f3cb2a14066dc89f9e177b656b9c866591df17a86420634bbc45b9ab4","observation_id":"db309a84-39d5-4e0c-a421-064e9c7db105","resolution":{"observed_at":"2026-07-07T19:34:06.761850Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13632","last_updated":"2023-10-20T16:31:06Z","snapshot_observed_at":"2026-08-03T11:55:34.023021Z","submitted_at":"2023-10-20T16:31:06Z","title":"Counting Divisors in the Outputs of a Binary Quadratic Form","version":1},"cited_work":{"arxiv_id":"2310.13632","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.13632","snapshot_observed_at":"2026-07-07T19:34:06.387830Z","title":"Counting Divisors in the Outputs of a Binary Quadratic Form","venue":"math.NT","work_id":"854bb7b3-c874-4741-8921-c0cc4ff104cd","year":2023},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"cited_paper":"/paper/2310.13632","citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:715400e830c18494d5eefb363f5ef07aa1e571552abae0ce1ebaaf4639c3f548","observation_id":"3e93ea9d-1650-4983-a6cb-50cbef1d3f90","resolution":{"observed_at":"2026-07-07T19:34:06.390078Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.794224Z","title":"ChartQA: A benchmark for question answering about charts with visual and logical reasoning","venue":null,"work_id":"4b29f930-4dbf-4c26-8c9c-fbea7e17a3bc","year":2022},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:3a83458709af20523dc85f1875cc611fdc3271db24badfabba963fb95d94725b","observation_id":"89ac7dbf-e0ff-4565-bce4-e53692618d26","resolution":{"observed_at":"2026-07-07T19:34:06.797214Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.717390Z","title":null,"venue":null,"work_id":"ad4ae6f1-eadf-482c-9190-c5852f545168","year":2021},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:f50f2c17dc5f5117f8d4d846e4659e0d50d4d3baf2b69407d12b76043cb1cc14","observation_id":"84a7feb8-b4fb-45f7-aa7c-dfff2a0fc1f6","resolution":{"observed_at":"2026-07-07T19:34:06.719372Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.798778Z","title":"CRC Press, 2014","venue":null,"work_id":"ae22c733-77cd-43e1-82b1-d5d7f1f7cd6f","year":2014},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:34d181dcf2d5cb0d85abbff6961ca970b46d1eab403b4b1914595e38444112ae","observation_id":"41115014-6cde-4e02-8c22-87a451d6c0cd","resolution":{"observed_at":"2026-07-07T19:34:06.801627Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.789549Z","title":"LocateAnything-3B: A visual grounding model","venue":null,"work_id":"f65c1c52-15e6-453a-8b75-64bba83949a1","year":null},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:28ab90f245e49783d886b890bccfccd4dd1aa8ffcc2fb145d64e962e742e8ba4","observation_id":"ad32b0a2-54c3-4115-9f4a-478a99c914f5","resolution":{"observed_at":"2026-07-07T19:34:06.792756Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.720648Z","title":"Chart-to-Text: Generating textual descriptions of charts","venue":null,"work_id":"580705be-25dd-47ab-bd9a-f65a1a547710","year":2020},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:6185ba18ffea10e3d445ded09064c879f89c2e53c18c699d0d41d5b7ad44bcd7","observation_id":"245c8dd1-b9cd-4d2f-a528-859a2fe81e26","resolution":{"observed_at":"2026-07-07T19:34:06.722953Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":"2303.08774","doi":"10.1002/tea.20265","metadata_source":"pith","pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"GPT-4 Technical Report","venue":"cs.CL","work_id":"b928e041-6991-4c08-8c81-0359e4097c7b","year":2023},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:bf8907a3150acea00d6d17d11e89c7ad67c002c1b728fa13c87e33ea2e6eb16c","observation_id":"1dbb3275-0f81-44e1-bc40-53e06ce96e68","resolution":{"observed_at":"2026-07-07T19:34:06.407051Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.784916Z","title":"SDXL: Improving latent diffusion models for high-resolution image synthesis","venue":null,"work_id":"a3bed837-1afd-4e1d-a621-cb7a9436b0ce","year":2024},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:c7941eee78a687a470c1e54b8d54ba16e01d7c82e6c3263e47b9fae5f5c51bab","observation_id":"0d4e3383-3bc1-4bf8-bb8d-24c8b2a42a91","resolution":{"observed_at":"2026-07-07T19:34:06.788107Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2204.06125","last_updated":"2022-04-13T01:10:33Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-04-13T01:10:33Z","title":"Hierarchical Text-Conditional Image Generation with CLIP Latents","version":1},"cited_work":{"arxiv_id":"2204.06125","doi":"10.48550/arxiv.2204.06125","metadata_source":"pith","pith_arxiv_id":"2204.06125","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Hierarchical Text-Conditional Image Generation with CLIP Latents","venue":"cs.CV","work_id":"0c6a768b-70b8-4242-bb0e-459f1008c9fc","year":2022},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"cited_paper":"/paper/2204.06125","citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:6abd224e5af0ef880d9ace1053c8cb518aeb8827274b8194a2305cf439f76dd5","observation_id":"bfcd80d0-2678-431f-be8f-8154e0697eac","resolution":{"observed_at":"2026-07-07T19:34:06.376190Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-14T13:50:02.412412+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-14T13:50:02.412412+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.00714","last_updated":"2024-10-28T16:37:57Z","snapshot_observed_at":"2026-07-06T18:55:41.459417Z","submitted_at":"2024-08-01T17:00:08Z","title":"SAM 2: Segment Anything in Images and Videos","version":2},"cited_work":{"arxiv_id":"2408.00714","doi":"10.1038/s41598-025-97590-3","metadata_source":"pith","pith_arxiv_id":"2408.00714","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 2: Segment Anything in Images and Videos","venue":"cs.CV","work_id":"acc13f66-d814-44f9-9688-375688bf2d4a","year":2024},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"cited_paper":"/paper/2408.00714","citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:1f22817a15b839e0a78106cf4e6c5c31c6c47e1c398f2d5dc391df894314d77d","observation_id":"b6deefb4-f147-4366-88b9-bb440bdabbd3","resolution":{"observed_at":"2026-07-07T19:34:06.403344Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-05-24T04:24:23.885301+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-24T04:24:23.885301+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T09:44:49.718281Z","title":"High-resolution image synthesis with latent diffusion models","venue":null,"work_id":"02478f33-0764-47e3-b50f-24dcbc5a5bc0","year":2022},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:8c272dfb5e845f882ced82d630db441885e8be9076f3c53e448a62a316311e84","observation_id":"9ae6322d-6b31-45ca-a3e3-16160ecac009","resolution":{"observed_at":"2026-07-07T19:34:06.744279Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.754442Z","title":"Photorealistic text-to-image diffusion models with deep language understanding","venue":null,"work_id":"131f9d64-cdb5-401a-ac05-71e506ab44e7","year":2022},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:983bacb3c52fa98001d420075fb55184d82eaebd9285c86bdcc122a446868228","observation_id":"bffa7a5d-dd25-4cfb-9221-7c4287ff7ffb","resolution":{"observed_at":"2026-07-07T19:34:06.756944Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.771529Z","title":"Toolformer: Language models can teach themselves to use tools","venue":null,"work_id":"29101b40-6d04-4188-90c0-1fc99be52a69","year":2023},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:2fca147798db78db0d645429de3c107c5365e5f3bf42a5326c0bc4b7bc86bba1","observation_id":"f012e20a-318f-4bd2-92f2-832f442f3e0b","resolution":{"observed_at":"2026-07-07T19:34:06.775142Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.16483","last_updated":"2023-11-27T15:20:23Z","snapshot_observed_at":"2026-08-06T14:11:19.997270Z","submitted_at":"2023-11-27T15:20:23Z","title":"ChartLlama: A Multimodal LLM for Chart Understanding and Generation","version":1},"cited_work":{"arxiv_id":"2311.16483","doi":null,"metadata_source":"pith","pith_arxiv_id":"2311.16483","snapshot_observed_at":"2026-07-07T19:34:06.392258Z","title":"Chartllama: A mul- timodal llm for chart understanding and generation","venue":"cs.CV","work_id":"f6c4f1ff-0ce7-47e8-82e5-5c24c2a3ed9f","year":2023},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"cited_paper":"/paper/2311.16483","citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:7b2639c743886a2fd64581630cc9e625363c1772bc73cd139248a4de4a113fe1","observation_id":"004a7901-3849-47e3-a941-39559dafdccc","resolution":{"observed_at":"2026-07-07T19:34:06.394815Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.709213Z","title":"Chain-of-thought prompting elicits reasoning in large language models","venue":null,"work_id":"1138e3c1-c843-4fc3-a15c-359dae88b20e","year":2022},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:48cef2d324bbf12c204c5bc7f49ac9debca7dcbb32b1f2a4b3d49518a097bb0f","observation_id":"55d0d7a1-2263-4b36-a549-b0a9b68e33af","resolution":{"observed_at":"2026-07-07T19:34:06.711848Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.824525Z","title":"MiMo-VL: Xiaomi MiMo vision-language model","venue":null,"work_id":"264fba15-bc43-480c-b163-95a4c2f10861","year":null},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:4971e6b0bfd39c88ba25dfeffd555845445630d6d05e136e21729a71aa59515b","observation_id":"2833e145-7e42-4507-9ec6-3a3cd5013d23","resolution":{"observed_at":"2026-07-07T19:34:06.827556Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.802791Z","title":null,"venue":null,"work_id":"56bab32c-8f73-4216-8982-767db37fe242","year":null},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:d87a264e00a042017967c52a4f76fe768f21ad0b7c86dd7fb224e261dcc0beea","observation_id":"ea4883c0-4c75-457e-9647-11178c97a973","resolution":{"observed_at":"2026-07-07T19:34:06.805860Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.811850Z","title":"LLM-oriented token- adaptive knowledge distillation","venue":null,"work_id":"e4535460-7662-4de5-ba99-2a2d8fc47274","year":null},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:845324743498b144faaa9404c5e57d3bc7c52330ccd48dbf83973334bd47873d","observation_id":"2f96abde-6127-4493-9693-0841d3e0fc3e","resolution":{"observed_at":"2026-07-07T19:34:06.814889Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1609/aaai.v40i16.40701","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.339495Z","title":null,"venue":null,"work_id":"225eb3e2-4dec-4376-86f3-d010332d103b","year":null},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:97711080cd11f71ee851fb4d3943e1108cee53ced109a8693b33dc563ef5e39b","observation_id":"d056bd35-82c6-4119-9796-6b17bf0a58f1","resolution":{"observed_at":"2026-07-07T19:34:06.342685Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-11T17:49:03.2912+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T17:49:03.2912+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2603.24690","last_updated":"2026-07-07T02:44:32Z","snapshot_observed_at":"2026-07-13T18:42:32.163637Z","submitted_at":"2026-03-25T18:09:33Z","title":"UniICL: Systematizing Unified Multimodal In-context Learning through a Capability-Oriented Taxonomy","version":3},"cited_work":{"arxiv_id":"2603.24690","doi":null,"metadata_source":"pith","pith_arxiv_id":"2603.24690","snapshot_observed_at":"2026-07-07T19:34:06.409175Z","title":"UniICL: Systematizing Unified Multimodal In-context Learning through a Capability-Oriented Taxonomy","venue":"cs.CV","work_id":"6624e9b9-9e22-4b9f-b4bb-72cf2d1f5517","year":2026},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"cited_paper":"/paper/2603.24690","citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:eb47d38d1807cd76aeba6cce8f6b97a26ed3e24ebe4db8dbdc51eef7d701814b","observation_id":"c075f7a1-ad10-4891-9855-286a7582edd3","resolution":{"observed_at":"2026-07-07T19:34:06.411651Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.733665Z","title":"UNINEXT: Universal instance perception as object-in-context prompting","venue":null,"work_id":"40c4996d-9c0d-47a3-b4de-74467fcae0cd","year":2023},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:5faa1f54e16e18f7487258b5fe7dab0583af89c38281565bd9ce648c3db6c7e5","observation_id":"3e19e448-1012-4b29-81f5-5cae83ef75e9","resolution":{"observed_at":"2026-07-07T19:34:06.735898Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T05:56:02.447528Z","title":"ReAct: Synergizing reasoning and acting in language models","venue":null,"work_id":"1ae5dc39-fc24-4009-b873-2b12af346d25","year":2023},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:cec99655e4011e6dd6a442d7b603d1f127305ff67045fb1f1cdfeb5e4aec2016","observation_id":"a5ec72a0-b7b9-4e1c-a691-dd2e57fae546","resolution":{"observed_at":"2026-07-07T19:34:06.705003Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.820847Z","title":"Adding conditional control to text-to-image diffusion models","venue":null,"work_id":"e70bfce6-599d-4e4c-b2eb-560528bcc0f5","year":2023},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:521c0096acbff19b0c32c2d8b469ef570b016462819d3d50786291cd832ce0eb","observation_id":"f55b13b3-dcd7-4d4a-9867-112d8cefcbba","resolution":{"observed_at":"2026-07-07T19:34:06.823384Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T19:34:06.706227Z","title":"Mindmap: A creative visual thinking tool for education.Journal of Educational Technology Systems, 2014","venue":null,"work_id":"1ee6138a-fd32-443c-94bb-bed0629a3bc5","year":2014},"citing_paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-07-07T19:32:20.095114Z"},"links":{"citing_paper":"/paper/2607.05290"},"observation_digest":"sha256:c85222cfdc7f1e29c4a3afd80162ee504add33c5c1e3d1b57d49d17fb6bb01f9","observation_id":"b92a58e9-49d6-4897-8119-b87c2eaff329","resolution":{"observed_at":"2026-07-07T19:34:06.708182Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2607.05290","last_updated":"2026-07-06T16:30:41Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-04T14:16:36.833173Z","submitted_at":"2026-07-06T16:30:41Z","title":"ChatImage: Navigating Long-Form LLM Answers through Interactive Images"},"reference_resolution":{"displayed":43,"state_counts":{"malformed_identifier":0,"metadata_mismatch":3,"parse_uncertain":0,"unresolved":2,"verified_exact":10,"verified_fuzzy":28},"total_outbound_references":43},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 43 of 43 outbound references and 0 inbound Pith citation observations for arXiv:2607.05290."}