{"as_of":"2026-08-23T00:05:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:5132a0d3e41856eccd13817139f1d1b6c028f01462f190525d66f4a42a5f07f3","coverage":[{"denominator":55,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":55,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T17:32:33.484173Z","state":"measured"},{"denominator":59,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":59,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-22T06:32:14.747728+00:00","state":"measured"},{"denominator":4,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":4,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T20:16:46.639351Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-14T22:03:02.530501Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.08907","snapshot_observed_at":"2026-08-15T20:16:46.639351Z","title":"Gaga: Towards interactive global geolocation assistant.arXiv preprint arXiv:2412.08907, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.13731","last_updated":"2026-06-24T08:21:04Z","snapshot_observed_at":"2026-08-20T05:42:03.654770Z","submitted_at":"2025-05-19T21:04:46Z","title":"GeoRanker: Distance-Aware Ranking for Worldwide Image Geolocalization","version":4},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-15T20:16:46.639351Z"},"links":{"cited_paper":"/paper/2412.08907","citing_paper":"/paper/2505.13731"},"observation_digest":"sha256:e1e14d7f0521a4213efb4da7057e2c395a8ddcf19786f2597375ca64e292b3a7","observation_id":"8c89c3a0-ae13-4bcf-9f6f-c5a3ab13aa51","resolution":{"observed_at":"2026-08-15T20:16:46.639351Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"cited_work":{"arxiv_id":"2412.08907","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2412.08907","snapshot_observed_at":"2026-07-03T02:17:04.954456Z","title":null,"venue":null,"work_id":"bd6959f4-1b27-4690-9a29-a7f27b2d1b55","year":2024},"citing_paper":{"arxiv_id":"2604.09025","last_updated":"2026-05-13T04:36:07Z","snapshot_observed_at":"2026-08-11T01:36:39.096057Z","submitted_at":"2026-04-10T06:43:48Z","title":"Skill-Conditioned Visual Geolocation for Vision-Language Models","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T17:49:50.853502Z"},"links":{"cited_paper":"/paper/2412.08907","citing_paper":"/paper/2604.09025"},"observation_digest":"sha256:81e69f381a89527ceed3f3b53ac550a37bfe5845986e723d7d6f4bc4c98c026f","observation_id":"8cf4701b-db38-4062-9ea6-89a8f45bd622","resolution":{"observed_at":"2026-07-03T02:17:04.954456Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"cited_work":{"arxiv_id":"2412.08907","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2412.08907","snapshot_observed_at":"2026-07-03T02:17:04.954456Z","title":null,"venue":null,"work_id":"bd6959f4-1b27-4690-9a29-a7f27b2d1b55","year":2024},"citing_paper":{"arxiv_id":"2604.09025","last_updated":"2026-05-13T04:36:07Z","snapshot_observed_at":"2026-08-11T01:36:39.096057Z","submitted_at":"2026-04-10T06:43:48Z","title":"Skill-Conditioned Visual Geolocation for Vision-Language Models","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-14T22:02:21.638064Z"},"links":{"cited_paper":"/paper/2412.08907","citing_paper":"/paper/2604.09025"},"observation_digest":"sha256:54da298b6eb986473418f53169ced1ed0579a25acd1d2e68f4fc5a39584b5b61","observation_id":"82601659-7706-4d78-8f36-8079aed23aa1","resolution":{"observed_at":"2026-07-03T02:17:04.954456Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.08907","snapshot_observed_at":"2026-08-01T23:46:29.140437Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.15255","last_updated":"2026-07-16T17:50:03Z","snapshot_observed_at":"2026-08-19T03:43:19.913460Z","submitted_at":"2026-07-16T17:50:03Z","title":"HoloGeo: Mitigating Landmark Bias in Geo-localization via Evidence-Driven Reasoning","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-01T23:46:29.140437Z"},"links":{"cited_paper":"/paper/2412.08907","citing_paper":"/paper/2607.15255"},"observation_digest":"sha256:c78af8d0c0b3819c898bb456b894462c6dfb9fa532fa5813fca4906a83d86aff","observation_id":"3e31d786-109e-4751-8ef4-8864f6edf202","resolution":{"observed_at":"2026-08-01T23:46:29.140437Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2412.08907/citation-record","integrity":"/paper/2412.08907/integrity","json":"/paper/2412.08907/citation-record.json","paper":"/paper/2412.08907"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.955863Z","title":"A survey of location prediction on twitter,","venue":null,"work_id":"b024b9cb-3ec0-425f-85cf-fe023a779e1b","year":2018},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.337142Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:fe69a4d9c1a2aa1cfedca56feadd4181724331786aba99e60bda5bd42f8a586d","observation_id":"09fbcd58-e34e-4aa5-97c3-af3ec2949e76","resolution":{"observed_at":"2026-08-11T17:32:33.958781Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.947907Z","title":"Metageo: A general framework for social user geolocation identification with few- shot learning,","venue":null,"work_id":"eb0ce3e1-cd15-496d-a758-3a40f109af2c","year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.340623Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:965ae6e19f2e2c44a7461d5bbe6de0140ed80a1ca71b1f02569bc830db9ab680","observation_id":"1e6298b9-b18c-4037-b26e-755908c9c55b","resolution":{"observed_at":"2026-08-11T17:32:33.950888Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.940454Z","title":"Cross-view image sequence geo- localization,","venue":null,"work_id":"b0b346c2-c858-4804-8a8d-19a77a8c0272","year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.343601Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:f0c233116bf26f368dc0d0223f532c2a8e4061591a5b83b61f0871c1e27fceed","observation_id":"74770304-33ea-4882-94a9-be4519f427db","resolution":{"observed_at":"2026-08-11T17:32:33.942970Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.933019Z","title":"TransGeo: Transformer Is All You Need for Cross-view Image Geo-localization,","venue":null,"work_id":"c44ff0d8-a20a-498e-b4a8-6f3be89d3c9c","year":2022},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.346334Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:079ebe442f24bdf16274d08d1a327ea3da9b05834146d6f425cd74be41ed843f","observation_id":"e890e9dc-1bca-4338-8dcd-a1774e4f5a86","resolution":{"observed_at":"2026-08-11T17:32:33.935659Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.925990Z","title":"CPlaNet: Enhancing Image Geolocalization by Combinatorial Partitioning of Maps,","venue":null,"work_id":"5c8f0160-995c-4921-a660-9d27c3bb1eda","year":2018},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.349206Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:9929fd3366d557e6491f87f4524517b70c2bcc3217ea70c6480f6fb2677eedad","observation_id":"5b369c81-0291-45c3-b5f1-5c68fb1dfa68","resolution":{"observed_at":"2026-08-11T17:32:33.928503Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.352019Z","title":"Geoclip: Clip- inspired alignment between locations and images for effective worldwide geo-localization,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.352019Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:b7f9a4d2248c39d191feedda626f80af6b0be3e0d86226145d18ac0df101c16a","observation_id":"e368e051-248e-4337-a3ff-afe046ae3dad","resolution":{"observed_at":"2026-08-11T17:32:33.352019Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.913826Z","title":"A latent variable model for geographic lexical variation,","venue":null,"work_id":"9fea131c-07f7-4a4c-808b-129e43e0c5b6","year":2010},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.354988Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:20473ab9859a7b66ae3816228076003e0b99fe0c43ccbfa7114e997eba829a61","observation_id":"fadad823-ef1d-4714-9e4c-c0ed1bc2d03d","resolution":{"observed_at":"2026-08-11T17:32:33.916655Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.905756Z","title":"Hierarchical geographical modeling of user locations from social media posts,","venue":null,"work_id":"5222aac2-cd38-435f-be06-c2a244107227","year":2013},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.357741Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:a97a71858874f7c66ea035227e10d1115a1aaa54079d673d5732763f52873313","observation_id":"9ba8449a-6a6e-4d46-a305-ce5c6c11be38","resolution":{"observed_at":"2026-08-11T17:32:33.908957Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.897703Z","title":"Text-based twitter user geolocation prediction,","venue":null,"work_id":"ef9453d6-b786-402a-8c6d-33ba82356963","year":2014},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.359988Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:13691e53c9ec815e11a903f9104679208aac69769841213bda4efe7a25a97e2e","observation_id":"5849383c-c7b5-441d-8890-8d8c2661a26d","resolution":{"observed_at":"2026-08-11T17:32:33.900746Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1704.04008","last_updated":"2017-04-27T01:18:58Z","snapshot_observed_at":"2026-08-17T13:23:39.338950Z","submitted_at":"2017-04-13T06:35:55Z","title":"A Neural Model for User Geolocation and Lexical Dialectology","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1704.04008","snapshot_observed_at":"2026-08-11T17:32:33.362202Z","title":"A neural model for user geolocation and lexical dialectology,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.362202Z"},"links":{"cited_paper":"/paper/1704.04008","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:985a504fe02b8fc5101ac4e2d3bf55e0c1fcdfa4f7fcfb0aba1db24c72761456","observation_id":"54e1617e-941e-4d3a-8d77-d875f42ad5a7","resolution":{"observed_at":"2026-08-11T17:32:33.362202Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.889390Z","title":"Kernel density estimation for text-based geolocation,","venue":null,"work_id":"2cfafc43-9f93-4a11-99b1-22379f6f1865","year":2015},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.364910Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:1d9d2e10a3217e7e77b77a63b6dfc49b2424c6d384df09b434f9c62f8e173f96","observation_id":"90f4d9d8-8ecd-41ec-bd26-ec4aa7d13340","resolution":{"observed_at":"2026-08-11T17:32:33.892323Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.881273Z","title":"Find me if you can: im- proving geographical prediction with social and spatial proximity,","venue":null,"work_id":"f0d35d52-bdf8-4524-b5c8-907dfa7bf7ed","year":null},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.367639Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:5908323c15df0ff78c9137bd96b231fea39eccbe6cc8a42ff373b436019b59dc","observation_id":"2cf34fd6-3ef9-4d3a-b86c-66cd4a53fbc5","resolution":{"observed_at":"2026-08-11T17:32:33.884228Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.873132Z","title":"Spot: Locating social media users based on social network context,","venue":null,"work_id":"9dcb6fbb-5927-4118-a58b-cab90521b359","year":2014},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.370195Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:b46c09ff1feade3a1e66254c303ac69d44b757b54d79a0cb3ed905c15adf2a15","observation_id":"f741f086-a611-4bfb-87aa-24c7227df05f","resolution":{"observed_at":"2026-08-11T17:32:33.876059Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1506.08259","last_updated":"2015-09-22T01:14:20Z","snapshot_observed_at":"2026-08-14T22:42:44.239410Z","submitted_at":"2015-06-27T04:51:18Z","title":"Twitter User Geolocation Using a Unified Text and Network Prediction Model","version":3},"cited_work":{"arxiv_id":"1506.08259","doi":null,"metadata_source":"pith","pith_arxiv_id":"1506.08259","snapshot_observed_at":"2026-08-11T17:32:33.708911Z","title":"Twitter User Geolocation Using a Unified Text and Network Prediction Model","venue":"cs.CL","work_id":"02599e82-6d9d-4108-ba0e-dbf1d80a62bc","year":2015},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.372292Z"},"links":{"cited_paper":"/paper/1506.08259","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:fee29f86ea6148666b594af85a1af0a6497971652fae4dd0e793d1c2e1c7142d","observation_id":"b2af9661-43e1-4b5e-baf7-55eace468856","resolution":{"observed_at":"2026-08-11T17:32:33.711743Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.08485","last_updated":"2023-12-11T17:46:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-04-17T17:59:25Z","title":"Visual Instruction Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.08485","snapshot_observed_at":"2026-08-11T17:32:33.375186Z","title":"Visual instruction tuning,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.375186Z"},"links":{"cited_paper":"/paper/2304.08485","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:f2a512f9af36e975974c7f2b67a01c8b627a482f2271caf6af0e61c312b69dbb","observation_id":"5c9525c6-8207-4379-b04e-71614242dddb","resolution":{"observed_at":"2026-08-11T17:32:33.375186Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.14238","last_updated":"2024-01-15T15:23:55Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-21T18:59:31Z","title":"InternVL: Scaling up Vision Foundation Models and Aligning for Generic Visual-Linguistic Tasks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.14238","snapshot_observed_at":"2026-08-11T17:32:33.378130Z","title":"Internvl: Scaling up vision foundation models and aligning for generic visual-linguistic tasks,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.378130Z"},"links":{"cited_paper":"/paper/2312.14238","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:4ef0d14d05ae52e662af6d08916b3c1aea5be373280bafdb9e5804a77569e51e","observation_id":"b1864b81-815b-437d-80c2-679dc3f7255e","resolution":{"observed_at":"2026-08-11T17:32:33.378130Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.14824","last_updated":"2023-07-13T05:41:34Z","snapshot_observed_at":"2026-08-12T12:24:23.815073Z","submitted_at":"2023-06-26T16:32:47Z","title":"Kosmos-2: Grounding Multimodal Large Language Models to the World","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.14824","snapshot_observed_at":"2026-08-11T17:32:33.381103Z","title":"Kosmos-2: Grounding multimodal large language models to the world,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.381103Z"},"links":{"cited_paper":"/paper/2306.14824","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:f0f831ad3823fdf633763e988dd359cedcbea0983a1d9c78a85fff7cd947f45e","observation_id":"c23ef287-5906-4f79-85e4-2f8ca7332fba","resolution":{"observed_at":"2026-08-11T17:32:33.381103Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.864975Z","title":"Openstreetview-5m: The many roads to global visual geolocation,","venue":null,"work_id":"dfe0a252-3cc2-4f88-9eb8-d32fea4de564","year":2024},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.384080Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:2d801f659e8a73d7c6cb2a8799dcc4da0bc4e8c7e66d1c6179debe13a4806c05","observation_id":"8453d228-8d2f-4c87-81fe-d2666c51171f","resolution":{"observed_at":"2026-08-11T17:32:33.868036Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.12793","last_updated":"2023-11-28T08:52:50Z","snapshot_observed_at":"2026-08-14T06:42:43.375489Z","submitted_at":"2023-11-21T18:58:11Z","title":"ShareGPT4V: Improving Large Multi-Modal Models with Better Captions","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.12793","snapshot_observed_at":"2026-08-11T17:32:33.386602Z","title":"Sharegpt4v: Improving large multi-modal models with better captions,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.386602Z"},"links":{"cited_paper":"/paper/2311.12793","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:1bd2e50a1645fbb91b14178d03374ccf43da21c4d939a047a1a8d28935d0550b","observation_id":"b46a8675-110c-4812-81eb-a6020ab9facd","resolution":{"observed_at":"2026-08-11T17:32:33.386602Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.856731Z","title":"DataComp: In search of the next generation of multimodal datasets,","venue":null,"work_id":"c1884603-0f76-46c2-b9f7-0e7b6ed51072","year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.389473Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:a894f233923f502e51f34541c4b133a4447891fb6b8ad60628494ce86c2f09be","observation_id":"e4ffe7a2-464e-4639-a6bd-7166a83a3f1b","resolution":{"observed_at":"2026-08-11T17:32:33.859732Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.05425","last_updated":"2023-06-08T17:59:56Z","snapshot_observed_at":"2026-08-16T15:25:23.093007Z","submitted_at":"2023-06-08T17:59:56Z","title":"MIMIC-IT: Multi-Modal In-Context Instruction Tuning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.05425","snapshot_observed_at":"2026-08-11T17:32:33.392215Z","title":"Mimic-it: Multi-modal in-context instruction tuning,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.392215Z"},"links":{"cited_paper":"/paper/2306.05425","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:3faeefe472ab84e0e563bc8867a40bc26db7663fce2e4134568ff1f19a80aeba","observation_id":"4d0b32a8-4952-4ee4-8afa-320e81f77688","resolution":{"observed_at":"2026-08-11T17:32:33.392215Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.395147Z","title":"Qlora: Efficient finetuning of quantized llms,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.395147Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:2cc2c2538f52f81914ac1539516d62abb9c052d7b53596684af789fbd2d7077a","observation_id":"76a63df7-c08e-4c5a-934b-793bc1860d40","resolution":{"observed_at":"2026-08-11T17:32:33.395147Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.844621Z","title":"Translocator: local realignment and global remapping enabling accurate translocation detec- tion using single-molecule sequencing long reads,","venue":null,"work_id":"a01dbbcb-8877-4091-b26d-1fb197d18a18","year":2020},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.397638Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:0d511efde42ba6a4ad55704784c58f818af48241f2b1f9fe6830134f34361593","observation_id":"7d5e9dd5-c3e2-4b14-8a8e-c161f571f877","resolution":{"observed_at":"2026-08-11T17:32:33.847627Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.400189Z","title":"Learning transferable visual models from natural language supervision,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.400189Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:452a0c04999aa4362874d50b9266306f134797a6d29af4d8dd632b5452ebbf72","observation_id":"c63c8ff8-7568-4533-88dc-bbded4ffc598","resolution":{"observed_at":"2026-08-11T17:32:33.400189Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.833513Z","title":"Pigeon: Predicting image geolocations,","venue":null,"work_id":"48423172-a1fb-4c13-99ec-1b7031a3d30e","year":2024},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.402777Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:87bcd3c2aff78db0222f10b268918960d498e7edb3ecb3adfa4066544c8c3463","observation_id":"d19e9fb5-cdf8-46e6-971a-c36c1291b0ff","resolution":{"observed_at":"2026-08-11T17:32:33.836093Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.15521","last_updated":"2022-11-28T16:34:40Z","snapshot_observed_at":"2026-08-16T16:12:44.837344Z","submitted_at":"2022-11-28T16:34:40Z","title":"G^3: Geolocation via Guidebook Grounding","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.15521","snapshot_observed_at":"2026-08-11T17:32:33.405340Z","title":"Gˆ 3: Ge- olocation via guidebook grounding,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.405340Z"},"links":{"cited_paper":"/paper/2211.15521","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:c8498c589f0cd941f55adf564a2a080262031ea64cacfa00280e453bcca3678c","observation_id":"d52db36e-4dc3-491f-a083-6ae34c9f3c8e","resolution":{"observed_at":"2026-08-11T17:32:33.405340Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.00275","last_updated":"2023-02-01T06:44:07Z","snapshot_observed_at":"2026-08-16T15:58:36.220567Z","submitted_at":"2023-02-01T06:44:07Z","title":"Learning Generalized Zero-Shot Learners for Open-Domain Image Geolocalization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.00275","snapshot_observed_at":"2026-08-11T17:32:33.408397Z","title":"Learning generalized zero- shot learners for open-domain image geolocalization,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.408397Z"},"links":{"cited_paper":"/paper/2302.00275","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:904e92a938cb62d3792bcdc244325e9a1ca6b2ef4a482ae66bee8891d0c96321","observation_id":"81502e4c-70f4-46ff-99b0-18d3bb8c63de","resolution":{"observed_at":"2026-08-11T17:32:33.408397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.826010Z","title":"Georeasoner: Geo-localization with reasoning in street views using a large vision-language model,","venue":null,"work_id":"97330ffb-cec2-4770-a4a1-851fdc06b1e4","year":null},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.411212Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:818e15617ac8804e65d69487dec3082bde8f3f4937aca12c3bc9e893fa29add5","observation_id":"b78743d7-1858-4331-871e-a74f62d848d5","resolution":{"observed_at":"2026-08-11T17:32:33.828675Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.416576Z","title":"Im2gps: estimating geographic information from a single image,","venue":null,"work_id":null,"year":2008},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.416576Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:e612c9a65fdeb85abc8e0541ffe7102e41e275c2cab30781335bd17f00daa4ae","observation_id":"f983ba4e-d520-483a-9f42-448b4c58af91","resolution":{"observed_at":"2026-08-11T17:32:33.416576Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.814007Z","title":"Revisiting im2gps in the deep learning era,","venue":null,"work_id":"9b5cacd9-0330-4b18-ad72-cf8e9a8835e9","year":2017},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.419779Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:74a5c773821242d663b209aa4202007f1635018ed3ec90c01adfc69e25230f09","observation_id":"04e9856b-4312-4f0c-8e58-3fae93a8e166","resolution":{"observed_at":"2026-08-11T17:32:33.817043Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.805026Z","title":"Interpretable semantic photo geolocation,","venue":null,"work_id":"2b70b51a-1c44-4b45-a43e-3e712047734e","year":2022},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.422315Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:9e179a89f6ac268037e2d33f7ae49973a6fd836a56bb805d61ab7285bc365bf6","observation_id":"53074e8a-0dfe-49ad-a6d6-af2f3f561b35","resolution":{"observed_at":"2026-08-11T17:32:33.808716Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.424921Z","title":"Where we are and what we’re looking at: Query based worldwide image geo-localization using hierarchies and scenes,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.424921Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:02d7e7a2fc365d4d1fdb276ece6fbc6ba7554a303b1c27a60ad9c5d6df3a200d","observation_id":"a5f57851-afe1-42dc-bcee-3456deb05d75","resolution":{"observed_at":"2026-08-11T17:32:33.424921Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.792717Z","title":"Google landmarks dataset v2-a large-scale benchmark for instance-level recognition and retrieval,","venue":null,"work_id":"372dd69b-5dee-4906-969c-2dfc0651ec02","year":2020},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.427488Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:0de0f70d38058281770771fb1a4077eeed333e50179a264f53056eea8bf976e6","observation_id":"67e9ba1e-0ff0-4b75-b5cc-47f2e79802c4","resolution":{"observed_at":"2026-08-11T17:32:33.795636Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.12966","last_updated":"2023-10-13T02:41:28Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-08-24T17:59:17Z","title":"Qwen-VL: A Versatile Vision-Language Model for Understanding, Localization, Text Reading, and Beyond","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.12966","snapshot_observed_at":"2026-08-11T17:32:33.430188Z","title":"Qwen-vl: A frontier large vision-language model with versatile abilities,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.430188Z"},"links":{"cited_paper":"/paper/2308.12966","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:a644b61ca8633ab0b9868277cff8133f17feed73d1273f3a76bdf3041b4856b0","observation_id":"694342d2-8ef4-4fe7-90fa-9ca81dd1590e","resolution":{"observed_at":"2026-08-11T17:32:33.430188Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.16821","last_updated":"2024-04-29T20:24:30Z","snapshot_observed_at":"2026-08-17T14:16:52.244007Z","submitted_at":"2024-04-25T17:59:19Z","title":"How Far Are We to GPT-4V? Closing the Gap to Commercial Multimodal Models with Open-Source Suites","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.16821","snapshot_observed_at":"2026-08-11T17:32:33.432981Z","title":"How far are we to gpt-4v? closing the gap to commercial multimodal models with open-source suites,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.432981Z"},"links":{"cited_paper":"/paper/2404.16821","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:457bcab8e5d2e1e9cefdd194f2604a8b62ffcaabddd99f99070f93bf047ee34f","observation_id":"a9b9e55f-04ae-4a9c-9dc3-11a61f254cb4","resolution":{"observed_at":"2026-08-11T17:32:33.432981Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.00890","last_updated":"2023-06-01T16:50:07Z","snapshot_observed_at":"2026-08-09T05:23:28.778360Z","submitted_at":"2023-06-01T16:50:07Z","title":"LLaVA-Med: Training a Large Language-and-Vision Assistant for Biomedicine in One Day","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.00890","snapshot_observed_at":"2026-08-11T17:32:33.435441Z","title":"Llava-med: Training a large language-and-vision assistant for biomedicine in one day,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.435441Z"},"links":{"cited_paper":"/paper/2306.00890","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:b595b63fc3f01b646445035ac7b288b801c21846e78d021c71de35ee69019da3","observation_id":"6b5c5e11-2ac1-4ade-b695-5e6791d23139","resolution":{"observed_at":"2026-08-11T17:32:33.435441Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.15826","last_updated":"2023-11-24T18:59:10Z","snapshot_observed_at":"2026-08-16T14:40:34.354977Z","submitted_at":"2023-11-24T18:59:10Z","title":"GeoChat: Grounded Large Vision-Language Model for Remote Sensing","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.15826","snapshot_observed_at":"2026-08-11T17:32:33.437987Z","title":"Geochat: Grounded large vision-language model for remote sensing,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.437987Z"},"links":{"cited_paper":"/paper/2311.15826","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:8a8456d527184c3ebe405a06bd39f863b709b739e4fa141b7b3e60c94ec55245","observation_id":"4a48eeff-d293-4c5b-b6d5-183e6edb6a26","resolution":{"observed_at":"2026-08-11T17:32:33.437987Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.12915","last_updated":"2025-03-04T18:51:48Z","snapshot_observed_at":"2026-08-16T01:01:16.814303Z","submitted_at":"2024-11-19T22:59:14Z","title":"VILA-M3: Enhancing Vision-Language Models with Medical Expert Knowledge","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.12915","snapshot_observed_at":"2026-08-11T17:32:33.440369Z","title":"Vila-m3: Enhancing vision-language models with medical expert knowledge,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.440369Z"},"links":{"cited_paper":"/paper/2411.12915","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:5c1e0f524e95d7d6857a48885d7abb8fde5fae50cc22479f0c8662eca9f842b7","observation_id":"aefab9d9-57d6-450c-b074-86525182277c","resolution":{"observed_at":"2026-08-11T17:32:33.440369Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-08-16T09:25:53.087782Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-08-11T17:32:33.442776Z","title":"An image is worth 16x16 words: Transformers for image recognition at scale,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.442776Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:5250f03a6666a3eb9106f75bed0997bb508e1b3048ce02a905a4046b7d6b80d7","observation_id":"c06a7f2b-3efa-4fbd-abc0-65707e1a0a15","resolution":{"observed_at":"2026-08-11T17:32:33.442776Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.784558Z","title":"Xtuner: A toolkit for efficiently fine-tuning llm,","venue":null,"work_id":"0a3feb91-c194-4f7e-860b-ee2e1562035a","year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.445575Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:44f9c1b44c6647bc5e0613d2f3b4a6c97289c5cc70b0641e1237d676b2d58f56","observation_id":"daef10b4-faa8-40d7-810c-46c07c988cd4","resolution":{"observed_at":"2026-08-11T17:32:33.787694Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.776763Z","title":"Lmdeploy: A toolkit for compressing, deploying, and serving llm,","venue":null,"work_id":"11cc9eda-f857-4f09-8e31-77820feb57b8","year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.448141Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:579d76dd1583bc2f29421a3cf2347a4ee9afde10570cd756f2f279ad0522dcb1","observation_id":"f545894c-f2b1-4c8f-8b57-b4914ad3d859","resolution":{"observed_at":"2026-08-11T17:32:33.779630Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2201.11903","last_updated":"2023-01-10T23:07:57Z","snapshot_observed_at":"2026-08-13T07:04:41.220509Z","submitted_at":"2022-01-28T02:33:07Z","title":"Chain-of-Thought Prompting Elicits Reasoning in Large Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2201.11903","snapshot_observed_at":"2026-08-11T17:32:33.450720Z","title":"Chain-of-thought prompting elicits reasoning in large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.450720Z"},"links":{"cited_paper":"/paper/2201.11903","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:4356c8b6553d44b589a5e153da39a7c163ab01ab1c3887f56be63ff3296e1f81","observation_id":"a1150d93-fdbf-44b0-a4b0-c247e972fa64","resolution":{"observed_at":"2026-08-11T17:32:33.450720Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.768684Z","title":"Openstreetview- 5m: The many roads to global visual geolocation,","venue":null,"work_id":"e4fff208-d424-4d59-b59e-014de7cfcfe7","year":2024},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.453624Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:5ebaa615c718f98ac4b37928ca3a4d5d4c118ec8e7766e6bd3f9ca4c86d1f662","observation_id":"545d4497-a3dc-46c4-b91a-b31413321f98","resolution":{"observed_at":"2026-08-11T17:32:33.771737Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.456262Z","title":"Language models as zero-shot planners: Extracting actionable knowledge for embodied agents,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.456262Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:941242d609864e0d4836343d50800cc06169211df7bb27679d4125c9f6cd41e9","observation_id":"a63ea552-a5bf-4b55-8131-7e59c7441834","resolution":{"observed_at":"2026-08-11T17:32:33.456262Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.06770","last_updated":"2024-11-11T23:05:04Z","snapshot_observed_at":"2026-08-18T08:11:45.716032Z","submitted_at":"2023-10-10T16:47:29Z","title":"SWE-bench: Can Language Models Resolve Real-World GitHub Issues?","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.06770","snapshot_observed_at":"2026-08-11T17:32:33.458731Z","title":"Swe-bench: Can language models resolve real-world github issues?, 2024,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.458731Z"},"links":{"cited_paper":"/paper/2310.06770","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:f4a7bd2d5552a46ed6756cd452b8a871484e4bcb15a59a9868401d64099df72f","observation_id":"edca6d91-10d9-4547-83a6-4a1e69a1f1d3","resolution":{"observed_at":"2026-08-11T17:32:33.458731Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.755700Z","title":"Vipergpt: Visual inference via python execution for reasoning,","venue":null,"work_id":"f88ce3f6-abdd-486c-9be9-c8e69f8b0bb9","year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.461721Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:e234234c6efa0ead309b3bc0c4def08a364da30b4e77c4f460af2c8f7c2ba74e","observation_id":"5e1a1211-39d7-407b-b355-5bda1339e08a","resolution":{"observed_at":"2026-08-11T17:32:33.759029Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.16291","last_updated":"2023-10-19T16:27:03Z","snapshot_observed_at":"2026-08-13T16:55:46.498156Z","submitted_at":"2023-05-25T17:46:38Z","title":"Voyager: An Open-Ended Embodied Agent with Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.16291","snapshot_observed_at":"2026-08-11T17:32:33.464238Z","title":"V oyager: An open-ended embodied agent with large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.464238Z"},"links":{"cited_paper":"/paper/2305.16291","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:b07fc8bd1783749722946bc7ac31347185f8cae9e1b498a56fd99c992c39dd0f","observation_id":"69d6b38d-1f80-4086-a1f8-eeab4b25978d","resolution":{"observed_at":"2026-08-11T17:32:33.464238Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.747711Z","title":"Llama 3 model card,","venue":null,"work_id":"05f85b65-f96f-40a0-91c3-fe358b9d81a3","year":2024},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.467192Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:eceb260950f174a891fd78f56c1126788a3250a6b1c05191bc72b129b8293389","observation_id":"8485378c-ce87-48fe-9b1c-013eb207ef9b","resolution":{"observed_at":"2026-08-11T17:32:33.750537Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.739369Z","title":"Geolocation estima- tion of photos using a hierarchical model and scene classification,","venue":null,"work_id":"94d6e90b-4581-4fa0-a306-c6fd2e53c013","year":2018},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.469759Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:34666436f51824e12cd22fe71217efd62ca48c7659ca5b1b005e23a30a208904","observation_id":"c735e2c7-369b-4378-b1b6-053e233a9a6a","resolution":{"observed_at":"2026-08-11T17:32:33.742419Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.15118","last_updated":"2024-02-18T23:44:05Z","snapshot_observed_at":"2026-08-16T14:24:22.676964Z","submitted_at":"2024-01-26T02:39:40Z","title":"GeoDecoder: Empowering Multimodal Map Understanding","version":2},"cited_work":{"arxiv_id":"2401.15118","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.15118","snapshot_observed_at":"2026-08-11T17:32:33.518989Z","title":"GeoDecoder: Empowering Multimodal Map Understanding","venue":"cs.CV","work_id":"d4417413-5564-4b12-ba20-a90c0f191ef0","year":2024},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.472387Z"},"links":{"cited_paper":"/paper/2401.15118","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:2834d55da32470d3b056328b79ca881911eabef9f19d5143ca3dc253b312fa95","observation_id":"21a2e0ce-ac70-4854-a850-eab16f50eeb6","resolution":{"observed_at":"2026-08-11T17:32:33.524341Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.01728","last_updated":"2024-01-29T06:27:53Z","snapshot_observed_at":"2026-08-16T11:06:17.042878Z","submitted_at":"2023-10-03T01:31:25Z","title":"Time-LLM: Time Series Forecasting by Reprogramming Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.01728","snapshot_observed_at":"2026-08-11T17:32:33.475256Z","title":"Time-llm: Time series forecasting by reprogramming large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.475256Z"},"links":{"cited_paper":"/paper/2310.01728","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:858e422defbe42d09b4a6187c909c4d8de02f1fbeb708e2c79476105f49ed585","observation_id":"65f7b459-7430-4341-9010-144decd1d4a7","resolution":{"observed_at":"2026-08-11T17:32:33.475256Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.731114Z","title":"Better and faster large language models via multi-token prediction,","venue":null,"work_id":"527113c3-41a6-4ef0-af1b-fe45362a1313","year":null},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.478193Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:a65b308d20e8c5aa8cbd42e6d143196f3c6c24a4e4338f843b50a80f87520939","observation_id":"4417a9df-12a3-4580-86cb-80a8ce9c8552","resolution":{"observed_at":"2026-08-11T17:32:33.734125Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.723406Z","title":"Measuring nominal scale agreement among many raters,","venue":null,"work_id":"dce7f826-be75-4ad2-915f-f7710638deaf","year":1971},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.484173Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:0292daddf66dd923e80e3fc9243db18f1bd905eb904cb5f9e9a415346a7003b0","observation_id":"f9b6c31d-25ab-4a83-bd17-369533b22569","resolution":{"observed_at":"2026-08-11T17:32:33.726347Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.19737","last_updated":"2024-04-30T17:33:57Z","snapshot_observed_at":"2026-08-16T23:36:44.216328Z","submitted_at":"2024-04-30T17:33:57Z","title":"Better & Faster Large Language Models via Multi-token Prediction","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.19737","snapshot_observed_at":"2026-08-11T17:32:33.481104Z","title":"Available: https://arxiv.org/abs/2404.19737","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.481104Z"},"links":{"cited_paper":"/paper/2404.19737","citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:cad1db1e535b63df36cbfa1c6270adedbc7297603b3f1022d885ca1ec63f2fe3","observation_id":"b8728bc1-0d66-4dbc-9b49-0f9e5600817d","resolution":{"observed_at":"2026-08-11T17:32:33.481104Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T17:32:33.413890Z","title":"Available: https://arxiv.org/abs/2406.18572","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant","version":3},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-11T17:32:33.413890Z"},"links":{"citing_paper":"/paper/2412.08907"},"observation_digest":"sha256:527381a7b3d289d33bc30a36864f506c1ef965085caca9e40e8fbcdda40386a8","observation_id":"de294b2a-540f-4509-a9e0-01fa84d4ae71","resolution":{"observed_at":"2026-08-11T17:32:33.413890Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2412.08907","last_updated":"2026-07-02T13:52:05Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-17T20:10:01.183255Z","submitted_at":"2024-12-12T03:39:44Z","title":"Towards Interactive Global Geolocation Assistant"},"reference_resolution":{"displayed":55,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":26,"verified_exact":2,"verified_fuzzy":27},"total_outbound_references":55},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"thesis":"As of 23 August 2026, this Paper Citation Record lists 55 of 55 outbound references and 4 inbound Pith citation observations for arXiv:2412.08907."}