{"as_of":"2026-08-08T11:17:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:8f88a503dffd30006d28818bc6dd1aabd483b5e86fb43f111c96ed01fe1e8186","coverage":[{"denominator":58,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":58,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-10T13:42:04.410771Z","state":"measured"},{"denominator":58,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":58,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2604.13941/citation-record","integrity":"/paper/2604.13941/integrity","json":"/paper/2604.13941/citation-record.json","paper":"/paper/2604.13941"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Deep learning reforms image matching: A survey and ou tlook","venue":null,"work_id":"b491c887-8821-4178-86df-1a50cb5313a1","year":2025},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:b4a67b7cae3826fcaea3c3a592e16bb85956d7c521b76919bdc66d68f61d665b","observation_id":"e770c480-a0e6-49b6-939a-eff9f8749054","resolution":{"observed_at":"2026-05-18T22:32:53.097924Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"EC-SfM : Efﬁcient covisibility-based structure-from-motion for b oth sequential and unordered images","venue":null,"work_id":"a1127ad1-3953-47e2-bbfd-95e3bd5e3250","year":2024},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:1cbd8bea12a0f6d55404c572290c189219c7fb6f6eeab2975aa87eeeabcdf647","observation_id":"62550ea6-af53-4cef-bc49-0da581d34902","resolution":{"observed_at":"2026-05-18T22:31:55.130734Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"PAS-SLA M: A visual SLAM system for planar-ambiguous scenes","venue":null,"work_id":"5d731d64-3f07-40b4-91d6-285a4efdf2fe","year":2026},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:e5e6bc098beff3be623e0714b484f8a41ed9aa7c406672e0010d75d2228733e7","observation_id":"89b7c68c-954b-40f7-aeb3-d5b6810bc578","resolution":{"observed_at":"2026-05-18T22:32:53.177432Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Distinctive image features from scale-inva riant keypoints","venue":null,"work_id":"5986048f-5460-48b5-bc15-143ff06f0cc5","year":2004},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:728e496cc5c618d3222c4c7c2c3675573ee9dbd53a2b2f05593bf73e393c9295","observation_id":"695d269c-7fc8-4415-8807-0a3ac6df1285","resolution":{"observed_at":"2026-05-18T22:31:55.119609Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"D2-Net: A trainable CNN for joint descripti on and detection of local features","venue":null,"work_id":"dda63390-6b1b-4034-b8ad-d961ffc4c1c4","year":2019},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:f4acf37b06e76cebe2b09e08c1a9f3734dea053431ed10c32f5c445140854dac","observation_id":"79e6dd47-6d13-400f-8010-f00797605c2e","resolution":{"observed_at":"2026-05-18T22:31:55.115721Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"SuperPoi nt: Self- supervised interest point detection and description","venue":null,"work_id":"60c0d2a3-ba27-4dd2-a9d6-4f74224a8f82","year":2018},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:09032eef69e30c88b52f3b218b4957f721e0b2c77128900f3ab5b0d3bdf5daa6","observation_id":"af7540fb-ad75-4a1c-88d3-74c8e98c9e5c","resolution":{"observed_at":"2026-05-18T22:31:55.126845Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"MatchMa mba: Correspondence pruning via selective state space model","venue":null,"work_id":"2e3d347b-1818-4bda-93c6-527bc1bf74e1","year":2025},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:1aa1ef488a4e9930cdb6374a80b195b365cd85c7afc302effb7e54eadac69dfa","observation_id":"8786c8f9-e69b-4233-a805-dde17f9bd90f","resolution":{"observed_at":"2026-05-18T22:32:53.156198Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"CGR-Net: Consistency guided ResFormer for two-view corre spondence learning","venue":null,"work_id":"1e8dc77d-9296-4af5-8a6c-03807be6898d","year":2024},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:c31e47c7ba826b111d3bde5afc35c57ab4c48044cabbc881b562fa707f2840a1","observation_id":"f4e2c590-5f6a-4c85-89f4-cec8cac847e5","resolution":{"observed_at":"2026-05-18T22:32:53.087628Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"MGCNet: Multi-granularity consensus network for remote s ensing image correspondence pruning","venue":null,"work_id":"3fd326c3-8306-49c3-b3d3-a1575b816dd8","year":2025},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:dd1f3c17c18f9586bf97e02bf4a64772530afdf972ff562953c1f8a3e5d33f67","observation_id":"54f7fc91-3bfe-4c55-a5d6-cb705ce68893","resolution":{"observed_at":"2026-05-18T22:32:53.062604Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"U-Match: Exploring hierarch y-aware local context for two-view correspondence learning","venue":null,"work_id":"85afb825-b621-41d7-b760-39a220b372ad","year":2024},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:03219b62cefe9352152953d59af25f715e430d6d4c0fc7eeba03e6540bbd0df7","observation_id":"3af73a66-2560-4f62-a195-daf8a3f87360","resolution":{"observed_at":"2026-05-18T22:31:55.111951Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"MIN IMA: Modality invariant image matching","venue":null,"work_id":"3caefbc9-391c-4f80-bd65-d1e65be4d478","year":2025},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:a0e888b6ef01c278a4db19ed8be59f7c96f7150701e76c55f6f77df11891b484","observation_id":"21ec676d-16d4-4cde-bf83-4c801fa21c84","resolution":{"observed_at":"2026-05-18T22:31:55.137781Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"ASLFeat: Learning local features of accurate s hape and localization","venue":null,"work_id":"ca65e122-76d5-4a26-90ec-dd2a1dc1df96","year":2020},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:6fec84e968d38833aabee330628d77474674706bf4af82015290776b306e36fe","observation_id":"aa4efce6-14df-4d89-9354-e2d66d885e04","resolution":{"observed_at":"2026-05-18T22:32:53.163188Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Super- Glue: Learning feature matching with graph neural networks","venue":null,"work_id":"ed6358c7-acdb-4280-b167-2c7b279fb611","year":2020},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:ec575294f0ceda020564b4a4fba84f5d5eebf41ee1928c0395324f390384247d","observation_id":"a98ce08f-9f49-4e7a-bf9c-1698127b5f60","resolution":{"observed_at":"2026-05-18T22:32:53.058358Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"LoFTR: Dete ctor-free local feature matching with transformers","venue":null,"work_id":"a4563005-5e3e-4e5b-af1a-b7c8e20589c9","year":2021},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:8238bc03d05d0cc733114bd2fcefa5a38cccffa4e2065a0a013f0e74c9421ad3","observation_id":"e1765369-c90c-4ee0-85e7-83f72d85acc2","resolution":{"observed_at":"2026-05-18T22:32:53.173859Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Ligh tGlue: Local feature matching at light speed","venue":null,"work_id":"d2698a77-c237-4623-93da-14141e14caeb","year":2023},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:5c4614586d85906140307d15c8da65f047bdaa8f157f11a2dddbcaac205adae1","observation_id":"13fcc463-c9db-4a56-94b9-bbf92578b32d","resolution":{"observed_at":"2026-05-18T22:31:55.123059Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"ORB: An efﬁcient alternative to SIFT or SURF","venue":null,"work_id":"76190a0f-84e6-4da8-851b-769422222b01","year":2011},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:b24fe5af6d66d3d77dab36c4e4a2eccc4a4e17da3bfd68308d2f26191e30b4f3","observation_id":"20b565da-154c-4eeb-a9b3-09fa342b061c","resolution":{"observed_at":"2026-05-18T22:32:53.138717Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"R2D2: Reliable and repeatable detector and descriptor","venue":null,"work_id":"1f6a76d9-16cb-40a4-92ca-c5abdaafa809","year":2019},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:c0b89559f732f9d0b93db31803ec3f8a8ded57131b4d68c951b6cdd4b422e1a6","observation_id":"b4d3cd98-e6c8-4375-ad22-4d31c79b8a1b","resolution":{"observed_at":"2026-05-18T22:31:55.134454Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Attention is all you need","venue":null,"work_id":"2f52feaa-f8f1-4286-abc0-dca8944784e8","year":2017},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:91471232a65662b87f5ae39163a39c4370eee76bc24b18ad4b681be9b8006270","observation_id":"f33aa160-3b95-4ddf-9b1e-1f33d30188bc","resolution":{"observed_at":"2026-05-18T22:32:53.090877Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Swin Transformer: Hierarchical vision Transformer using shifted win- dows","venue":null,"work_id":"d4d8fd84-702f-4d39-8894-b47192dc606c","year":2021},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:d03b77895eceb7daaa0a7b3205cd288cb021362d6d8e3b8359d4bd40071bdaa0","observation_id":"324330b6-52a7-4632-acb3-d1f8944e9cc3","resolution":{"observed_at":"2026-05-18T22:32:53.076920Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Match- Former: Interleaving attention in Transformers for featur e matching","venue":null,"work_id":"07f9ffdf-7760-4289-8375-370dc3baa436","year":2022},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:531df6598af07d26d69d8fdad42a7772caf2d08c20522f46e7eef820dd61afcc","observation_id":"919ef73a-10f2-4e6b-9c1f-cbdd97b462c2","resolution":{"observed_at":"2026-05-18T22:32:53.054643Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"An image patch is a wave: Phase-aware vision MLP","venue":null,"work_id":"be7c1229-a736-4850-a8e7-febb035d5341","year":2022},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:8910059fe195ecaa22ba241d5efb9ed51023ca55d2925a39f4a8729fe4dbe051","observation_id":"5d02cd5e-8d25-4e08-b2de-b90e7acca1af","resolution":{"observed_at":"2026-05-18T22:32:53.124518Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Learning feature matching via matchabl e keypoint- assisted graph neural network","venue":null,"work_id":"6b7ab7bb-b95b-42fe-baec-5837a35ddd79","year":2025},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:ac20eda1216bb80ee207744c6a7fb746cbdf776a954e3aff4e1db111b984705d","observation_id":"bdec6fa8-cfdd-4164-9d7e-f7b1d97ead38","resolution":{"observed_at":"2026-05-18T22:32:53.149616Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Guide local f eature matching by overlap estimation","venue":null,"work_id":"eb0616af-97cd-47b0-bc35-3eb44d8d1999","year":2022},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:6d1ef696ea482e8b4c83ce49ae8353b1faa99688feffb11a8fba1f16d29e796c","observation_id":"d328a92b-86eb-4cfe-a9f7-f0cee74c4b17","resolution":{"observed_at":"2026-05-18T22:32:53.192611Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Le arning accurate dense correspondences and when to trust them","venue":null,"work_id":"97abca61-51d0-4a10-bad7-35c526050acb","year":2021},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:b97e67e61d71f18add52ae3865582c5aa2d41b2bd0b5eefefb84b05c33537281","observation_id":"3154f19e-b00f-499e-8fa1-80db4d438b98","resolution":{"observed_at":"2026-05-18T22:32:53.121040Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"DeepMatcher: A deep transformer-based network for robust and accurate loc al feature matching","venue":null,"work_id":"32fd899f-7487-405e-963a-484992d0d31e","year":2024},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:09937d7f978b3d142e6cf2adb481609f87d12f968acf623ab8a8574adf379bc3","observation_id":"5ce0987b-37c8-47a9-9bfe-fdfd33c45559","resolution":{"observed_at":"2026-05-18T22:32:53.100110Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"VD-Matcher: A very deep local feature matcher w ith weight recycling and keypoint detection","venue":null,"work_id":"8aae10b3-238f-41a8-8883-93ad0a9e5e51","year":2025},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:0ca91221192766730671cbc106f69b84b7c5570a646ed6c644dee89e99734953","observation_id":"df2524c6-3703-491d-b968-3ac4402e86c4","resolution":{"observed_at":"2026-05-18T22:32:53.102056Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Adaptiv e spot-guided Transformer for consistent local feature matching","venue":null,"work_id":"15b76d8a-69d4-40da-9166-6e06efce5711","year":2023},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:320dacc762b747df7cb26980af7819b481620af97a7b468ea72760738f170318","observation_id":"87df8324-724a-41b3-9755-43fda28a3948","resolution":{"observed_at":"2026-05-18T22:32:53.113443Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"ContextDesc: Local descriptor augmentation with cross- modality context","venue":null,"work_id":"283490c1-28e0-497b-8973-38b856fc08b7","year":2019},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:3882b1de6b69ce29f88a0021c374897d9325a070f6294a2a69db0011f6b20a00","observation_id":"6bb2a39b-4388-4b10-a55c-99b67e90351a","resolution":{"observed_at":"2026-05-18T22:32:53.166944Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Attention weighted local descriptors","venue":null,"work_id":"86fd5560-4faa-491f-b0d3-f61c602fa8f2","year":2023},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:01cca084c88e24fc94e8e5e6f27a2c7cb08a6917412d9944c111ffbea56b9749","observation_id":"47d3affa-5ddf-43a8-bff9-12b4e3851995","resolution":{"observed_at":"2026-05-18T22:31:55.141291Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"OAMa tcher: An overlapping areas-based network with label credibility for robust and accurate feature matching","venue":null,"work_id":"2a8a72bf-57fa-4bcc-a259-0e269f01d39f","year":2024},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:e660b313b3dd61848ecca2777a7dd82aefad20d0ad47bb0ffdcb831b62261d67","observation_id":"c1d3ad0b-396b-470a-b248-5c30f0a5036b","resolution":{"observed_at":"2026-05-18T22:32:53.089403Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Adaptive assignment for geometry aware local f eature matching","venue":null,"work_id":"33b49b0a-d5a7-4bb9-8820-fd2d1d2a04a2","year":2023},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:3dbffe93886fe115962b3c27edd8e2c8cd5f6023a3c595965646a3ab593c51b5","observation_id":"48f0b426-8866-474e-8448-073b02009dd0","resolution":{"observed_at":"2026-05-18T22:32:53.189223Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Scene-aware feature mat ching","venue":null,"work_id":"2cf985b1-7120-498f-b6bd-efdd70ae114c","year":2023},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:9cd4f3f0111e2c1f83a77b6ad4bd2b5792fdf3618b9c1f3b5777a1d56fa16bfa","observation_id":"75aa28c9-71e6-4b49-b062-e56cc8b626c6","resolution":{"observed_at":"2026-05-18T22:32:53.073448Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"CoMatch: Dynam ic covisibility-aware transformer for bilateral subpixel-l evel semi-dense image matching","venue":null,"work_id":"c140f25e-80b9-4b25-90c3-3435ceb22b75","year":2025},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:a31406ba00a3d54ac2bc89fb18694caf20b1e76ad05f46adff39c9d3c84d8f21","observation_id":"e4acdbb5-4cc0-4ebc-a170-05c26b2b329f","resolution":{"observed_at":"2026-05-18T22:32:53.159906Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Object retrieval with large vocabularies and fast spatial matchin g","venue":null,"work_id":"1dd838f8-e039-4710-b71d-5b5cee4ade68","year":2007},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:32dfa3cf1d15b3d8e4a19e26061086cec45272c1fe1cc28a8eac9ac8b42643b6","observation_id":"2ba12198-115a-4e5b-b5ac-a79e58cb9446","resolution":{"observed_at":"2026-05-18T22:31:55.148670Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"MegaDepth: Learning single-view depth pre- diction from internet photos","venue":null,"work_id":"c285c7a3-02d9-4c2f-9e3c-9a02b1118beb","year":2018},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:ac081b0d268dc5a1b3f2aa4863cb712cf23ed3fc60e8589a3094ae8d3a47bdfb","observation_id":"c6846692-952c-4585-9dc7-7eb41461da7e","resolution":{"observed_at":"2026-05-18T22:31:55.156180Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1412.6980","last_updated":"2017-01-30T01:27:54Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2014-12-22T13:54:29Z","title":"Adam: A Method for Stochastic Optimization","version":9},"cited_work":{"arxiv_id":"1412.6980","doi":"10.1002/mrm.28086","metadata_source":"pith","pith_arxiv_id":"1412.6980","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Adam: A Method for Stochastic Optimization","venue":"cs.LG","work_id":"1910796d-9b52-4683-bf5c-de9632c1028b","year":2014},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"cited_paper":"/paper/1412.6980","citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:800ccbb44e004fdca532762c3ede17713b7c56a707dca4161b491d6b67539af3","observation_id":"348f8073-b6c2-4804-9434-4f12027bcd21","resolution":{"observed_at":"2026-05-10T13:45:28.115565Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"HP atches: A benchmark and evaluation of handcrafted and learned local d escriptors","venue":null,"work_id":"c6ae2889-46f8-437a-ba9e-8cd62b3031f8","year":2017},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:dd1f1684be7812aa25b5c543947d32dcd89baf21ae6ed3af29c2afe506a6e7f5","observation_id":"25c2980c-4c69-4249-8eb7-75a68911c38c","resolution":{"observed_at":"2026-05-18T22:32:53.135394Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Patch2Pix: Epi polar-guided pixel-level correspondences","venue":null,"work_id":"4e6b2f0c-5691-49be-bbf7-368f4f7cc3de","year":2021},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:048f5f93896285eb10d6cf3b477311c8caca08a96d93c8540afe60a539f61873","observation_id":"d09c71d3-7438-4dfd-a582-82fdf49bba66","resolution":{"observed_at":"2026-05-18T22:32:53.083837Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"NCNet: Neighbourhood consensus networks for estimating i mage cor- respondences","venue":null,"work_id":"ef0c0586-4325-457f-86f4-5e630cdfb48c","year":2022},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:44dac234dc5d51b284d0d71d84e72f89655d6630f887fb2cce6c4bf6f62c91cf","observation_id":"9b50bd05-1c63-4ace-b0c2-6f6f8162d6c9","resolution":{"observed_at":"2026-05-18T22:31:55.145029Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Learnin g feature descriptors using camera pose supervision","venue":null,"work_id":"74c67ea3-0e5d-4bad-9959-5f816a3394d0","year":2020},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:c92d65d596a9a3b275daeac72bf6511622a955c07a8d2df92925151859152a9f","observation_id":"eac514fb-e545-41c1-9603-fda6c6d0c2ba","resolution":{"observed_at":"2026-05-18T22:32:53.146174Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Revisiting Oxford and Paris: Large-scale image retrieval benchmarkin g","venue":null,"work_id":"46d0956e-31d3-4bfb-8520-00474f15fac5","year":2018},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:daa31fcd36e912b5fbe7a4b14e5c17ff8d38ce811115b588452bfa4ef8f8da29","observation_id":"3bfa8ec4-5d99-4a90-b927-723c213fe31e","resolution":{"observed_at":"2026-05-18T22:32:53.185701Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Learning to Find Good Correspondences","venue":null,"work_id":"8d0c1e0e-f542-48fe-99fd-932de6b7046e","year":2018},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:6be880aaaa08c4b35e8b98c87bbe35cc562c81c4aa97d25b8bcaa9df8c9673a2","observation_id":"b8d65238-7ebd-4ded-a1f9-33b6f2281bdb","resolution":{"observed_at":"2026-05-18T22:32:53.117117Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"OANet: Learning two-view corresponde nces and geometry using order-aware network","venue":null,"work_id":"698d0ec3-adbb-4cb8-8ea8-574a6cb67047","year":2022},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:2374479542c65f1000ee6b8163906f84d320fbe998aa531e86271b378d6b30d7","observation_id":"eb87a686-8ec3-4897-8c2a-247c06b96948","resolution":{"observed_at":"2026-05-18T22:32:53.046164Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"YFCC100M: The new data in multimedia research","venue":null,"work_id":"89028814-3727-453d-9b5c-336224486b1a","year":2016},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:1ac9a9c58057a05fd853cc93c82dcaefe4452816098d9613ebf91a4702cd1351","observation_id":"3f77ed3e-a5ec-44c9-a190-558a5cb38fc0","resolution":{"observed_at":"2026-05-18T22:32:53.080186Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Learning to match features with seeded graph match ing network","venue":null,"work_id":"38bc09e6-efac-493d-80f1-112a4c6892cd","year":2021},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:e6d8631c3e26c40fa6b948e744f33da093ba12dfafa1f4038f51abcc719566e6","observation_id":"60a86359-14f7-4172-acad-b458c50c1418","resolution":{"observed_at":"2026-05-18T22:32:53.109604Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"DenseGAP: Gr aph- structured dense correspondence learning with anchor poin ts","venue":null,"work_id":"35bd15e2-5d4f-4ad7-b677-aeb3a70dd7de","year":2022},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:690055babfd201d6f0566d99326bb4e882332a0f59956262cbf8f8c19a1b7b7e","observation_id":"aa040b38-8aab-4a3b-bbd0-aaa732352bea","resolution":{"observed_at":"2026-05-18T22:32:53.142330Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"ClusterGNN: Cluster-based coarse-to-ﬁne graph neural ne twork for efﬁcient feature matching","venue":null,"work_id":"2f7151b4-1ea4-4bdd-814d-fcb27974a7c1","year":2022},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:c83df7de9e2fe9a5ba1aa117635c74dbbba14aeace3e253e08220808b1f586e1","observation_id":"44be1d70-5285-44fb-865c-586017a02e7f","resolution":{"observed_at":"2026-05-18T22:32:53.105782Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.07193","last_updated":"2024-02-02T10:24:09Z","snapshot_observed_at":"2026-08-06T05:58:29.182448Z","submitted_at":"2023-04-14T15:12:19Z","title":"DINOv2: Learning Robust Visual Features without Supervision","version":2},"cited_work":{"arxiv_id":"2304.07193","doi":"10.48550/arxiv.2304.07193","metadata_source":"pith","pith_arxiv_id":"2304.07193","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DINOv2: Learning Robust Visual Features without Supervision","venue":"cs.CV","work_id":"26b304e5-b54a-4f26-be7e-83299eca52e4","year":2023},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"cited_paper":"/paper/2304.07193","citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:8acebf20483d28abc9c6c8374e8fb9348c19d981754bfae6254a36b2b7e2a066","observation_id":"8465b4dd-8d95-4e80-8b5e-24ed2feafcfe","resolution":{"observed_at":"2026-05-10T13:45:28.110990Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"ScanNet: Richly-annotated 3D reconstruction s of indoor scenes","venue":null,"work_id":"b700edde-23b7-41bb-9604-1b4b77731f8a","year":2017},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:782495d3c3614a9bce3827671531d25ab839e60cfcb4b79c3165a267e343f7da","observation_id":"bec8e3a0-87c0-483c-b4c8-a8946a950396","resolution":{"observed_at":"2026-05-18T22:32:53.066121Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"InLoc: Indoor visual localization with dense matching and view synthesis","venue":null,"work_id":"9511d708-12e8-4a03-bfc4-b0dbb1c382ac","year":2018},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:7e6b73aa7eb28e03260bdd6e7bc35d75a50a6a02d8cbcf4c850e949ea26d24bb","observation_id":"93d824a1-3eac-41ae-a35b-febe3b11b507","resolution":{"observed_at":"2026-05-18T22:32:53.094328Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"DiffGlue: Diffusion-aided image fe ature match- ing","venue":null,"work_id":"b1388b30-b0fe-40d6-bdc0-4664cfd74bc9","year":2024},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:85d10bcd57bfbd52515f17e852f818b057311971f3cf19df765138317d660e90","observation_id":"79a9ec6c-27aa-4f08-88c3-b3f0b7744164","resolution":{"observed_at":"2026-05-18T22:32:53.181000Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Handcrafted outlier detection revisited","venue":null,"work_id":"73cca240-3b5e-4c38-ad58-256834af487a","year":2020},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:620cb65c0aa958ad6dac42bd505ffe2bb82d3ff90caf3c8bec17f565e0e880b0","observation_id":"55ffc986-3bc9-4e92-b884-c4eee0591d0c","resolution":{"observed_at":"2026-05-18T22:31:55.159534Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"ResMatch: R esidual attention learning for feature matching","venue":null,"work_id":"0ebcde72-405a-4540-82d7-831f1c0a443e","year":2024},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:acb89ec14ca387c2f5859d4b0353bac4b9cc7b0810dbb621a29917118ed60f7d","observation_id":"f3ecca1a-ebf7-4e38-8b38-b036fda606c5","resolution":{"observed_at":"2026-05-18T22:32:53.131739Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Benchmarking 6DOF outdoor visual localization in changin g condi- tions","venue":null,"work_id":"c0b47d8c-44cc-4c3e-a6ac-8eb3cd10a61b","year":2018},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:d4ad175ae36e9d1a60a5f92c3e70cd800fc0c3e999689e4e3e47fb996143564f","observation_id":"f9b9caf3-bae7-4f52-8129-d50ec63161f8","resolution":{"observed_at":"2026-05-18T22:32:53.152881Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"F rom coarse to ﬁne: Robust hierarchical localization at large scale","venue":null,"work_id":"e7291145-f4b9-44f0-866d-994850334afe","year":2019},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:376b27867f49fbb592ba21c7aa03756190c65240cd3fb49e11be5ff54dbfad06","observation_id":"9a4846d3-2722-401e-b91d-6833d76f11f4","resolution":{"observed_at":"2026-05-18T22:32:53.071540Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"ASpanFormer: Detector-free image matching with adaptive span Transformer","venue":null,"work_id":"982d3bff-2212-40f3-8063-40554233e258","year":2022},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:ff7c1d5893c1f9456e395e25b6c5bd69220e18fcd5e516bc92d600f2988a40ee","observation_id":"b4afb58f-4532-4960-9471-268852a49ad8","resolution":{"observed_at":"2026-05-18T22:32:53.170310Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"He has published over 50 papers in journals and conferences including IEEE TPAMI/TIP , IJCV , ICCV , and ECCV","venue":null,"work_id":"2e0c96ad-a3d5-415f-a60f-84c21c0ea58a","year":1998},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:6de75308798565b5e3f26d7d8763ff61f200c3f1170cf8e175b8b3bf2ca8f1ff","observation_id":"1c7fde63-20e1-43e7-acf0-adfef1b684c3","resolution":{"observed_at":"2026-05-18T22:32:53.128249Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"He is a co-author of the book An Introduction to the Intelligent Transportation Systems (China Communications Press, Beijing, 2008)","venue":null,"work_id":"8c875f7b-89a3-454e-8819-7ed724aace12","year":2008},"citing_paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-05-10T13:42:04.410771Z"},"links":{"citing_paper":"/paper/2604.13941"},"observation_digest":"sha256:ac002c29986340687edf422207aba5432601b3466682be1972ea44ead408e569","observation_id":"89907305-4537-40bf-ae05-d7641cd20b0a","resolution":{"observed_at":"2026-05-18T22:31:55.152441Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2604.13941","last_updated":"2026-04-15T14:52:10Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T23:01:50.745555Z","submitted_at":"2026-04-15T14:52:10Z","title":"SceneGlue: Scene-Aware Transformer for Feature Matching without Scene-Level Annotation"},"reference_resolution":{"displayed":58,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":2,"verified_fuzzy":56},"total_outbound_references":58},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 58 of 58 outbound references and 0 inbound Pith citation observations for arXiv:2604.13941."}