{"as_of":"2026-08-16T12:31:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:26bd243da7648c07d13bef1bcc56d595a85312f88abe0bbaed73f633c1cc938c","coverage":[{"denominator":57,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":57,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T18:44:53.970024Z","state":"measured"},{"denominator":58,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":58,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-16T06:30:59.297886+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-10T15:34:11.684735Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-11T10:16:06.094774Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"cited_work":{"arxiv_id":"2412.07612","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2412.07612","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Viewdelta: Text-prompted change detection in unaligned im- ages","venue":null,"work_id":"aca80962-fd63-47bc-9b89-7a0b8691cc01","year":2024},"citing_paper":{"arxiv_id":"2604.11402","last_updated":"2026-08-06T11:18:21Z","snapshot_observed_at":"2026-08-15T18:27:59.825815Z","submitted_at":"2026-04-13T12:43:35Z","title":"SCD4VPR: Multi-modal Scene Change Detection for Long-term Visual Place Recognition Database Update","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-05-10T15:34:11.684735Z"},"links":{"cited_paper":"/paper/2412.07612","citing_paper":"/paper/2604.11402"},"observation_digest":"sha256:5ae1d90ce79a78acf1aedf1ce201a98311ce623db14c6f27945e1a3ad40c64e1","observation_id":"f35835f5-af60-4e6a-86c6-c78f8f6ec1d6","resolution":{"observed_at":"2026-05-11T10:16:06.098447Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2412.07612/citation-record","integrity":"/paper/2412.07612/integrity","json":"/paper/2412.07612/citation-record.json","paper":"/paper/2412.07612"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.537956Z","title":"Street-view change detection with deconvolutional networks","venue":null,"work_id":"a0f936ff-f46c-452b-81de-0a0949dcacf6","year":2016},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.791587Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:ba29ee13a54e39aadf8c6fafc7ed82bc4243e6fbd7c0ed087818102b1be9da7a","observation_id":"3857d45c-2cf9-4fb5-aeec-1790b62ee044","resolution":{"observed_at":"2026-08-11T18:44:54.542112Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.528086Z","title":"Transformers in remote sens- ing: A survey","venue":null,"work_id":"d0776b5b-cf60-4c3e-a932-88a8ecdc73a6","year":2023},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.795809Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:c140f53e4729204162f3212f6377b296d61bee88da21dfac2da78d0fdd103624","observation_id":"a251d58a-fff2-477e-bef9-1328d9a30a65","resolution":{"observed_at":"2026-08-11T18:44:54.530832Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.519180Z","title":"Deep learning for change detec- tion in remote sensing: a review","venue":null,"work_id":"3dbff920-fb7d-41d7-917d-e8744372b292","year":2023},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.799124Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:73abb77339162b863a3c816b77026963e6cd3f1b4ee060d09374b0143d23b617","observation_id":"0476aeb9-128b-453f-8c6f-8817a02e96ca","resolution":{"observed_at":"2026-08-11T18:44:54.522277Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.511036Z","title":"A transformer-based siamese network for change detection","venue":null,"work_id":"2b9abf2a-c7d6-4bcb-921b-3e0fcea2d3a8","year":2022},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.803120Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:292e59b4b42899f0c63bff43f3673e6d4d07526d11b48efa7df8556ee9f4a8f0","observation_id":"38e606e8-044d-450d-a352-32ae5890774a","resolution":{"observed_at":"2026-08-11T18:44:54.513777Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.501114Z","title":"End-to- end object detection with transformers, 2020","venue":null,"work_id":"756ae248-4963-4eee-b7c4-61c172921391","year":2020},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.806911Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:a0185973ca791ab123645cbff2f5eb47ab742fd1d0a9cddea8a4372ca3ade968","observation_id":"4bb5b3bd-102a-4dda-a6dd-504560439c45","resolution":{"observed_at":"2026-08-11T18:44:54.504371Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.492258Z","title":"A spatial-temporal attention- based method and a new dataset for remote sensing image change detection","venue":null,"work_id":"c820cf95-9315-4d20-bf49-54d9039f2c94","year":2020},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.809928Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:6054e607ef9bca5d321950f701f16eabf150d69ba7cd5d417eadea9ac2e0de45","observation_id":"256c2f8e-c5ab-4ae3-85e5-1635c59042a0","resolution":{"observed_at":"2026-08-11T18:44:54.495244Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.483419Z","title":"Remote sensing im- age change detection with transformers","venue":null,"work_id":"5df1dc8e-39fd-4ff7-af97-c78903ec229c","year":2021},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.812891Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:c837cea7fa16c44d906c0a686235cbf4c784b428b76952ebfc8b8189f2df4762","observation_id":"022b5b28-61f7-4e30-8767-bcb8d3213362","resolution":{"observed_at":"2026-08-11T18:44:54.486327Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.474003Z","title":"Changemamba: Remote sensing change de- tection with spatio-temporal state space model","venue":null,"work_id":"1ceaf72f-c681-46fb-b5e5-4365b9d721ef","year":2024},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.815343Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:6047f46429057a8b37fdc7be69e8069ea98845745e9d43fa554b5b02d6f8bd5f","observation_id":"d832fbc2-be8d-4211-9fa8-613f882bb7b2","resolution":{"observed_at":"2026-08-11T18:44:54.477194Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.464448Z","title":"Dr- tanet: Dynamic receptive temporal attention network for street scene change detection","venue":null,"work_id":"54d4b9c9-6913-4141-82f2-6ac1cc585ad6","year":2021},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.817744Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:b371698c3068341ad607ba406d460d75bb0cf4f9dad9025a3bb30b85debd05ea","observation_id":"845a8202-5b7a-4332-bb50-488b86b4ca47","resolution":{"observed_at":"2026-08-11T18:44:54.467865Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.454765Z","title":"When deep learning meets metric learning: Remote sensing image scene classification via learning discrimina- tive cnns","venue":null,"work_id":"c60c3a00-86c8-40d7-b9f6-b45dc19f8131","year":2018},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.820591Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:e7a529086147a8d30b870beef6e33eebe243b64c4bd4c6222aa01befd4948b19","observation_id":"a778831f-3c93-4b57-a54c-88ab36fce773","resolution":{"observed_at":"2026-08-11T18:44:54.458318Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.445110Z","title":"Change detection methods for remote sensing in the last decade: A comprehensive review","venue":null,"work_id":"0cafe25d-6d52-4ec7-8948-32a886d074db","year":2024},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.823465Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:b09f152f365e92014dcb9a66ea0a77884cc12546734f56ff4f9097cd2789443d","observation_id":"6fc41144-08ce-4e97-a2d4-606658deebcc","resolution":{"observed_at":"2026-08-11T18:44:54.448587Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.435068Z","title":"Re- gion filling and object removal by exemplar-based image in- painting","venue":null,"work_id":"0899ceae-c694-4280-9c73-84c937a1bed8","year":2004},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.826226Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:24796a960a66abeb25a139dc7b06ccd16d2e817b8810b8af67c014d2fc36c750","observation_id":"e556daee-8133-4954-a770-62afd21ef0b9","resolution":{"observed_at":"2026-08-11T18:44:54.438654Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.424837Z","title":"Fully convolutional siamese networks for change detection, 2018","venue":null,"work_id":"10f25182-4dac-4291-ba98-f5620406f44a","year":2018},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.829817Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:63a01e059f867c6713612055c6586596387908f2f90a08558477f0b496917259","observation_id":"22ac15c4-b09a-4fd6-b610-b2919c6f1381","resolution":{"observed_at":"2026-08-11T18:44:54.428048Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-08-16T09:25:53.087782Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-08-11T18:44:53.833692Z","title":"An image is worth 16x16 words: Transformers for image recognition at scale","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.833692Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:9cbf30f3ccacbeb7027b53160526422343afa5fd7aa543e575b3ea19c62d5c25","observation_id":"1e7447d9-a718-4423-b051-53711b5effd3","resolution":{"observed_at":"2026-08-11T18:44:53.833692Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.415307Z","title":"Foreground gating and background refining network for surveillance object detec- tion","venue":null,"work_id":"165edfe9-89a2-4d7e-923c-fc58e609594a","year":2019},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.836561Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:e9b6122a54cf6020137e47915439712d4848bf6836d89be6d6a32f56381c241e","observation_id":"ee47717d-bfd4-4103-8bc2-9387abf61938","resolution":{"observed_at":"2026-08-11T18:44:54.418724Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.406230Z","title":"Dat- acomp: In search of the next generation of multimodal datasets","venue":null,"work_id":"3bfcd919-5eea-4948-8dd4-9e4b32867253","year":2024},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.839735Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:2ccd5722638620991c76369a51ca0339cc08e7405e44cbc160993da1277cced1","observation_id":"697e74d7-f77b-4477-a175-ba6bd67c8d3a","resolution":{"observed_at":"2026-08-11T18:44:54.408913Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.396842Z","title":"A framework for the detection and attribution of biodiversity change","venue":null,"work_id":"0299ce17-41fa-465a-b4be-39f9c928483c","year":2023},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.843160Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:bd3fc686ff1c5aa0d4470b06b1136d2d1b8c3cd96bc1999008e9012eb74819a9","observation_id":"ea56e77a-599b-4637-8371-2c6a940bca37","resolution":{"observed_at":"2026-08-11T18:44:54.400089Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1911.09296","last_updated":"2019-11-21T05:30:13Z","snapshot_observed_at":"2026-08-13T19:55:55.116447Z","submitted_at":"2019-11-21T05:30:13Z","title":"xBD: A Dataset for Assessing Building Damage from Satellite Imagery","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1911.09296","snapshot_observed_at":"2026-08-11T18:44:53.845800Z","title":"Patel, Richard Hosfelt, Sandra Sajeev, Eric T","venue":null,"work_id":null,"year":1911},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.845800Z"},"links":{"cited_paper":"/paper/1911.09296","citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:277781081c290c43a5be3dd317396a364d863879162ab7996a778ae90fb29e71","observation_id":"f03068b1-4f9c-4ef6-ae2f-27e321275507","resolution":{"observed_at":"2026-08-11T18:44:53.845800Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.387402Z","title":"Fast flood extent mon- itoring with sar change detection using google earth engine","venue":null,"work_id":"6c6e0c10-7c36-4ec0-9793-fdddb24e7922","year":2023},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.849060Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:2b340ce3d5f3fc35932e3dc00d165b1d777b6b221c225eeac990a88f0590cfa0","observation_id":"e05f7715-3e6e-4843-a963-2314df80bb17","resolution":{"observed_at":"2026-08-11T18:44:54.390608Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.376864Z","title":"A long-term analysis of urbanization pro- cess, landscape change, and carbon sources and sinks: A case study in china’s yangtze river delta region","venue":null,"work_id":"45cd4399-1dea-4cdf-bb79-2de2327e523c","year":2017},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.852062Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:a100c7be973c924f560a387045ee244803e54e2b4b8ac8c4421c02d326fdb79a","observation_id":"99cda395-396f-4780-a6ff-87b0dffaeda2","resolution":{"observed_at":"2026-08-11T18:44:54.379991Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.366440Z","title":"Image repairing: Robust im- age synthesis by adaptive nd tensor voting","venue":null,"work_id":"bfd1ada9-7d5b-453d-9e36-4a7cf19f7224","year":2003},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.855154Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:d5abacd0e9b3da92e52bf80973ba64154c170a160ac8e2229e85be789a60fc5d","observation_id":"b8a661c0-22c0-420d-9c0f-e548e542bbe0","resolution":{"observed_at":"2026-08-11T18:44:54.369859Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.357166Z","title":"A survey on deep learning-based change detection from high- resolution remote sensing images","venue":null,"work_id":"7d2fe780-d37f-415d-bc52-bcfc795e5d62","year":2022},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.858065Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:36db5b6d0201fa72f3567a518707b9b87ed54d9deec16137b827b04b482f7435","observation_id":"ce7bc80e-2a93-4b5e-9b18-a4f307b0ab98","resolution":{"observed_at":"2026-08-11T18:44:54.359920Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.347039Z","title":"T-rex2: Towards generic object detec- tion via text-visual prompt synergy, 2024","venue":null,"work_id":"8105fa65-a5d1-4178-8e0e-3d1131c5aae7","year":2024},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.861562Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:733294ac1980a1ab92835c11006d6404524e62be3e6b5530dd16e2d7a02bd2e4","observation_id":"8fc65936-6c36-4e44-846b-ab83b0327211","resolution":{"observed_at":"2026-08-11T18:44:54.350813Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.335955Z","title":"Towards generalizable scene change detection","venue":null,"work_id":"812ae472-5b19-42e9-9161-a9745b90f2f2","year":2025},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.864746Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:74738d4afc58b3a469143beb0cc8fd05035ccbc6014a73813440aace4faa54ab","observation_id":"f4d4902c-7753-4779-a05c-0937b1856fc6","resolution":{"observed_at":"2026-08-11T18:44:54.339181Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1412.6980","last_updated":"2017-01-30T01:27:54Z","snapshot_observed_at":"2026-08-14T18:51:16.666127Z","submitted_at":"2014-12-22T13:54:29Z","title":"Adam: A Method for Stochastic Optimization","version":9},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1412.6980","snapshot_observed_at":"2026-08-11T18:44:53.867811Z","title":"Adam: A method for stochastic opti- mization","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.867811Z"},"links":{"cited_paper":"/paper/1412.6980","citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:f3b6a7cf11f6ed505d2bf2412481ddd0cbb4dc017858c5c1d897d95d96d2d1f8","observation_id":"ae1fa982-2cce-413d-bfdb-b1a7729c31b0","resolution":{"observed_at":"2026-08-11T18:44:53.867811Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.325447Z","title":"Segment any- thing","venue":null,"work_id":"8efce28f-5f4e-49ec-8422-ce61a6a8c7d6","year":2023},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.871713Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:ac6adfbf0f698c7a49c6c45178a17def904bd72af5ec8b27af772b0578412eb5","observation_id":"8c92b88d-fa9e-4c66-91a7-6a70c968ec63","resolution":{"observed_at":"2026-08-11T18:44:54.328647Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.312996Z","title":"Tran- sunetcd: A hybrid transformer network for change detec- tion in optical remote-sensing images","venue":null,"work_id":"c3f6ad7a-0040-4744-9426-ff205ca58a3c","year":2022},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.874956Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:b02d15c68d54dc719573280b23399328997a2404174c6c0fc67bfa354cce6c72","observation_id":"0d897b41-5a54-4a92-9fc5-9c8a20d7ac89","resolution":{"observed_at":"2026-08-11T18:44:54.316661Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.16850","last_updated":"2025-03-04T02:16:30Z","snapshot_observed_at":"2026-08-12T22:37:52.640972Z","submitted_at":"2024-09-25T11:55:27Z","title":"Robust Scene Change Detection Using Visual Foundation Models and Cross-Attention Mechanisms","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.16850","snapshot_observed_at":"2026-08-11T18:44:53.878480Z","title":"Robust scene change detection using visual foun- dation models and cross-attention mechanisms","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.878480Z"},"links":{"cited_paper":"/paper/2409.16850","citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:bbb54301a3d420dd48e07475c4eaeaf16ca4baf6bd940c5f2dd58d2a6c56ab0f","observation_id":"f428a05b-b721-4382-a24a-0628c96fd59a","resolution":{"observed_at":"2026-08-11T18:44:53.878480Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:53.882617Z","title":"Llava-next: Im- proved reasoning, ocr, and world knowledge, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.882617Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:1aa42f3c6c0e9b96e89f5ba39fd0ab9cca3d52dd32207fca4720416b01f19e75","observation_id":"2629f5c9-9a88-497e-a70d-a41174423206","resolution":{"observed_at":"2026-08-11T18:44:53.882617Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.05499","last_updated":"2024-07-19T06:00:41Z","snapshot_observed_at":"2026-07-06T15:00:58.804337Z","submitted_at":"2023-03-09T18:52:16Z","title":"Grounding DINO: Marrying DINO with Grounded Pre-Training for Open-Set Object Detection","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.05499","snapshot_observed_at":"2026-08-11T18:44:53.886317Z","title":"Grounding dino: Marrying dino with grounded pre-training for open-set object detection","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.886317Z"},"links":{"cited_paper":"/paper/2303.05499","citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:cbb56b2ea1da4967339739dab20950435ca76f2dd0960f3ddc03e872336c52ff","observation_id":"ce69971c-b84f-40e6-ae7a-4d47ab488c3b","resolution":{"observed_at":"2026-08-11T18:44:53.886317Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.297325Z","title":"Global effects of land use on local terrestrial biodiversity","venue":null,"work_id":"384f6b31-f386-4d73-8e24-90f4dd44bcda","year":2015},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.890242Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:dc17d7037f7f62e35fd81122eb32d2f597e4c168df2f795ab95c087967e32260","observation_id":"57343970-541d-4af0-970c-cdd989ae442d","resolution":{"observed_at":"2026-08-11T18:44:54.300574Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.285406Z","title":"Dinov2: Learning robust visual features with- out supervision, 2024","venue":null,"work_id":"4b31ed66-4194-469f-a93a-608f0d4488e3","year":2024},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.893452Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:05d58a34f160268c49a254f7ac3c7be2cb30c35e62a8aa6564009457541911b9","observation_id":"28749197-2541-4c48-9c70-65626de61fe2","resolution":{"observed_at":"2026-08-11T18:44:54.289640Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.275183Z","title":"Detecting urban changes with recurrent neural networks from multi- temporal sentinel-2 data, 2019","venue":null,"work_id":"7ad68d19-cdc0-4393-a394-3d36dfa058ce","year":2019},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.895770Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:a975132e05e745c91ca7f630c77294c2dbf607199043fb57aee8e9214259415c","observation_id":"a8a7d0b1-23d5-410c-94d6-398374e54647","resolution":{"observed_at":"2026-08-11T18:44:54.278761Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:53.898293Z","title":"Learning transferable visual models from natural language supervi- sion","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.898293Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:0799a91317f5a43a0b061c8d3bd9a450414b09ba0a3a36f37f2e6c28a72095c7","observation_id":"81e1da03-66e6-4045-9d6f-d262e458d13c","resolution":{"observed_at":"2026-08-11T18:44:53.898293Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.259684Z","title":"Zero: Memory optimizations toward training trillion parameter models","venue":null,"work_id":"9383e11c-9f65-43c7-8e9a-37941512a04e","year":2020},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.901122Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:b66dab3a24a75b8006e040cfb4ab5f348b068cacf603f4dc393d203b4795c11c","observation_id":"ae4c44ae-8891-4769-b5da-a79e30e71e12","resolution":{"observed_at":"2026-08-11T18:44:54.262760Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:53.904559Z","title":"Sam 2: Segment anything in images and videos,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.904559Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:70e8e046910dbd249bbd3a135c788ef20818dac3e9ba9903e8ab20f8b91cb595","observation_id":"1d15faca-9aea-43de-bcec-be540a2814d6","resolution":{"observed_at":"2026-08-11T18:44:53.904559Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.243640Z","title":"Grounding dino 1.5: Ad- vance the ”edge” of open-set object detection, 2024","venue":null,"work_id":"bac0ac1b-01a8-4817-baf5-d2ec9bd9d646","year":2024},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.907280Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:d8f8fb2545b91ff8e84987b418be850852843db26c23ab80a301149be6936d0a","observation_id":"15af9905-d2f1-4b5f-93ca-fc2113c6a4b7","resolution":{"observed_at":"2026-08-11T18:44:54.246911Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:53.910547Z","title":"Grounded sam: Assembling open-world models for diverse visual tasks,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.910547Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:06aa3a09479b0e244df0f0dd538aba06e01b5acf7aaa10bb958c2815e07123f0","observation_id":"a9bba190-e3aa-46cb-961d-0619fcc8be03","resolution":{"observed_at":"2026-08-11T18:44:53.910547Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.224750Z","title":"The change you want to see","venue":null,"work_id":"ebea8a0b-6019-47f4-87ed-55e23efc0321","year":2023},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.913119Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:911ef4fb695087ac90f3c4df79d1a656e28890bba29dd07836f66b1a21553dbb","observation_id":"32c8fd15-edc2-4b78-8a9b-fa2ea3a93be7","resolution":{"observed_at":"2026-08-11T18:44:54.227645Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.216372Z","title":"The change you want to see (now in 3d)","venue":null,"work_id":"d1780e9c-c867-42f9-b155-e6164786955c","year":2023},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.915559Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:bc4ad3ed57607fcfff945b3ee530fc86092618c9d82ff55ef05b6078d2a9c7ce","observation_id":"43539cfd-54f3-4e54-aab0-8f8681fd885c","resolution":{"observed_at":"2026-08-11T18:44:54.219248Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.207033Z","title":"Using high-resolution satellite images for post-earthquake building damage assessment: a study fol- lowing the 26 january 2001 gujarat earthquake","venue":null,"work_id":"4403dd72-ed47-4cba-8d8f-90ea7347689a","year":2001},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.917925Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:c7fc2b79c8bf629615fdb48f9e5c9fbc05dabe74abcacb5054fb2318882302c9","observation_id":"b20b07b7-1b65-4427-8968-b710c679dee3","resolution":{"observed_at":"2026-08-11T18:44:54.209982Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.197755Z","title":"Weakly supervised silhouette-based semantic scene change detec- tion, 2022","venue":null,"work_id":"c9527489-cea6-4492-b494-1d0427afdb29","year":2022},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.921614Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:a826465926ebe25d0aec5fca03d2280cbee635f308135b01af5fac326856578d","observation_id":"a69388bc-3416-427d-88d1-73e74ddc8fe7","resolution":{"observed_at":"2026-08-11T18:44:54.200701Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.188221Z","title":"A deeply supervised attention metric-based network and an open aerial image dataset for remote sensing change detection","venue":null,"work_id":"1fdf6be3-aeb9-471f-a27b-0eb45d6b661e","year":2021},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.924862Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:4e978a25cc8b52d6f6129ab0f7f9eeeddd5041afdb1b4972835753ce568d641c","observation_id":"da122b78-46b0-41d1-b985-437c12155b32","resolution":{"observed_at":"2026-08-11T18:44:54.191419Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.178632Z","title":"Change detection based on artificial intel- ligence: State-of-the-art and challenges","venue":null,"work_id":"c9d7d2ac-bacc-49af-a93d-04eb202e9756","year":2020},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.927472Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:a0df3de41a024a08cfe3d1c2bc492868bdd5af24731f36186cd1240c01af0df5","observation_id":"230e1ccf-2431-4a82-92e1-deb2e38ab2a1","resolution":{"observed_at":"2026-08-11T18:44:54.182225Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.166657Z","title":"Review article digital change detection techniques using remotely-sensed data","venue":null,"work_id":"17200ad6-0611-4dbf-b6fb-04a34f05c91d","year":1989},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.930083Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:cd57e4fc53dee4515163822ebb1fa5194e022286300f9a1cb052441add65a588","observation_id":"bfbddc6d-614b-4c68-b220-98b2357b94a4","resolution":{"observed_at":"2026-08-11T18:44:54.171042Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.154878Z","title":"Automatic post-disaster damage mapping using deep-learning tech- niques for change detection: Case study of the tohoku tsunami","venue":null,"work_id":"c0ced3df-3a96-48c8-9125-8d244a7bb3f0","year":2019},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.932707Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:491f946a720f86cba4bbe562c1a6df7646f62930d23cd15a0569cb2e2aca2054","observation_id":"9d72154b-641b-441c-b625-8a4d0fe6c86d","resolution":{"observed_at":"2026-08-11T18:44:54.158617Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.15389","last_updated":"2023-03-27T17:02:21Z","snapshot_observed_at":"2026-08-16T07:24:08.156932Z","submitted_at":"2023-03-27T17:02:21Z","title":"EVA-CLIP: Improved Training Techniques for CLIP at Scale","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.15389","snapshot_observed_at":"2026-08-11T18:44:53.935741Z","title":"Eva-clip: Improved training techniques for clip at scale","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.935741Z"},"links":{"cited_paper":"/paper/2303.15389","citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:2bd08c76ac0818e4fa2a725ee901d2eaa1ed5e0ea0f960cbd8c182bac0796aa7","observation_id":"22b4f897-54fd-4c81-8aa1-1271a4b495cd","resolution":{"observed_at":"2026-08-11T18:44:53.935741Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2109.07161","last_updated":"2021-11-11T01:58:26Z","snapshot_observed_at":"2026-08-13T18:10:58.709543Z","submitted_at":"2021-09-15T08:54:29Z","title":"Resolution-robust Large Mask Inpainting with Fourier Convolutions","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2109.07161","snapshot_observed_at":"2026-08-11T18:44:53.938502Z","title":"Resolution-robust large mask inpainting with fourier convolutions","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.938502Z"},"links":{"cited_paper":"/paper/2109.07161","citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:8f57f9146a4c76b7a98f89b0caee4cd08d99834987bcea0e4c74fd600c3f5a36","observation_id":"b3983b6c-3cbd-46a3-af19-d021be724e59","resolution":{"observed_at":"2026-08-11T18:44:53.938502Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.07557","last_updated":"2023-02-14T07:40:09Z","snapshot_observed_at":"2026-08-13T15:23:24.342529Z","submitted_at":"2022-06-15T14:16:30Z","title":"How to Reduce Change Detection to Semantic Segmentation","version":2},"cited_work":{"arxiv_id":"2206.07557","doi":null,"metadata_source":"pith","pith_arxiv_id":"2206.07557","snapshot_observed_at":"2026-08-11T18:44:53.995006Z","title":"How to Reduce Change Detection to Semantic Segmentation","venue":"cs.CV","work_id":"3e7e8fb7-f8cf-49de-9654-b52f4865edc3","year":2022},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.942046Z"},"links":{"cited_paper":"/paper/2206.07557","citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:3e91300f22aea1cfd410740bd9e2a00635f10749e77cbe1b120bd3d16445937d","observation_id":"1cf7aa3b-4954-4c1d-ba86-831ca381cebe","resolution":{"observed_at":"2026-08-11T18:44:54.000608Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.143535Z","title":"Building damage detection using u-net with attention mechanism from pre-and post-disaster remote sensing datasets","venue":null,"work_id":"6fb6baf5-24c0-45a8-b76e-f9c8da3e2db0","year":2021},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.946077Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:aa8ff6c46e392f14bd696ad4e742b59cfcbb86bc80a9b8897984b19397c30b6f","observation_id":"6bdc0e16-5a8d-40bc-b557-64736d1315ac","resolution":{"observed_at":"2026-08-11T18:44:54.147151Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.132996Z","title":"Robust image forgery detection over online social network shared 10 images","venue":null,"work_id":"16a867d7-970e-4314-ae5e-909eb3a69b3a","year":2022},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.949611Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:bcc35d937927e34584ee1277af9b84f2a3c088d01c89b499df984267eb15ec7d","observation_id":"5abfdcc1-d25c-4c42-b575-b2717fc1134a","resolution":{"observed_at":"2026-08-11T18:44:54.136535Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.121792Z","title":"Semantic change detection with asymmetric siamese networks, 2021","venue":null,"work_id":"b0e4ec2b-8cb7-46eb-88b6-243640e0f8d5","year":2021},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.953592Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:720b85bf8d94099ba002f46318cb1ef1159dcb869c0870606acef605a43b13aa","observation_id":"0d409a2d-985e-4f23-91ff-53c556a01a10","resolution":{"observed_at":"2026-08-11T18:44:54.124683Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.111462Z","title":"Pedestrian be- havior modeling from stationary crowds with applications to intelligent surveillance","venue":null,"work_id":"5524fd74-8758-4ddc-9f39-7d27f84303c8","year":null},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.957437Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:63024f29fe223f8ab08f3dd15d75624fd7bf76edf8118fa279082297e07ebb34","observation_id":"5f6b2eec-7a4c-4e5f-88c4-36e030c39704","resolution":{"observed_at":"2026-08-11T18:44:54.114451Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:53.960709Z","title":"Sigmoid loss for language image pre-training","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.960709Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:4e959f7edb6c8c67f895951d0a5a62214ecf16a01dc51fe86e30f2c482eca6f6","observation_id":"c666b8c1-d5dc-44e6-8d21-784a13b3c7f4","resolution":{"observed_at":"2026-08-11T18:44:53.960709Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.096452Z","title":"Swinsunet: Pure transformer network for remote sensing im- age change detection","venue":null,"work_id":"c36df27f-0902-44f7-b6d3-cc51d268b723","year":2022},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.963499Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:3e850001d514e4d62c0b0363136c323ea53e5211df01fef62396519a92d799d6","observation_id":"fd09d01b-53bd-4b46-ab44-9577305fb1f3","resolution":{"observed_at":"2026-08-11T18:44:54.099688Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.087486Z","title":"Changen2: Multi-temporal re- mote sensing generative change foundation model","venue":null,"work_id":"4281b7c2-9003-4f0b-896d-ef97ceb65fbe","year":2024},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.966439Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:a4a2617d683ef7eb2658e866a0bc16029f725416e710b681b21169adc91bb941","observation_id":"5af3397e-4d18-4bfb-8f32-1ec22595e82a","resolution":{"observed_at":"2026-08-11T18:44:54.090502Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T18:44:54.076441Z","title":"A review of multi- class change detection for satellite remote sensing imagery","venue":null,"work_id":"72895048-d1eb-4e43-85c6-3d1d1962ada3","year":2024},"citing_paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning","version":3},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-11T18:44:53.970024Z"},"links":{"citing_paper":"/paper/2412.07612"},"observation_digest":"sha256:7d43d8ece325a8bb3f209df793fa6917af9a96ecee071812eb4e79aea18de9de","observation_id":"fae30ec8-e54a-4578-b0e1-adcfbc154416","resolution":{"observed_at":"2026-08-11T18:44:54.079967Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2412.07612","last_updated":"2025-08-13T16:49:44Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-11T18:41:24.356227Z","submitted_at":"2024-12-10T15:51:17Z","title":"ViewDelta: Scaling Scene Change Detection through Text-Conditioning"},"reference_resolution":{"displayed":57,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":12,"verified_exact":1,"verified_fuzzy":44},"total_outbound_references":57},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"thesis":"As of 16 August 2026, this Paper Citation Record lists 57 of 57 outbound references and 1 inbound Pith citation observation for arXiv:2412.07612."}