{"as_of":"2026-08-17T15:56:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:0abffe82aa3cbf492bb9653d2606a09baec9b70b951789245ef487b40344666d","coverage":[{"denominator":61,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":61,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T23:59:35.792565Z","state":"measured"},{"denominator":63,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":63,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-17T06:30:58.91139+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T10:52:48.564988Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-12T12:16:16.734455Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.04611","snapshot_observed_at":"2026-08-01T11:10:14.106468Z","title":"Bridgedepth: Bridging monocular and stereo reasoning with latent alignment,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.19986","last_updated":"2026-07-22T10:19:43Z","snapshot_observed_at":"2026-08-12T17:06:57.150430Z","submitted_at":"2026-07-22T10:19:43Z","title":"STEREOFLOW: Progressive Stereo Matching with StereoDiT and Transition Flow Matching","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-01T11:10:14.106468Z"},"links":{"cited_paper":"/paper/2508.04611","citing_paper":"/paper/2607.19986"},"observation_digest":"sha256:480523c0faca456160a2e43ee274cfbbe3f39db8e0cfa15fd3f36ea140be650b","observation_id":"ec333798-3ec4-48c7-bc61-1845f70e7550","resolution":{"observed_at":"2026-08-01T11:10:14.106468Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"cited_work":{"arxiv_id":"2508.04611","doi":"10.48550/arxiv.2508.04611","metadata_source":"pith","pith_arxiv_id":"2508.04611","snapshot_observed_at":"2026-08-12T12:16:16.734455Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","venue":"cs.CV","work_id":"a713c35a-96e9-4e2d-8deb-84e5f5aba3b9","year":2025},"citing_paper":{"arxiv_id":"2608.11075","last_updated":"2026-08-11T15:37:49Z","snapshot_observed_at":"2026-08-15T17:03:03.322415Z","submitted_at":"2026-08-11T15:37:49Z","title":"Static in Frames, Dynamic in Events: Rethinking Features in Event Cameras as Motion Cues","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-12T10:52:48.564988Z"},"links":{"cited_paper":"/paper/2508.04611","citing_paper":"/paper/2608.11075"},"observation_digest":"sha256:d384362d136e8aef91f104f603c9357b5d303e1191cdfd738e41cec8c563d12c","observation_id":"7acd7e3f-c87a-4f5b-a539-0806c6f15799","resolution":{"observed_at":"2026-08-12T10:53:37.123878Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2508.04611/citation-record","integrity":"/paper/2508.04611/integrity","json":"/paper/2508.04611/citation-record.json","paper":"/paper/2508.04611"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:44.238858Z","title":"Multi- view depth estimation by fusing single-view depth prob- ability with multi-view geometry","venue":null,"work_id":"0b56f17b-26f6-4433-b806-71ad966e19ee","year":2022},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:29.889459Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:77c0ea62957c09d4b9a8e19f3ad457ab00afb9de3ea768cbf7c8c5286a526280","observation_id":"afcd6446-2bfb-4dbc-8018-46056d4fa50d","resolution":{"observed_at":"2026-08-05T23:59:44.289388Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.04472","last_updated":"2025-05-07T11:31:28Z","snapshot_observed_at":"2026-08-15T02:38:58.371690Z","submitted_at":"2024-12-05T18:59:58Z","title":"Stereo Anywhere: Robust Zero-Shot Deep Stereo Matching Even Where Either Stereo or Mono Fail","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.04472","snapshot_observed_at":"2026-08-05T23:59:30.000269Z","title":"Stereo anywhere: Robust zero-shot deep stereo matching even where either stereo or mono fail","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:30.000269Z"},"links":{"cited_paper":"/paper/2412.04472","citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:182b997e7be62b3a80e1b70b8b198bca5b797a6c34280647f672ecf98b2b6a73","observation_id":"668515b1-6458-4cff-8030-58ffc0079477","resolution":{"observed_at":"2026-08-05T23:59:30.000269Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.02073","last_updated":"2025-04-21T12:09:08Z","snapshot_observed_at":"2026-08-02T06:32:26.688405Z","submitted_at":"2024-10-02T22:42:20Z","title":"Depth Pro: Sharp Monocular Metric Depth in Less Than a Second","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.02073","snapshot_observed_at":"2026-08-05T23:59:30.108081Z","title":"Depth pro: Sharp monocular metric depth in less than a second","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:30.108081Z"},"links":{"cited_paper":"/paper/2410.02073","citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:a3aaaa943c644fd849d02011667080e8c89caba20fbe8346b749518d96c79f09","observation_id":"e9be24d5-f1a7-4565-8999-db84e6f16bd0","resolution":{"observed_at":"2026-08-05T23:59:30.108081Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:30.198456Z","title":"Pyramid stereo matching network","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:30.198456Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:95d9fb56f6fe1f1f3e42c588f1c110602abab1e707fdaeb241da43585cf7c209","observation_id":"a24c22ff-1fb1-4e5e-b1d7-5abd2657b5f0","resolution":{"observed_at":"2026-08-05T23:59:30.198456Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.12426","last_updated":"2025-01-05T13:20:07Z","snapshot_observed_at":"2026-08-17T15:20:44.039347Z","submitted_at":"2024-11-19T11:26:21Z","title":"Motif Channel Opened in a White-Box: Stereo Matching via Motif Correlation Graph","version":2},"cited_work":{"arxiv_id":"2411.12426","doi":null,"metadata_source":"pith","pith_arxiv_id":"2411.12426","snapshot_observed_at":"2026-08-05T23:59:36.218190Z","title":"Motif Channel Opened in a White-Box: Stereo Matching via Motif Correlation Graph","venue":"cs.CV","work_id":"41f9dc86-32d4-43ef-b544-633ee5c6253b","year":2024},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:30.285602Z"},"links":{"cited_paper":"/paper/2411.12426","citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:7d00b4a36a79312e25053ddfd61472ee6599e0d6ed54972db9a6342f85ddf5cf","observation_id":"b4ab422a-6779-4676-988a-f7496d937a9a","resolution":{"observed_at":"2026-08-05T23:59:36.374186Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:30.348568Z","title":"Monster: Marry monodepth to stereo unleashes power","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:30.348568Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:8adc107fba791d104da2add6075d5451f6b6c26a42718a0948fa7d07f83b8fe5","observation_id":"8b84238e-ba0f-4182-a057-7b79fc40b69c","resolution":{"observed_at":"2026-08-05T23:59:30.348568Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:44.138078Z","title":"Learning depth with convolutional spatial propagation network","venue":null,"work_id":"8278ff01-d104-4212-a7f9-ec8a536913ee","year":null},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:30.471146Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:cd4b07c25ace47dcc1f6ad113e22b4715ac583f24e0b1c67f809addaaf5cd896","observation_id":"5633c501-51f7-4781-ae8a-f5aa602bea56","resolution":{"observed_at":"2026-08-05T23:59:44.187718Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:43.946572Z","title":"Hierarchical neural architecture search for deep stereo matching","venue":null,"work_id":"8d2a9c9b-bb0f-4427-a537-24033468762a","year":2020},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:30.580375Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:c10b07659b88735b19c90a3de2f05d2e3fb1e104f2356614e2cd2eb80239488d","observation_id":"24a23c28-2059-4bec-bf01-95fef1cc7549","resolution":{"observed_at":"2026-08-05T23:59:44.024903Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-08-16T09:25:53.087782Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-08-05T23:59:30.671480Z","title":"An image is worth 16x16 words: Trans- formers for image recognition at scale","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:30.671480Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:cd1fbed516d0bf05ebc0c88948fabdbfa5451527fa86b8897b4aace8a62a9096","observation_id":"af656bd6-9871-443f-8d8b-9080dec5a89b","resolution":{"observed_at":"2026-08-05T23:59:30.671480Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:43.756231Z","title":"Depth map prediction from a single image using a multi-scale deep net- work","venue":null,"work_id":"f25c8ff1-a3aa-4f5d-a371-e8a1e387fa30","year":2014},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:30.742517Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:913cc8a9581af7def0c7735394b9d7feef66ad493cac8f1c64dab2d07f50bc16","observation_id":"ab28d294-e467-4355-b4fc-86dbe9d74f0c","resolution":{"observed_at":"2026-08-05T23:59:43.836778Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:30.816464Z","title":"Geowiz- ard: Unleashing the diffusion priors for 3d geometry esti- mation from a single image","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:30.816464Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:c7aaeda39770e8e69f2c642cec34a27133854e7ddb4ad48a954392dd5c65b998","observation_id":"66b03683-5e7d-4543-a2c6-675f8e0d7ffd","resolution":{"observed_at":"2026-08-05T23:59:30.816464Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:43.624656Z","title":"Are we ready for autonomous driving? the kitti vision benchmark suite","venue":null,"work_id":"f7ace400-5c02-431f-9db8-c075c0609721","year":2012},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:30.881125Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:bfe0b1df63193984d47b8dba8fb7b7a6f95adb58f9d091a696b491243cf7564a","observation_id":"d1649be2-f62c-4fab-80bb-508af28b4a08","resolution":{"observed_at":"2026-08-05T23:59:43.678559Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:43.473174Z","title":"Neural markov random field for stereo matching","venue":null,"work_id":"41e8582e-64d8-4723-bba1-a01c17b87bf8","year":2024},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:30.949394Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:9121ad37c76aad4560e68c08022a29bc295bb416c942f6cf32ba3f92e7f30f8a","observation_id":"174cc6db-30ac-4001-8ecb-c14ace44c070","resolution":{"observed_at":"2026-08-05T23:59:43.528874Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:43.326943Z","title":"Context-enhanced stereo transformer","venue":null,"work_id":"aaa16501-92af-4ab7-ac47-3233655c96ad","year":2022},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:31.008882Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:6514cbe794b1170009c6e20a91421f0e0a3016e07201737435f15f4ee5ff0777","observation_id":"f27056fa-ac31-4b61-be96-0cd1ab2a4861","resolution":{"observed_at":"2026-08-05T23:59:43.408874Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:43.218070Z","title":"Group-wise correlation stereo network","venue":null,"work_id":"1598e8fb-5751-42f1-be8c-453a2e1611bf","year":2019},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:31.080310Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:3eb6194cf5777672619779f7731dc7ecfed39bc7605f3dcc63498766aed8c1d5","observation_id":"be31e709-0d82-49b3-a671-68f2ff286556","resolution":{"observed_at":"2026-08-05T23:59:43.264389Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18124","last_updated":"2025-01-18T20:14:54Z","snapshot_observed_at":"2026-08-16T13:15:00.975439Z","submitted_at":"2024-09-26T17:58:55Z","title":"Lotus: Diffusion-based Visual Foundation Model for High-quality Dense Prediction","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18124","snapshot_observed_at":"2026-08-05T23:59:31.170775Z","title":"Lotus: Diffusion-based visual foundation model for high-quality dense prediction","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:31.170775Z"},"links":{"cited_paper":"/paper/2409.18124","citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:1945c77adc9812b877ddccbf8a3bd96391e18b7ff59b7ad2afe656ef48cda870","observation_id":"14e7924d-0ebf-4ea3-9a6b-20e8df6ad62f","resolution":{"observed_at":"2026-08-05T23:59:31.170775Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:31.247652Z","title":"Deep residual learning for image recognition","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:31.247652Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:6857a34b53f3e5449b5ce9000f4b2dc247a66c11fbd772d4398103391fef16e2","observation_id":"b9f2c839-7c7a-4764-9989-fbb819887f4a","resolution":{"observed_at":"2026-08-05T23:59:31.247652Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1606.08415","last_updated":"2023-06-06T01:53:32Z","snapshot_observed_at":"2026-08-13T19:48:28.322536Z","submitted_at":"2016-06-27T19:20:40Z","title":"Gaussian Error Linear Units (GELUs)","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1606.08415","snapshot_observed_at":"2026-08-05T23:59:31.335763Z","title":"Gaussian error linear units (gelus)","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:31.335763Z"},"links":{"cited_paper":"/paper/1606.08415","citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:2b00a4bdbad8d1ed722a38c23eb016a79c461bff930fffda45661a7abdae9897","observation_id":"d847719a-8e73-48e3-9906-960a82b9f3df","resolution":{"observed_at":"2026-08-05T23:59:31.335763Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:43.018191Z","title":"Stereo processing by semiglobal match- ing and mutual information","venue":null,"work_id":"1e163f1c-9ee1-4796-a569-d6e56f52bdea","year":2007},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:31.438021Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:67a1616181f8943ebd9b79a4c69970e15e66b71620ccf994debc406dd31601aa","observation_id":"9a08ecb6-c80d-4692-87b9-ca0783dd70ad","resolution":{"observed_at":"2026-08-05T23:59:43.116017Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:42.765413Z","title":"Learning accurate 3d shape based on stereo polarimetric imaging","venue":null,"work_id":"a5b7eab2-bdea-4dc6-b8af-67f27cc6e1fa","year":2023},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:31.554308Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:800648299aae26a4d96b5a13abcf84e2ee24d951c710c03fb30fb7870d6b65cc","observation_id":"be8fb9a6-43c6-450e-8b7e-41649d5a4883","resolution":{"observed_at":"2026-08-05T23:59:42.910805Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09466","last_updated":"2025-04-23T10:38:09Z","snapshot_observed_at":"2026-08-16T15:35:34.338608Z","submitted_at":"2025-01-16T10:59:29Z","title":"DEFOM-Stereo: Depth Foundation Model Based Stereo Matching","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09466","snapshot_observed_at":"2026-08-05T23:59:31.664672Z","title":"Defom-stereo: Depth foundation model based stereo matching","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:31.664672Z"},"links":{"cited_paper":"/paper/2501.09466","citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:01c18f4a0e19b3bf3607e07ddae6519df43e8af725c1c8cb1ca54e81e2418add","observation_id":"87821c19-df6c-4eb0-a60d-1b4fa0d12ace","resolution":{"observed_at":"2026-08-05T23:59:31.664672Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:42.529446Z","title":"Uncertainty guided adaptive warping for robust and efficient stereo matching","venue":null,"work_id":"0f940f4d-faf4-4e50-9ad5-e260c0e14b2d","year":2023},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:31.757981Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:acf452518991b2b4f997cb8f53812e675d205b434329ed184faa0d2f04ca614e","observation_id":"48fee5dc-32c6-47d0-8145-1e9cc7068d4b","resolution":{"observed_at":"2026-08-05T23:59:42.662459Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:42.243212Z","title":"Repurpos- ing diffusion-based image generators for monocular depth estimation","venue":null,"work_id":"624eb83c-caf5-43a0-90fc-b414cbff6325","year":2024},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:31.847725Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:0eba6069712335f65da23fc72990ff7cc42e3ee4bbe6ee67087e9c48309acc05","observation_id":"c4d2d4ac-3225-4324-b914-c7166b88e43b","resolution":{"observed_at":"2026-08-05T23:59:42.374874Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:42.092086Z","title":"End-to-end learning of geometry and context for deep stereo regression","venue":null,"work_id":"a5c51320-aa10-4c9a-8cf5-3c8396d2a750","year":2017},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:31.904884Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:495775859f575c58e9952aaf36bc68949d07adcd6d987b9d563417af40151ddb","observation_id":"a7929370-8e90-4488-b2d3-68cf2bf6099d","resolution":{"observed_at":"2026-08-05T23:59:42.145617Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1907.10326","last_updated":"2021-09-23T10:23:51Z","snapshot_observed_at":"2026-08-14T16:02:58.912669Z","submitted_at":"2019-07-24T09:31:24Z","title":"From Big to Small: Multi-Scale Local Planar Guidance for Monocular Depth Estimation","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1907.10326","snapshot_observed_at":"2026-08-05T23:59:31.963659Z","title":"From big to small: Multi-scale local planar guidance for monocular depth estimation","venue":null,"work_id":null,"year":1907},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:31.963659Z"},"links":{"cited_paper":"/paper/1907.10326","citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:77811596b8206ddfd8db1efec9f4c742df430f1d200f4b160077c15c5982fe37","observation_id":"1e13453e-bc70-4481-8169-712e0e549340","resolution":{"observed_at":"2026-08-05T23:59:31.963659Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:41.958784Z","title":"Practical stereo matching via cascaded recurrent net- work with adaptive correlation","venue":null,"work_id":"7f06ea4d-c2b8-4fd7-b978-d414f559ea40","year":2022},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:32.038384Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:e0bc83b9798f0232660ac48b5fed13a4e021bc8f6bc35d7f10be77e7167275a9","observation_id":"1a2652fa-d54d-4228-9769-9f29ee62c101","resolution":{"observed_at":"2026-08-05T23:59:42.017670Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:41.837101Z","title":"Los: Local structure-guided stereo matching","venue":null,"work_id":"9b8677d0-c56d-487c-b4db-56ea2d888f0b","year":2024},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:32.151109Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:eaef6f37c088a5e571e43e0fb23bf152acb82933fc7a1f5d415e8284b362c793","observation_id":"1be67994-4e38-4084-a9fd-00742500fa61","resolution":{"observed_at":"2026-08-05T23:59:41.880211Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:41.651581Z","title":"Revisiting stereo depth estimation from a sequence- to-sequence perspective with transformers","venue":null,"work_id":"bf93d146-3380-4cc3-8fa1-e0e0940981e9","year":2021},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:32.241042Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:561279b67a6ee258fc571af0f73830face28ddd264daffe4c3bc623f529467ca","observation_id":"7d7f9b69-d3be-46e2-8d5c-dc3ef8fcba5e","resolution":{"observed_at":"2026-08-05T23:59:41.751897Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:41.520357Z","title":"Raft-stereo: Multilevel recurrent field transforms for stereo matching","venue":null,"work_id":"91e54869-4d5b-45d7-8439-e233ce1619c6","year":2021},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:32.330705Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:ddfc94685228fd1bedaf7fd12ed09f9d4bfa9ce51ff80c82c0beafd613a26b35","observation_id":"beb6a81b-e729-427f-97a9-0854ebba6755","resolution":{"observed_at":"2026-08-05T23:59:41.580870Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:32.441710Z","title":"Swin transformer: Hierarchical vision transformer using shifted windows","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:32.441710Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:6f728cf82caa22ad673416f539fb2e8fd744de69de3de410638e22a26967add8","observation_id":"33654a50-de64-42d5-b594-c68c447b67c6","resolution":{"observed_at":"2026-08-05T23:59:32.441710Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:32.538032Z","title":"A convnet for the 2020s","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:32.538032Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:4f6b4f3544ad682cabdebc3f423b2eef8972598ee877115c613c9297bb71be49","observation_id":"85ae3263-d2e0-4c80-b895-0312c7c8e155","resolution":{"observed_at":"2026-08-05T23:59:32.538032Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:41.346896Z","title":"A large dataset to train convolutional networks for disparity, optical flow, and scene flow estimation","venue":null,"work_id":"5d170efc-b6ad-41b9-90ce-8d2ae90a3bd2","year":2016},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:32.626222Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:b0ecb03dd42f4f6ea78af65492e44024679023aa2a5e1460a037db8f498be840","observation_id":"2661ae07-c13c-41f7-b38b-e5bd6ad94826","resolution":{"observed_at":"2026-08-05T23:59:41.420093Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:41.190499Z","title":"Object scene flow for au- tonomous vehicles","venue":null,"work_id":"b78984ed-fcc4-4409-82a6-22b1705e8705","year":2015},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:32.687074Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:fa740a5ce92165b4e24e94e162d7013a8e529fbcdf3f190fe39865457cef57be","observation_id":"ac053ff7-425f-427b-9726-2989998ec923","resolution":{"observed_at":"2026-08-05T23:59:41.262407Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.07193","last_updated":"2024-02-02T10:24:09Z","snapshot_observed_at":"2026-08-17T13:03:40.359628Z","submitted_at":"2023-04-14T15:12:19Z","title":"DINOv2: Learning Robust Visual Features without Supervision","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.07193","snapshot_observed_at":"2026-08-05T23:59:32.769828Z","title":"Dinov2: Learning robust visual features without supervision","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:32.769828Z"},"links":{"cited_paper":"/paper/2304.07193","citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:48150ec2d782ed4d48b93a9840322ce01ff76d35d65679aa370b0b6c942b2672","observation_id":"7659e319-eee5-4393-8437-af27f64d2312","resolution":{"observed_at":"2026-08-05T23:59:32.769828Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:41.055031Z","title":"Unidepth: Universal monocular metric depth estimation","venue":null,"work_id":"0678c6c1-a2db-4027-af6d-3dc6f41b9b50","year":2024},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:32.864948Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:19ef3cf8eb9a9227464c80b8c992d2446b494769cf4702587801ee4135be894e","observation_id":"788f7f99-12c0-4cb4-87e7-fd8e50fcec47","resolution":{"observed_at":"2026-08-05T23:59:41.097059Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:40.918644Z","title":"UniDepthV2: Universal monocular metric depth estimation made simpler, 2025","venue":null,"work_id":"6bcc4f2e-b5d0-4b0d-9425-44d5423f31f6","year":2025},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:32.968991Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:8437f5edc3cf993b22e4bc10c24fd369c9b0da7accb2c6d5ec551ebc6de70670","observation_id":"c5672e86-eb29-4bda-8a32-4c6be9d74ebe","resolution":{"observed_at":"2026-08-05T23:59:40.980269Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:40.812021Z","title":"Towards robust monocular depth estimation: Mixing datasets for zero-shot cross-dataset transfer","venue":null,"work_id":"7bf74160-cc61-4201-9e94-ffa8c142c1e2","year":2020},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:33.074533Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:2bd5319c3cf8f55014def1d8770fea601f87d61065b048914639cd621998906e","observation_id":"4976ac76-e972-45cd-a207-0b38824edcc7","resolution":{"observed_at":"2026-08-05T23:59:40.859205Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:40.590045Z","title":"Vi- sion transformers for dense prediction","venue":null,"work_id":"b13792ae-5a9b-4c5e-8931-bf81945785c7","year":2021},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:33.139978Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:1773d6f165ad411dfd60d77b945639e047164f4514f98f423ff3bcd94eba9721","observation_id":"b88c5481-307c-4375-9b3f-c7feefd2fd39","resolution":{"observed_at":"2026-08-05T23:59:40.681616Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:40.407203Z","title":"Masked representation learn- ing for domain generalized stereo matching","venue":null,"work_id":"1a9d441d-6a45-47a3-bdd4-b5887b15f4b3","year":2023},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:33.197235Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:b822a8ddbc4308613f2424e410aebab569bc63bc74b044ecac7daa0d92e1ba82","observation_id":"61ae375c-ab95-4d38-bb1b-23dbdacba553","resolution":{"observed_at":"2026-08-05T23:59:40.495497Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:33.261656Z","title":"High-resolution image synthesis with latent diffusion models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:33.261656Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:8f93329f70a7d380374ce2e978b398434d8cca21520c0fc61f3d3c8e5eb596ba","observation_id":"8a9eac1a-cccc-4bb8-9d2b-010dd87bc4be","resolution":{"observed_at":"2026-08-05T23:59:33.261656Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:40.099947Z","title":"High-resolution stereo datasets with subpixel-accurate ground truth","venue":null,"work_id":"c0dba0a0-3a40-4dc8-8636-0ddc0d9133b9","year":2014},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:33.265977Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:88bb13bce3667ba11251a47002c208d296f6f0e6c6eb792dfdd6b8063930b003","observation_id":"1be7d3d0-bc85-4163-b9e4-869979e61e1f","resolution":{"observed_at":"2026-08-05T23:59:40.246801Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:39.848648Z","title":"A multi-view stereo benchmark with high- resolution images and multi-camera videos","venue":null,"work_id":"c536f135-ecb1-4feb-9458-50be484ded47","year":2017},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:33.280957Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:e30bcf61bd340ed9410f5b28bd9a10ebcbf88b92fefffd8081a1b4ea114d637b","observation_id":"ebeb9d55-8589-4138-9565-42dbf86f63e5","resolution":{"observed_at":"2026-08-05T23:59:39.972725Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:39.671871Z","title":"Cfnet: Cascade and fused cost volume for robust stereo matching","venue":null,"work_id":"fad56098-e8d7-4b15-8056-3dc148cb8b55","year":2021},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:33.361668Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:b511faaa3bc5068e165fce7b6b6a07f5edab451e1b3ea470c30b811dd34bc4c6","observation_id":"3195866a-af1e-4979-9e9c-5013dd72e271","resolution":{"observed_at":"2026-08-05T23:59:39.732326Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:39.376129Z","title":"Pcw-net: Pyramid combination and warping cost volume for stereo matching","venue":null,"work_id":"db712908-c91a-4bf0-b572-073ed12db4d9","year":null},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:33.498978Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:b147873fd9ff54c0e6cf65732bd01018e815965629d26c46e8d07cf1c342f671","observation_id":"302da1de-037d-4bb1-a95a-eb283ee9fd1f","resolution":{"observed_at":"2026-08-05T23:59:39.502213Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:39.191442Z","title":"Raft: Recurrent all-pairs field transforms for optical flow","venue":null,"work_id":"4376b351-e6c5-4cfc-b7e9-34b85647d5ad","year":2020},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:33.613409Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:0a83d22ec4269a9593fccd99b01c42b27b1355a261b98274cd0fbb0b815a8522","observation_id":"a04f572b-73a7-4b74-9321-6c885e9793ec","resolution":{"observed_at":"2026-08-05T23:59:39.289950Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:39.003196Z","title":"Moge: Unlocking accurate monocular geometry estimation for open-domain images with optimal training supervision","venue":null,"work_id":"bc717654-5b20-4b88-a349-8cf172522c53","year":2025},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:33.721314Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:80504573cddfca0f0e49bd050eac153bcd94781e2e740d11b7475320f342cde2","observation_id":"5ee608de-0777-48c6-a4c2-abc9422eda44","resolution":{"observed_at":"2026-08-05T23:59:39.090750Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:38.794929Z","title":"Selective-stereo: Adaptive frequency information selection for stereo matching","venue":null,"work_id":"2dfa3dd7-16a9-4ae1-b022-e8e89e0ea55b","year":2024},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:33.876376Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:74b217ae9e562b53e7c9cb48b1bd7939fb8ab8a246db62700af885c2c2104d03","observation_id":"47c70ea2-17c5-43fd-8155-7af5d428d417","resolution":{"observed_at":"2026-08-05T23:59:38.878819Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:38.620016Z","title":"Croco v2: Improved cross-view completion pre- training for stereo matching and optical flow","venue":null,"work_id":"86bdf53c-c3f1-4e71-9013-496731031416","year":2023},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:33.946950Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:ab0916e6828b1b4257b6bdf586b5583f5b79a30cd725c49bfb01329b2fed4168","observation_id":"606e4c40-1cf3-4823-b456-0844f8b94a46","resolution":{"observed_at":"2026-08-05T23:59:38.720595Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09898","last_updated":"2025-04-04T00:51:17Z","snapshot_observed_at":"2026-08-14T13:53:43.298863Z","submitted_at":"2025-01-17T01:01:44Z","title":"FoundationStereo: Zero-Shot Stereo Matching","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09898","snapshot_observed_at":"2026-08-05T23:59:34.069620Z","title":"Foundationstereo: Zero- shot stereo matching","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:34.069620Z"},"links":{"cited_paper":"/paper/2501.09898","citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:7f0dcf6bc43b3e7f7263e48fcbb5c2cac07bfccef2655bd49202d9913b889bbf","observation_id":"4f3c0e39-fadb-4e61-9e58-00272a6f7427","resolution":{"observed_at":"2026-08-05T23:59:34.069620Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:38.437819Z","title":"Atten- tion concatenation volume for accurate and efficient stereo matching","venue":null,"work_id":"a950c1ad-0bff-499b-8082-627cd8c08bb6","year":2022},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:34.236171Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:fab310b82e4405f5f7d122647f99c07c1ed5b0bc6c09e2ed7a7e7941c5f586e6","observation_id":"21979546-604c-48f2-8280-3c300edf9ca8","resolution":{"observed_at":"2026-08-05T23:59:38.510780Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:38.240946Z","title":"Iterative geometry encoding volume for stereo matching","venue":null,"work_id":"37f15082-daad-41bf-951d-6110657fa197","year":2023},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:34.367791Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:ad73e0677416c47ab1cb25727dfbf69375b9dcda3b9450755ead6b3a2d279a65","observation_id":"dd3f180d-a6ad-422a-b141-99bb3df65aae","resolution":{"observed_at":"2026-08-05T23:59:38.317701Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.00638","last_updated":"2025-05-11T09:54:21Z","snapshot_observed_at":"2026-08-16T13:22:21.789198Z","submitted_at":"2024-09-01T07:02:36Z","title":"IGEV++: Iterative Multi-range Geometry Encoding Volumes for Stereo Matching","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.00638","snapshot_observed_at":"2026-08-05T23:59:34.495081Z","title":"Igev++: iterative multi-range geometry encoding volumes for stereo matching","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:34.495081Z"},"links":{"cited_paper":"/paper/2409.00638","citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:a52d94cefa87db76239a7b56a4bb06862272143409c361c517b12df7da7e548e","observation_id":"c11b9ada-2d6a-4340-8f06-f6068ee60ced","resolution":{"observed_at":"2026-08-05T23:59:34.495081Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:38.035521Z","title":"Depth anything: Unleashing the power of large-scale unlabeled data","venue":null,"work_id":"32ad514b-9c6a-4ae7-b4a2-fd6d6623c0a7","year":2024},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:34.618845Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:b96166004134d336ae5e4e24ff6549fc06047f1e0e63b16848f18475ece05512","observation_id":"c59b1de6-296d-4a9d-bcb4-e068d46809bd","resolution":{"observed_at":"2026-08-05T23:59:38.153404Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:37.842556Z","title":"Depth any- thing v2","venue":null,"work_id":"b939d33e-37d5-4f3c-9945-cf3c5339bf5a","year":2024},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:34.779838Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:f9c86a46ada95fbea624287b776d61f61ffd672ee19e133794e5619790dec7cb","observation_id":"d3b4dda8-53a0-4961-b39e-6d786b0f18fd","resolution":{"observed_at":"2026-08-05T23:59:37.947003Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:37.665252Z","title":"Metric3d: Towards zero-shot metric 3d prediction from a single image","venue":null,"work_id":"c2d7aa68-6179-4af6-9ebb-aa46653ac832","year":2023},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:34.964117Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:729a578c0aa0c8a9e016db7ecfac9daf5ab39fc3f4a31581e59654c6a0777d46","observation_id":"1902cbe3-03e2-4f25-a6d5-6d4499d5716b","resolution":{"observed_at":"2026-08-05T23:59:37.753961Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:37.455929Z","title":"Open chal- lenges in deep stereo: the booster dataset","venue":null,"work_id":"c3331cce-069a-44ab-8733-0ac43eaa5608","year":2022},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:35.100209Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:dac01df71f7b617676acc4cf60169b56ac23156e3689701eec0838268d58313d","observation_id":"6a2a5934-e8cd-45f5-88c7-214151e1600a","resolution":{"observed_at":"2026-08-05T23:59:37.578908Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:37.279605Z","title":"Stereo matching by training a convolutional neural network to compare image patches","venue":null,"work_id":"e0a46813-7e32-4c91-8441-3a3d2035a5b6","year":2016},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:35.244021Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:066db8198539ccf57c7205b6fb2107da12f16872659683aa7c19986f851313a5","observation_id":"dc5ad3ed-0f61-454a-bf3c-b0f7def922b5","resolution":{"observed_at":"2026-08-05T23:59:37.345190Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:37.046573Z","title":"Ga-net: Guided aggregation net for end-to- end stereo matching","venue":null,"work_id":"d64203f8-2d09-440d-b530-e0f4e1af067d","year":2019},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:35.348206Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:c3b5888917be91ec41f530ffeaddc187f5d7736516eb5c7bf7fc2048a3397afb","observation_id":"a939be36-0352-4241-81bc-cccdc4a78957","resolution":{"observed_at":"2026-08-05T23:59:37.163304Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:36.865426Z","title":"Domain-invariant stereo matching networks","venue":null,"work_id":"db3eae48-e0d2-4347-8af6-8e0159f9690a","year":2020},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:35.493849Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:3e26a8130b997ce5524025c428e361a2d7258e3b64878c13c0d1f0f7e87f3f64","observation_id":"1de722fe-a9ed-4c22-a668-00b6e6b459b4","resolution":{"observed_at":"2026-08-05T23:59:36.967954Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:36.658857Z","title":"Learning representations from foundation models for domain generalized stereo matching","venue":null,"work_id":"377fd508-d253-4db2-ab64-42d95f02f94e","year":null},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:35.668023Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:91e4bfe5353ac97b5ca017821ed2bd5a78255fb96566acaed7326f509bf9bb95","observation_id":"8deb7726-5123-47c8-97eb-f3eb203c9836","resolution":{"observed_at":"2026-08-05T23:59:36.763834Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:36.488218Z","title":"High-frequency stereo match- ing network","venue":null,"work_id":"8f87f9cb-924a-4cae-b683-9a5738235926","year":2023},"citing_paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:35.792565Z"},"links":{"citing_paper":"/paper/2508.04611"},"observation_digest":"sha256:6ac01b37543a66cb78be9122b7a146d70d2ec0f50d32f2e59233b1de0568e5ef","observation_id":"30147161-57e7-47e8-bc0a-2642e50494b7","resolution":{"observed_at":"2026-08-05T23:59:36.555846Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2508.04611","last_updated":"2025-08-13T12:52:56Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-13T15:37:05.031836Z","submitted_at":"2025-08-06T16:31:22Z","title":"BridgeDepth: Bridging Monocular and Stereo Reasoning with Latent Alignment"},"reference_resolution":{"displayed":61,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":17,"verified_exact":1,"verified_fuzzy":43},"total_outbound_references":61},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"thesis":"As of 17 August 2026, this Paper Citation Record lists 61 of 61 outbound references and 2 inbound Pith citation observations for arXiv:2508.04611."}