{"as_of":"2026-08-06T21:31:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:7c3ecfa3e521aed02fb409419df98d911f8a7d6627ef6ff079bb57bcab7056e1","coverage":[{"denominator":42,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":42,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-29T12:45:53.128640Z","state":"measured"},{"denominator":42,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":42,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-06T06:34:29.942622+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2605.28477/citation-record","integrity":"/paper/2605.28477/integrity","json":"/paper/2605.28477/citation-record.json","paper":"/paper/2605.28477"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Holistic 3d scene understanding from a single image with implicit representation,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:2a7b2664cfe83dc50f4896b9a03b48591211ee1b53a8ae8fd823443f7b04f8a0","observation_id":"a53ef8f5-496f-41d4-803a-09a7563f4ae9","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"DepthSplat: Connecting gaussian splatting and depth,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:777d546e5a33b320a808cd2722648f1e5c24cc4d2c1b6f5beecd4e12a92b707d","observation_id":"38ee9065-01bb-4b6c-bc5c-387e12b3ec8b","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Visual attention- based self-supervised absolute depth estimation using geometric priors in autonomous driving,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:496d1c5f88215b6c1c16dd6217e91dde2713e8f5dd3c75caec21440a13b04fa2","observation_id":"3000e62a-4909-487b-81c0-b26f5ec6484c","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Orb-slam2: An open-source slam system for monocular, stereo, and rgb-d cameras,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:f54930320e525648f0955dc64fd6d41a8b3bbed1c23bf9c10d6ee85a8efc18c4","observation_id":"1fc100fd-476e-4fe0-b572-a382e2d454c6","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Adabins: Depth estimation using adaptive bins,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:f4e78ecdf389e03e38f570fce731cb9c960e674207eebd08e66decdb44bac939","observation_id":"32eec0cc-6b34-45cd-8824-08d65dd57e60","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Depth anything: Unleashing the power of large-scale unlabeled data,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:0b8fc9011a2a9ba3f2afd8c855334b26de040081ef2aae94ddaa3bf0d1ceb5ac","observation_id":"93cfb36d-5717-4d0c-9d3a-383b310255ff","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"3D packing for self-supervised monocular depth estimation,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:e9172e3f366f42597cf63b30d154fa08c5aa8aef48a848d676546cc39b2080ae","observation_id":"2d0fa5c1-7a98-4ef6-9ea4-463e9a77c4bf","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Digging into self-supervised monocular depth estimation,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:215fc758e0b4c9b5561e53965cc807abccae389c8ec35fdcf7a9819be034126f","observation_id":"48cd0afc-fcc4-4983-9a13-1b68370e92ec","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Unsupervised learning of depth and ego-motion from video,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:ed115080f02cb9418ede648d84fc7ac5fd8e5d64cd6363e7a41b852f955a5414","observation_id":"86e18f97-0adc-497e-9c65-d96a6f6c2c59","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Vision meets robotics: The KITTI dataset,","venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:9c6d76f0e9569731d12d2ddcc17485ae90282ccff2f4aff17de19c7d07800c66","observation_id":"8c922516-9941-4b8c-b946-9913ad33a964","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"The cityscapes dataset for semantic urban scene understanding,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:8579c15d3ac4ef0794214e223102c8d846784e628dbf3d0efd449c767337f7b9","observation_id":"de7fa7ec-62c6-4eb9-a570-9601b1e58d73","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Indoor segmen- tation and support inference from rgbd images,","venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:72380264c286f568740abbd9026000034631c6edf3481277f8311dde53afca16","observation_id":"53606309-3f4e-4bc7-bd77-db1cb2349ad4","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"MonoViT: Self-supervised monocular depth estimation with a vision transformer,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:2493043ccc5cb6045edb12f31d43d6d7012ad93a15920ddc8253db359cdbacd6","observation_id":"bde57f93-795d-4b8d-b8e5-de7e3f364cc0","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Camera height doesn’t change: Unsu- pervised training for metric monocular road-scene depth estimation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:dcd72d28421aa2c17d4a70ae4d2ca891cd052a02ee9f5e7191e685e6b8cf1cd9","observation_id":"6cc140cc-f482-4396-a3d0-c216d00ef01d","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"R4dyn: Exploring radar for self-supervised monocular depth estimation of dynamic scenes,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:bbe477cba2f7d3706178c7389c4ebe98edb9b8a76e376ec67bcfe06507c3b71f","observation_id":"09c3e69e-0eee-4041-b6c3-8c75413ff3b6","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Robust monocular depth estimation under challenging conditions,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:3344e1979c2c464bdf525c15dc6166d7240f34fb05b068de62ba4a6c730c8ca1","observation_id":"b150d7c8-491f-44c6-965b-da36b7f91210","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Learning depth from monocular videos using direct methods,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:088b389f69ef0410bc9209e94934806112ab15b3fe905c87096c730a7c578013","observation_id":"464c387f-f417-4ce5-a1e8-4b229758fef5","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Towards better generalization: Joint depth-pose learning without posenet,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:c8d71a14175cfe4a78c6de337172dd68e5a90266254d4bd4c174a8790ebbd350","observation_id":"7b70eb97-5930-4651-929f-6258d6c61db0","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"DualRefine: Self- supervised depth and pose estimation through iterative epipolar sampling and refinement toward equilibrium,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:93004c20747334eb547fc4e8ab68b4c74a278dd84ca30e5b6eed0577f5741488","observation_id":"1175d42a-a338-46d6-8dbe-2078e35d59df","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Are we ready for autonomous driving? the KITTI vision benchmark suite,","venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:6bd3d197397b58649e3dadbb3e0c3c653b3174238898447767e509764e2a2aac","observation_id":"4f13fa43-1be2-4ff9-9677-50cb10a82fea","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Unsupervised learning of monocular depth estimation and visual odom- etry with deep feature reconstruction,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:401a8f5efb5c557149575c17dabf03936a70414efead4a23d726daec6ba2c4f6","observation_id":"9254728d-a928-4b25-b951-ccfa2986b196","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Mono-vifi: A unified learning framework for self-supervised single and multi-frame monocular depth estimation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:b14d9984f0865cf3f1d564eb11e7a35b447a619d98b580bd0497867dcf6ee2ef","observation_id":"8d9ee535-ed99-4408-be7f-0cf2831a0846","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Channel-wise attention-based net- work for self-supervised monocular depth estimation,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:44e217c79e653378783bac45b27b9169982be0454dd845c02c2277aa7d39e721","observation_id":"3c8d0679-5d2a-494e-a953-438fa166838b","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Monoindoor++: Towards better practice of self-supervised monocular depth estimation for indoor envi- ronments,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:00342520aa1c064e4172abbc7b2fcf2ba64ec82716cc16239f217d9405a7accc","observation_id":"27933bb4-a7db-4711-b9f3-e9cb68cca474","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Superglue: Learning feature matching with graph neural networks,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:e23f97d12db9a717ce903de9dee23dd4971423c8d2cccffa2daeb0fa77f81179","observation_id":"fbe63889-0507-449d-972c-7bfb72066e63","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"LM-Reloc: Levenberg-Marquardt based direct visual relocalization,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:934fc906e37b0526749a37dd6e9ca5aa57fee7d47dfac71783b113158a6029b7","observation_id":"a500177a-52d2-46b2-9faa-54aac4b55987","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Back to the feature: Learning robust camera localization from pixels to pose,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:efb90c0fa41c615f7a474f9d1343eba41a7d28bfdcea4d83a871d1c37fd7a514","observation_id":"7cd75cc2-ae21-4434-9ab1-32594b7a3c50","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Relative pose estimation through affine corrections of monocular depth priors,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:8a12022d683ad97e7affcc84f2dada14baa94a8d47d0212ecb267f9bc32eeadc","observation_id":"db3240e3-a656-4ad3-b811-b92512be4773","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2001.10773","last_updated":"2020-01-29T12:13:20Z","snapshot_observed_at":"2026-07-06T08:53:28.193420Z","submitted_at":"2020-01-29T12:13:20Z","title":"Virtual KITTI 2","version":1},"cited_work":{"arxiv_id":"2001.10773","doi":null,"metadata_source":"pith","pith_arxiv_id":"2001.10773","snapshot_observed_at":"2026-07-07T22:24:11.153404Z","title":"Virtual KITTI 2","venue":"cs.CV","work_id":"c0d9c030-aa25-44e7-9cc4-72d7403f1447","year":2020},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"cited_paper":"/paper/2001.10773","citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:c1408449f2bd826bc2291812f3f648eff35952cfe93cc341b4319a91978d7af7","observation_id":"baa7842c-dce3-40ec-909d-0cf04a9d0b32","resolution":{"observed_at":"2026-06-29T12:53:26.920321Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Cnn-slam: Real-time dense monocular slam with learned depth prediction,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:c685ad54688cdfe0df903c7a64f30df4b614c12781a6a99e3ed4fe8c6da4e4ac","observation_id":"760079fa-338a-44f8-a2fe-c02b02439e44","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Sparsity invariant cnns,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:276675c5c0ac2e30a3e7bfab2bf75fb35663726a95b7c40ac4281d3cee9b4b20","observation_id":"30c11ed7-6542-41c6-95a8-8dd855ce4212","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Deep residual learning for image recognition,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:82e26cd5a2b7897dd82d83c6a2e626d72ab3c9b5fbab2834aaa07518cac7793b","observation_id":"127b069f-b7c0-4c43-a987-e93674ef67cb","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Hybrid-grained feature aggregation with coarse-to- fine language guidance for self-supervised monocular depth estimation,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:c8675611b01a74570782ee4b8a4d1c7f036be8198de29927c802f77c04593760","observation_id":"1dc2bb8a-5c38-49ef-a5a4-6f3688afbb77","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1412.6980","last_updated":"2017-01-30T01:27:54Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2014-12-22T13:54:29Z","title":"Adam: A Method for Stochastic Optimization","version":9},"cited_work":{"arxiv_id":"1412.6980","doi":"10.1002/mrm.28086","metadata_source":"pith","pith_arxiv_id":"1412.6980","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Adam: A Method for Stochastic Optimization","venue":"cs.LG","work_id":"1910796d-9b52-4683-bf5c-de9632c1028b","year":2014},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"cited_paper":"/paper/1412.6980","citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:f883357c7e012fb4390f1049ee4dd9ee2e109924901a04f56d400b026e6f5691","observation_id":"df7e5cce-a00f-4b72-ae71-f03ea77a525c","resolution":{"observed_at":"2026-06-29T12:53:26.917844Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Predicting depth, surface normals and semantic labels with a common multi-scale convolutional architecture,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:1bb070e33891c379fd1d521d6e5d615334135451c4cb450f6f110e44c8d757d2","observation_id":"cb849848-7f56-4ff9-a3ba-fa2c496b7c50","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"The temporal opportunist: Self-supervised multi-frame monocular depth,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:a3141b2fff5ae17ada7955da74a7e597a089dca971349f10c3c407cde47a5003","observation_id":"9c61daba-04ec-4e3a-9cca-5767f23ed5ed","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Auto- rectify network for unsupervised indoor depth estimation,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:4b3f7383193b98d62f700aa1cfcaf03bde1124869a8186afacffef5682c3f8f6","observation_id":"f4bdc1ea-2390-4bf6-93ff-f5232b25adb1","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Visual odometry revisited: What should be learnt?","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:d324a3755a6be5b7eb5ab3fe9eb114ad3185d161c83eb8c37c82e057c69e3169","observation_id":"daf21c3d-e774-4cf4-b1a6-e29fdb19f21a","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Unsupervised scale-consistent depth and ego-motion learning from monocular video,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:011cb0a3498fd3489d975509b62ddbe674dee620d9a71fc9a95981d7b4fb6d96","observation_id":"e0ca3a09-c5ee-4880-b3d6-01d998006ab4","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Grounding image matching in 3d with mast3r,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:ad5131a616b8bc28a6cc979060c42083217bee504301422f4e6cbd3323e38a94","observation_id":"938ff93b-2d1e-4243-96e0-be86407553e6","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Depthcrafter: Generating consistent long depth sequences for open- world videos,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:0cc92e3ee4599203cba8a034de7c497dcae40700dc832380f51b224916dddfbe","observation_id":"89675dd9-8470-490e-ae93-c4df1ba7d6ac","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T12:45:53.128640Z","title":"Video depth anything: Consistent depth estimation for super-long videos,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-06-29T12:45:53.128640Z"},"links":{"citing_paper":"/paper/2605.28477"},"observation_digest":"sha256:131f483edcbb60d6ece6750e1e6af1f0182228b12f2578829fe5dc48c778bd76","observation_id":"6ab5ef5a-4787-4ada-80c1-c946452d4fb8","resolution":{"observed_at":"2026-06-29T12:45:53.128640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2605.28477","last_updated":"2026-05-27T13:38:25Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T23:38:04.621658Z","submitted_at":"2026-05-27T13:38:25Z","title":"SA4Depth: Consistent Pose-Depth Scale Alignment for Self-Supervised Monocular Depth Estimation"},"reference_resolution":{"displayed":42,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":40,"verified_exact":2,"verified_fuzzy":0},"total_outbound_references":42},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"thesis":"As of 6 August 2026, this Paper Citation Record lists 42 of 42 outbound references and 0 inbound Pith citation observations for arXiv:2605.28477."}