{"as_of":"2026-08-18T16:04:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:0ffbe1b8998df70f6fca6236097119676c80d9ee6f1fec75b55eb58beda13ab2","coverage":[{"denominator":67,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":67,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T04:42:45.859616Z","state":"measured"},{"denominator":68,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":68,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-18T06:34:40.430872+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-08T17:29:15.428146Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-11T17:31:07.484648Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"cited_work":{"arxiv_id":"2506.09919","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.09919","snapshot_observed_at":"2026-07-01T02:17:15.958537Z","title":"Metrichmr: Metric human mesh recovery from monocular images","venue":null,"work_id":"365cd723-6396-45db-b974-9ba952d74a11","year":2025},"citing_paper":{"arxiv_id":"2605.04728","last_updated":"2026-05-06T10:23:10Z","snapshot_observed_at":"2026-08-13T18:26:49.103097Z","submitted_at":"2026-05-06T10:23:10Z","title":"Anny-Fit: All-Age Human Mesh Recovery","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-05-08T17:29:15.428146Z"},"links":{"cited_paper":"/paper/2506.09919","citing_paper":"/paper/2605.04728"},"observation_digest":"sha256:82fe42cfc0b933a8e9467050297113d745c4505f30a31197fdf07c8614edccd3","observation_id":"6e2adacb-b2a4-445a-8d63-6a9bde456d57","resolution":{"observed_at":"2026-07-01T02:17:15.958537Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.09919/citation-record","integrity":"/paper/2506.09919/integrity","json":"/paper/2506.09919/citation-record.json","paper":"/paper/2506.09919"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.439397Z","title":"2d human pose estimation: New benchmark and state of the art analysis","venue":null,"work_id":"6804dcae-60ba-4df8-8447-fbe010b2c59a","year":2014},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.703699Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:ce0b460d38c191aa40442d08dda686521d51779fe95083dff81e3bdd69229de5","observation_id":"be7cb1ae-192d-441a-a490-3121c586de5d","resolution":{"observed_at":"2026-08-07T04:42:46.441812Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.432153Z","title":"Multi-hmr: Multi-person whole-body hu- man mesh recovery in a single shot","venue":null,"work_id":"c6b6e73f-104a-4e33-af22-50c2a11ca545","year":2024},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.706771Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:8b142b031250ad810f674937dc6d1a8c4c6c9f60646e18aec742e1235cf27b01","observation_id":"f56d4078-56d4-4fcc-b659-9d0b65423fdc","resolution":{"observed_at":"2026-08-07T04:42:46.434656Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2302.12288","last_updated":"2023-02-23T19:13:10Z","snapshot_observed_at":"2026-07-06T14:55:15.719380Z","submitted_at":"2023-02-23T19:13:10Z","title":"ZoeDepth: Zero-shot Transfer by Combining Relative and Metric Depth","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.12288","snapshot_observed_at":"2026-08-07T04:42:45.709232Z","title":"Zoedepth: Zero-shot trans- fer by combining relative and metric depth.arXiv preprint arXiv:2302.12288, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.709232Z"},"links":{"cited_paper":"/paper/2302.12288","citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:5899aa5119c8e9a3c65a3840610561d0d6952252a0c78f7c68e9b94cc225c9c9","observation_id":"b9c1a2fa-5060-4980-b920-f166e31914a5","resolution":{"observed_at":"2026-08-07T04:42:45.709232Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:45.712205Z","title":"Bedlam: A synthetic dataset of bodies exhibit- ing detailed lifelike animated motion","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.712205Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:a10f27e46ba1bb9355b29ba0d6228ac22b78cfb1a33956def077d5e4d84a4ae3","observation_id":"2bac904b-58ab-416b-817a-5b2b6ae36de4","resolution":{"observed_at":"2026-08-07T04:42:45.712205Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.420471Z","title":"https://www.blender.org","venue":null,"work_id":"4a89de04-dc8b-4aae-baaf-012e7be8b5b4","year":2025},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.714553Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:decf00e1393a0834604167e31f102d7a66602c2bc9d011d9390847d5e7722326","observation_id":"f534173b-d50c-4cd1-baee-e092c0e0fa08","resolution":{"observed_at":"2026-08-07T04:42:46.422731Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.413113Z","title":null,"venue":null,"work_id":"808edafe-11b4-46bb-94c7-7f2afae8f03c","year":2016},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.716978Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:44e1130f5227b8c8440cea2fda4759d86d60c0e072832c769533a9925803bf7a","observation_id":"900b11ef-1d4a-4c2a-9966-28d9d7439a78","resolution":{"observed_at":"2026-08-07T04:42:46.415667Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.405651Z","title":"Smpler-x: Scaling up expressive human pose and shape estimation.Advances in Neural In- formation Processing Systems, 36, 2024","venue":null,"work_id":"a6da72dd-efc0-4a1d-99b2-c45143824b4f","year":2024},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.719467Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:03c0e9b2c6f1b698f7f4776f38afa6bedc92e89b4699a5175c3383e1e19dd734","observation_id":"24bf5637-5939-42d9-b525-f8c7b0d93a85","resolution":{"observed_at":"2026-08-07T04:42:46.408174Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:45.721601Z","title":"Human3r: Everyone every- where all at once.arXiv preprint arXiv:2510.06219, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.721601Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:60c6cee118af8bf80a60191ee3cf80a6ffc1a987ba746224770dd883ab2380eb","observation_id":"3ca2d415-7beb-4e5b-a727-650d62f147f3","resolution":{"observed_at":"2026-08-07T04:42:45.721601Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.398449Z","title":"Accurate 3d body shape regression using metric and semantic attributes","venue":null,"work_id":"4f5673e3-406b-49b6-bc97-501e6e724224","year":2022},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.723898Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:4653aed53618fe1ff60704174460fd5e7b41b458d96d4bf91f674bd598aee580","observation_id":"2f21318b-b549-42e9-a939-c1cd5d889805","resolution":{"observed_at":"2026-08-07T04:42:46.400884Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.390823Z","title":"Deepseekmoe: Towards ultimate ex- pert specialization in mixture-of-experts language models","venue":null,"work_id":"d926cae5-8f23-4f65-b2a6-327b1d4681f6","year":2024},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.726243Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:206e658c95fe55d8d020476e995f16cf2e5119b7d511ee1e1fdd557ed68fc444","observation_id":"c29ae74e-1e8a-4906-a731-3912bc98890a","resolution":{"observed_at":"2026-08-07T04:42:46.393439Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.383776Z","title":"An image is worth 16×16 words: Transformers for image recognition at scale","venue":null,"work_id":"769194da-2a74-4687-bc5a-84a1db39b0a4","year":2021},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.728546Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:97e9c1b714f28ec1f33f84d1f4b9b1a0c0ca4af338b378c1a7e8e59e386dac52","observation_id":"42923f30-74c5-428a-b5e3-fc6a6379c9be","resolution":{"observed_at":"2026-08-07T04:42:46.386230Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.375951Z","title":"Tokenhmr: Advancing human mesh recov- ery with a tokenized pose representation","venue":null,"work_id":"4bd8afa5-5035-43df-998d-46bef8704c69","year":2024},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.730942Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:7124331a6831beb90be8b47b5deb9e8bb4285129242313282d069e5c892d120d","observation_id":"5e2bf442-fd35-4268-ab0e-a8cfda717ac2","resolution":{"observed_at":"2026-08-07T04:42:46.378906Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:45.733245Z","title":"Switch transformers: Scaling to trillion parameter models with sim- ple and efficient sparsity.Journal of Machine Learning Re- search, 23(120):1–39, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.733245Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:632e0dd6d3f36959848b831d67e5909e2eea45fe4fbee817132001d2c2ab0472","observation_id":"989f7b5e-1606-4868-9250-1c0506d6528e","resolution":{"observed_at":"2026-08-07T04:42:45.733245Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:45.735477Z","title":"3d-front: 3d furnished rooms with layouts and semantics","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.735477Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:d6a93c3d9743c6767eb7bfc25821f7c0e72920baeb98cfc2ee22502c671f29fb","observation_id":"d3e7f77b-d70e-4f08-8a8c-c849b0971315","resolution":{"observed_at":"2026-08-07T04:42:45.735477Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.360259Z","title":"Humans in 4d: Re- constructing and tracking humans with transformers","venue":null,"work_id":"5de9ef93-51ea-4991-a6e0-0abaa5494637","year":2023},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.738029Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:1f7ceea25903b86fe0499d10a7ed465a6a46a32275a1919021563e740d348dd1","observation_id":"c199ef6c-71ad-4076-9714-97c117e19b36","resolution":{"observed_at":"2026-08-07T04:42:46.362701Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.15732","last_updated":"2024-11-23T06:06:14Z","snapshot_observed_at":"2026-08-16T21:44:56.744200Z","submitted_at":"2024-10-21T07:51:17Z","title":"ViMoE: An Empirical Study of Designing Vision Mixture-of-Experts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.15732","snapshot_observed_at":"2026-08-07T04:42:45.740335Z","title":"Vimoe: An empirical study of designing vision mixture-of-experts.arXiv preprint arXiv:2410.15732, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.740335Z"},"links":{"cited_paper":"/paper/2410.15732","citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:24126e82328176a364b670865f35b6dcc70a076ebf59d3a74a944863e7566a0b","observation_id":"695335bc-caa1-4ad5-8110-fda779536e24","resolution":{"observed_at":"2026-08-07T04:42:45.740335Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.353540Z","title":"Perspose: 3d human pose estima- tion with perspective encoding and perspective rotation","venue":null,"work_id":"8cb5fff4-9178-42a5-a3dc-55ae407d850c","year":2025},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.742715Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:d2ea119a70bd046a651be2407ebbbdb69da57e9fd619082558f3363648fbe2e2","observation_id":"a8832697-0185-43f0-a3b0-c438abaa250a","resolution":{"observed_at":"2026-08-07T04:42:46.355923Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.346729Z","title":"Resolving 3d human pose ambiguities with 3d scene constraints","venue":null,"work_id":"6b0e7ce5-c3f5-4800-91c8-9df5c7fa8b3a","year":2019},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.744998Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:7805e4a0e24d6c8bd678058a5532a0c732f37377f5e386405736e367d1657c74","observation_id":"6e7fbe59-54a2-41c8-aa4b-8da79b53e2b4","resolution":{"observed_at":"2026-08-07T04:42:46.349135Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.339743Z","title":"Pow3r: Empowering un- constrained 3d reconstruction with camera and scene priors","venue":null,"work_id":"70bd2858-40dc-4cfa-aaa6-79fc80e8fb1f","year":2025},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.747183Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:acf1239e4fe54d867b76fb0d1551a79ac54ae58822bba9179b803e0a1ae2e2de","observation_id":"8c1f932b-0fe9-4ede-b0a9-d61105c4b0e1","resolution":{"observed_at":"2026-08-07T04:42:46.342111Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.332756Z","title":"End-to-end recovery of human shape and pose","venue":null,"work_id":"b250fe78-44cc-4cd8-a51d-8d4e91bf4a00","year":2018},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.749587Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:af0f596948c2476ba351efad6fc247e3279127ed18234a2ffe8c82b74aa8db7e","observation_id":"24df35c5-d0eb-44ff-9397-cde70b2204d8","resolution":{"observed_at":"2026-08-07T04:42:46.335206Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.326034Z","title":"Emdb: The electromagnetic database of global 3d human pose and shape in the wild","venue":null,"work_id":"31011184-13bf-4f1a-9cad-109f958fff3a","year":2023},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.752155Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:67dd5e6011617c753067feedf9c6e219b2131e28aa49e5245116734cb9b825ff","observation_id":"a9c28b58-d9c9-4e0f-a6fb-32a288926f48","resolution":{"observed_at":"2026-08-07T04:42:46.328313Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.13414","last_updated":"2026-01-23T18:59:33Z","snapshot_observed_at":"2026-08-17T16:28:17.444550Z","submitted_at":"2025-09-16T18:00:14Z","title":"MapAnything: Universal Feed-Forward Metric 3D Reconstruction","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2509.13414","snapshot_observed_at":"2026-08-07T04:42:45.754510Z","title":"Mapanything: Universal feed-forward metric 3d re- construction.arXiv preprint arXiv:2509.13414, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.754510Z"},"links":{"cited_paper":"/paper/2509.13414","citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:74807b410bca77c38e80d9ce3adca9d30da1eec26e14decb47ec3bf28b29247e","observation_id":"a651b0f6-1005-49b3-b341-8a26b3ea95b9","resolution":{"observed_at":"2026-08-07T04:42:45.754510Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.318133Z","title":"Beyond weak perspective for monocular 3d human pose estimation","venue":null,"work_id":"7757396d-6986-4260-8afb-e94ca0fb7194","year":2020},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.757903Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:d4a858f3ba33b692437364538b566870405c197846aaca952c1c53079e659398","observation_id":"8c523c4e-1e08-411f-b21b-50e4fc577d9a","resolution":{"observed_at":"2026-08-07T04:42:46.320857Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.310757Z","title":null,"venue":null,"work_id":"1a16f41a-e380-435f-a4dc-0d739cb05cb5","year":2020},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.760556Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:caf470837a7a2346d582b56154b230e008f245e1a2dac1405ae4c2f1f1b06de4","observation_id":"07231fe7-6b50-4881-b294-2ced62c5ca24","resolution":{"observed_at":"2026-08-07T04:42:46.313075Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.303128Z","title":"Spec: Seeing people in the wild with an estimated camera","venue":null,"work_id":"3cda2daf-45c9-4cfc-bd81-bfdfa400557b","year":2021},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.762916Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:8c2d0f3538a681aaacd16d55731060ba46980e927ebe9fdef80c3a39ed7013ab","observation_id":"bf2093df-d4be-4828-a3ef-2efa5f154760","resolution":{"observed_at":"2026-08-07T04:42:46.305991Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.295715Z","title":"Learning to reconstruct 3d human pose and shape via model-fitting in the loop","venue":null,"work_id":"1429c597-62b9-4485-a80c-bc2088afa924","year":2019},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.765511Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:fe70a4420fe9957b2abdcd5ac7b0494ac6464c670f2e1551289f6ca987bceaa1","observation_id":"4884148e-a0d4-4867-b390-97a9586c7e91","resolution":{"observed_at":"2026-08-07T04:42:46.298266Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.288126Z","title":"Coin: Control-inpainting diffusion prior for human and camera motion estimation","venue":null,"work_id":"2c2f0ed6-c118-4630-a5f8-ebd98008645d","year":2024},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.767908Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:a17c4a3c7c41403ce1e33048b1ba228079fcf4561fe14aafca8796c885538744","observation_id":"0e7a1110-21ae-4e97-95bf-94be633afd7d","resolution":{"observed_at":"2026-08-07T04:42:46.290887Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.280477Z","title":"Cliff: Carrying location information in full frames into human pose and shape estimation","venue":null,"work_id":"6c993788-251d-4a2e-9bb2-5ef16a57c74f","year":2022},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.770524Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:df14f0505a2681baedc2b3efe50c03205e5988f689629c1f21fac79d2d927859","observation_id":"73e360f3-b866-4fca-a84f-b7c1b73447ae","resolution":{"observed_at":"2026-08-07T04:42:46.283016Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:45.772873Z","title":"Microsoft coco: Common objects in context","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.772873Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:cc418393ea26d13ec3088998effb0caeda83ddbbec293b10fa619b620bd7b3f5","observation_id":"78b76d16-e77c-4e67-8090-e4f49cb09379","resolution":{"observed_at":"2026-08-07T04:42:45.772873Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04434","last_updated":"2024-06-19T06:04:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-05-07T15:56:43Z","title":"DeepSeek-V2: A Strong, Economical, and Efficient Mixture-of-Experts Language Model","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.04434","snapshot_observed_at":"2026-08-07T04:42:45.775047Z","title":"Deepseek-v2: A strong, economical, and efficient mixture-of-experts language model.arXiv preprint arXiv:2405.04434, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.775047Z"},"links":{"cited_paper":"/paper/2405.04434","citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:18ac5cd8343b80c33c6324e73c25e0759aa43e602c23baaba67014afc7736757","observation_id":"db0cf834-fdce-49e5-81a8-c408da1ea1b0","resolution":{"observed_at":"2026-08-07T04:42:45.775047Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.268212Z","title":"Smpl: A skinned multi- person linear model.ACM Transactions on Graphics, 34(6): 1–16, 2015","venue":null,"work_id":"1a0e934c-5fe9-43e6-aa91-767709d49162","year":2015},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.777597Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:c2e91c8d91b21329f49f078d0fe5bac185ed1b719b3846afc9ff092773c75876","observation_id":"7ffac3f1-9b7e-4049-9d2d-21988db1b8fd","resolution":{"observed_at":"2026-08-07T04:42:46.270964Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.260764Z","title":"Nerf: Representing scenes as neural radiance fields for view syn- thesis.Communications of the ACM, 65(1):99–106, 2021","venue":null,"work_id":"e653185c-5060-46ff-b13c-7a1a06befc94","year":2021},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.779713Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:2e541ec5c3caee3270846bb916e9feeeac5e67b440f6fbfe16788c489c3a3f5e","observation_id":"32258920-16c3-4640-b4a9-5b4321e96bdc","resolution":{"observed_at":"2026-08-07T04:42:46.263378Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.253629Z","title":null,"venue":null,"work_id":"dd83c73b-9691-46ad-b4af-862f05f4f110","year":null},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.782022Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:1580f4ce7525d727cb4bd40cd58c50540861c1cc3fe09d8181fbeba0f40312bb","observation_id":"4786cd0e-2ff9-4960-8114-df1f2c88a7b3","resolution":{"observed_at":"2026-08-07T04:42:46.255942Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.03745","last_updated":"2024-05-13T16:20:29Z","snapshot_observed_at":"2026-08-16T15:26:08.420515Z","submitted_at":"2023-06-06T15:04:31Z","title":"Soft Merging of Experts with Adaptive Routing","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.03745","snapshot_observed_at":"2026-08-07T04:42:45.784299Z","title":"Soft merging of experts with adaptive routing.arXiv preprint arXiv:2306.03745, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.784299Z"},"links":{"cited_paper":"/paper/2306.03745","citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:99a51874dd6abb511e0426f6bbe650e1101735a4602db113baf7811ea18d6b7f","observation_id":"48f511f2-2856-4029-aaf3-3d658afcec93","resolution":{"observed_at":"2026-08-07T04:42:45.784299Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.246264Z","title":null,"venue":null,"work_id":"cb60cdff-31e7-44b6-b761-2a53a46fe11a","year":2025},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.786611Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:448d8e22aaf5806da726021f3c3742a94f70006e64c987eca5d901f589a91c76","observation_id":"ed7e4c87-5226-43e6-af4d-3e2526f32b4e","resolution":{"observed_at":"2026-08-07T04:42:46.248795Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:45.788863Z","title":"Expressive body capture: 3d hands, face, and body from a single image","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.788863Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:2d1bc89fdb2177802dcf263d2fc416c154586d9137df2013fb4c86149ab1254e","observation_id":"8a4c1eca-df41-4d6b-9c11-d27b43749114","resolution":{"observed_at":"2026-08-07T04:42:45.788863Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.18229","last_updated":"2024-11-27T11:07:27Z","snapshot_observed_at":"2026-08-16T20:24:32.770604Z","submitted_at":"2024-11-27T11:07:27Z","title":"SharpDepth: Sharpening Metric Depth Predictions Using Diffusion Distillation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.18229","snapshot_observed_at":"2026-08-07T04:42:45.791253Z","title":"Sharpdepth: Sharpening metric depth predictions using diffusion distillation.arXiv preprint arXiv:2411.18229, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.791253Z"},"links":{"cited_paper":"/paper/2411.18229","citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:0a675239e38d6dc4c6a95e1cf0a4e915351028e0bc0a534a7c608c49f43d1243","observation_id":"c39f6727-c0a0-4340-87d7-a4d9745b4948","resolution":{"observed_at":"2026-08-07T04:42:45.791253Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.234298Z","title":"Unidepth: Universal monocular metric depth estimation","venue":null,"work_id":"cc6b17ab-ecd9-4c76-bc8b-9553ca7be0f9","year":2024},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.793895Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:e9da5816a9a645f0a69d7fa51330a9cb5d238ee578d1ffaf0ff554886ba2e7e1","observation_id":"5a8fe38b-d432-4393-8dcb-49f746dbb7ee","resolution":{"observed_at":"2026-08-07T04:42:46.236788Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.227072Z","title":"Neural localizer fields for continuous 3d human pose and shape estimation","venue":null,"work_id":"49ce479a-ddb7-478f-8a93-d9cf0c9d738f","year":2025},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.796040Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:4aef4d60f4d73960350b37b11eb9e26a75717f340e7449adc25c4ae7c0b10581","observation_id":"5a4ebb36-8be3-4b91-ab3f-5c1924085ed9","resolution":{"observed_at":"2026-08-07T04:42:46.229584Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1701.06538","last_updated":"2017-01-23T18:10:00Z","snapshot_observed_at":"2026-08-13T11:35:07.866136Z","submitted_at":"2017-01-23T18:10:00Z","title":"Outrageously Large Neural Networks: The Sparsely-Gated Mixture-of-Experts Layer","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1701.06538","snapshot_observed_at":"2026-08-07T04:42:45.798276Z","title":"Outra- geously large neural networks: The sparsely-gated mixture- of-experts layer.arXiv preprint arXiv:1701.06538, 2017","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.798276Z"},"links":{"cited_paper":"/paper/1701.06538","citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:1395670be4304572d7c576085649aeb1bb22156d7ff1376ab972ac4ea0c5f381","observation_id":"3fbe2e64-6f87-4737-bf23-47e10c906fa1","resolution":{"observed_at":"2026-08-07T04:42:45.798276Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.219487Z","title":"World-grounded human motion recovery via gravity-view coordinates","venue":null,"work_id":"4f901ac0-2d02-4bb3-bae7-c2309798bbf2","year":2024},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.800545Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:dffb3e724a2f94bbdf270f113bf16d65f45bdf4612a82275869619b5b92847cf","observation_id":"c0aceb4f-9061-4519-b1ad-ea680a9d707a","resolution":{"observed_at":"2026-08-07T04:42:46.222225Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.212643Z","title":"Wham: Reconstructing world-grounded humans with accu- rate 3d motion","venue":null,"work_id":"de2ed5be-05e3-4d04-9d36-e78add01b7f3","year":2024},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.802608Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:db8d27cc5a008b58250fdb408618aadd8f6a07e832495500832e88dd7810b9b1","observation_id":"7864b64e-e783-41cf-9eef-4e681bb5faa5","resolution":{"observed_at":"2026-08-07T04:42:46.215083Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.205555Z","title":"Putting people in their place: Monocular regression of 3d people in depth","venue":null,"work_id":"0974f8d2-da60-432c-b495-0aab7d208b95","year":2022},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.804781Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:2f44fa609bfeab9c84595e686de54011d3efe70bf608b389670941935ea2c84d","observation_id":"6646b29a-dbec-47ef-8b02-bcdeb42663c5","resolution":{"observed_at":"2026-08-07T04:42:46.208122Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.198339Z","title":"Trace: 5d temporal regression of avatars with dynamic cam- eras in 3d environments","venue":null,"work_id":"1e0f3ed7-371f-45c4-a24b-3dfcfcd3a8b1","year":2023},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.806935Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:b37bf47868562c32ec68a3513050f953ae59f408c7b351ee4346567b135aa1da","observation_id":"9ef128d2-a89e-4fff-a341-cbe6485905f8","resolution":{"observed_at":"2026-08-07T04:42:46.200897Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.190873Z","title":"Droid-slam: Deep visual slam for monocular, stereo, and rgb-d cameras.Advances in neu- ral information processing systems, 34:16558–16569, 2021","venue":null,"work_id":"eb20198e-c2bc-4c40-9660-c04e62e82336","year":2021},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.809064Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:dbd0bf52142149451169e39c2392d437a807ee898be655fc74728162e77cca6a","observation_id":"c12d98f8-8962-4eb9-953f-d7e659bfcc05","resolution":{"observed_at":"2026-08-07T04:42:46.193758Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.183771Z","title":null,"venue":null,"work_id":"fed3cb0b-8a12-4c7d-b30e-d39d793f1941","year":2025},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.811069Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:ee53400e3547f61339ca240b061bdcba4ef28913429a802da53aaaaa1dd51566","observation_id":"644c63d8-a54e-4661-ae24-c3d5a92b63f7","resolution":{"observed_at":"2026-08-07T04:42:46.186113Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.176438Z","title":"Anycalib: On- manifold learning for model-agnostic single-view camera calibration","venue":null,"work_id":"c41d6e0f-441a-4826-b2fa-7852e045f740","year":2025},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.813245Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:045eb8572631c18890fb2963a06539951c6669dc0ac1fc04dbc3b049f5988c3e","observation_id":"306e01bc-62d1-4deb-97cd-729701fc183b","resolution":{"observed_at":"2026-08-07T04:42:46.178898Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.168895Z","title":"Recovering ac- curate 3d human pose in the wild using imus and a moving camera","venue":null,"work_id":"cf77208a-900c-4376-9fb1-1808d9ef0f2a","year":2018},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.815558Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:ccaacc2638196e6622f5b07df55879482d241c18115decdaafc0e82b5f18adff","observation_id":"66f19371-940d-424b-9227-3b9fd41bc4cd","resolution":{"observed_at":"2026-08-07T04:42:46.171493Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.160939Z","title":"Vggt: Visual geometry grounded transformer","venue":null,"work_id":"e2ce7a47-02dd-4bde-840d-634957cfafa8","year":2025},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.817884Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:483862025f6be68c5b3d22f56cda832b2934ba980e009ded59c6c6753a2a6787","observation_id":"469f907b-489f-4c16-adc1-1dc63befca61","resolution":{"observed_at":"2026-08-07T04:42:46.163868Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.153139Z","title":"Continuous 3d perception model with persistent state","venue":null,"work_id":"026e589e-5d2e-4205-a789-7f6ac8e711da","year":2025},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.820047Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:36f3c543a4cd22c307209f98d120cc038f2fd32009a19b27cb66e801692ba77e","observation_id":"b07852ec-ce29-4d5a-a7d6-0f5a7a914d71","resolution":{"observed_at":"2026-08-07T04:42:46.155612Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.08640","last_updated":"2024-12-11T18:59:08Z","snapshot_observed_at":"2026-08-16T13:00:15.524765Z","submitted_at":"2024-12-11T18:59:08Z","title":"BLADE: Single-view Body Mesh Learning through Accurate Depth Estimation","version":1},"cited_work":{"arxiv_id":"2412.08640","doi":null,"metadata_source":"pith","pith_arxiv_id":"2412.08640","snapshot_observed_at":"2026-08-07T04:42:45.888576Z","title":"BLADE: Single-view Body Mesh Learning through Accurate Depth Estimation","venue":"cs.CV","work_id":"2d28f19a-ab7d-411d-af34-175d295301e0","year":2024},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.822469Z"},"links":{"cited_paper":"/paper/2412.08640","citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:210702b110e95cf78f0c7dc9d91520e50bb88ca00aa7814c7c01ce4031780ef4","observation_id":"944882db-20d2-459b-887e-d07fa09940b0","resolution":{"observed_at":"2026-08-07T04:42:45.893233Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.145613Z","title":"Zolly: Zoom focal length correctly for perspective- distorted human mesh reconstruction","venue":null,"work_id":"58d88b47-abc7-42d3-a312-793f73044b9c","year":2023},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.825172Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:53cba4207d2168d9eb1187a51049f3da0fd32cac1575c321003c5e6779dd3e3a","observation_id":"1528197e-49e3-40f3-abb5-73d7f5d9199d","resolution":{"observed_at":"2026-08-07T04:42:46.148181Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.138071Z","title":"Tram: Global trajectory and motion of 3d humans from in- the-wild videos","venue":null,"work_id":"2ac82ffc-3f69-4000-ab47-6fd79c1f251e","year":2024},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.827427Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:8a9c7d6c7e6ffa49ee0d79ceeafb7b244241aba55e4e82d886679f631a1c4cba","observation_id":"786d8899-8962-41a1-8f55-c456af7787ad","resolution":{"observed_at":"2026-08-07T04:42:46.140667Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.130602Z","title":"Prompthmr: Promptable human mesh recovery","venue":null,"work_id":"5df0dee1-2ffd-40f4-98d4-c6880e8633b8","year":2025},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.829776Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:34ba85a81f70d05304020b59467cf523c058bbfa8ece93b37b4210c6c0158a95","observation_id":"ff7d0531-bd6b-4b18-9745-da5914bb85cd","resolution":{"observed_at":"2026-08-07T04:42:46.133190Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1711.06475","last_updated":"2017-11-17T09:58:20Z","snapshot_observed_at":"2026-08-15T23:49:50.399859Z","submitted_at":"2017-11-17T09:58:20Z","title":"AI Challenger : A Large-scale Dataset for Going Deeper in Image Understanding","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1711.06475","snapshot_observed_at":"2026-08-07T04:42:45.832197Z","title":"Ai challenger: A large-scale dataset for going deeper in image understanding.arXiv preprint arXiv:1711.06475, 2017","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.832197Z"},"links":{"cited_paper":"/paper/1711.06475","citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:478eca6d57c8a2400a58093fd88030e54adf1a4f47277c4adf7c8267e15c45d7","observation_id":"1aa37d4e-83c3-4594-bb89-7c8f27fedc16","resolution":{"observed_at":"2026-08-07T04:42:45.832197Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.122889Z","title":"Vit- pose: Simple vision transformer baselines for human pose estimation.Advances in neural information processing sys- tems, 35:38571–38584, 2022","venue":null,"work_id":"df06c25a-139a-45e9-93e8-6fdaf6f29a9e","year":2022},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.834637Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:86077b6d5379f15bd2863c2ccda1f83537d9387f772e27be1dd61092e9f75778","observation_id":"b25764b7-058e-4b81-974b-518f580e1c28","resolution":{"observed_at":"2026-08-07T04:42:46.125617Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:45.837010Z","title":"Depth any- thing v2.Advances in Neural Information Processing Sys- tems, 37:21875–21911, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.837010Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:15db2f70005838f4cf130fa456b10af51b642df10448a63ea89a89f3d438a4ae","observation_id":"cbc71bfd-4f92-4961-abda-19d4fd600181","resolution":{"observed_at":"2026-08-07T04:42:45.837010Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.110951Z","title":"Decoupling human and camera motion from videos in the wild","venue":null,"work_id":"897c4ffd-5ef1-49b1-8a27-e0c0497b74c4","year":2023},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.839211Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:c0886f4a8c8cec694513898112cec394600b64464853eddd5a555faafe2b3a57","observation_id":"fe6889a6-235c-4bed-b7dd-ac3a2f849cb9","resolution":{"observed_at":"2026-08-07T04:42:46.113506Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.103491Z","title":"Metric3d: Towards zero-shot metric 3d prediction from a single image","venue":null,"work_id":"75718068-a207-4b80-9f3d-b5e0930c2e87","year":2023},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.841442Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:d7ee900ef45145d097c5d50132221fa9ae63cb29eb934599ef18a57f395fc3fb","observation_id":"f467d747-453d-49ad-b9f6-dfb1191f32b2","resolution":{"observed_at":"2026-08-07T04:42:46.106321Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.096446Z","title":"Function4d: Real-time human vol- umetric capture from very sparse consumer rgbd sensors","venue":null,"work_id":"54371aaf-024d-4368-a28d-6ce48e5e5c80","year":2021},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.843658Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:3ddd6b68772287f87bac621c94f76c8a523b109b4d0e67d7d10dace905a3bf16","observation_id":"308f5d25-6b2b-4e22-92d7-9946ad3c45ad","resolution":{"observed_at":"2026-08-07T04:42:46.098948Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.089060Z","title":"Glamr: Global occlusion-aware human mesh recov- ery with dynamic cameras","venue":null,"work_id":"ad541fb4-fd3c-4f1b-940b-01721eba97bc","year":2022},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.845925Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:89a0cc34a2f22bfea23b6798e5193717f37751b726e516b541fd5b9e0f22aa8e","observation_id":"c06391ca-1ba1-4fc1-a06b-ac437f1d237b","resolution":{"observed_at":"2026-08-07T04:42:46.091593Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.081166Z","title":"Not all tokens are equal: Human-centric visual analysis via token clustering transformer","venue":null,"work_id":"4fd6ac8b-64c6-4317-96d5-5f720aa16e9c","year":2022},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.848082Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:637e81de457257f5775289a2f1ae956c39a360165442e7375fc08072fe37192a","observation_id":"82dc4ebb-eae6-4969-8ec4-510c2d74ed17","resolution":{"observed_at":"2026-08-07T04:42:46.083763Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.072507Z","title":"Pymaf: 3d human pose and shape regression with pyramidal mesh alignment feedback loop","venue":null,"work_id":"1938a9a7-653a-4eea-8ba5-c580b849c98d","year":null},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.850192Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:2558fd1da65adedfffaacc31d21b13d0cd33b1069b8e251831f13c09557b9e13","observation_id":"77551ad2-cd78-40aa-a915-acf10aa3144b","resolution":{"observed_at":"2026-08-07T04:42:46.075871Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.064890Z","title":"Pymaf-x: To- wards well-aligned full-body model regression from monoc- ular images.IEEE Transactions on Pattern Analysis and Ma- chine Intelligence, 45(10):12287–12303, 2023","venue":null,"work_id":"d63c68ac-7235-4f96-81fd-4cd481bd4bff","year":2023},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.852554Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:01892a2293f3838cc1cd1f75736ceb2e097238a4c99873a0db7b4372dcab4097","observation_id":"16f3826a-44c9-4141-ad6a-23d4be0baf77","resolution":{"observed_at":"2026-08-07T04:42:46.067689Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.057245Z","title":"Metric from human: Zero-shot monoc- ular metric depth estimation via test-time adaptation","venue":null,"work_id":"3ff98c8d-84b1-424d-b0f3-172081d738a2","year":2024},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.854769Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:e6ebc76728a346058eaa8247bb0084f6c392a0949189d57f176dcb2796249bf1","observation_id":"65c2424f-b382-4c64-98f3-d67003c99000","resolution":{"observed_at":"2026-08-07T04:42:46.059997Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.049773Z","title":"Camera, Image and Metrics In human mesh recovery, full perspective projection [2, 17, 28, 53] has gained increasing attention","venue":null,"work_id":"58299794-9b3c-4d89-a774-808ea8eb8fe3","year":null},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.856929Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:59a8324f434c9518e6269573fd3cdfa659564fdcb73b8e71d8686ecd76cbbe0a","observation_id":"97027569-3749-4383-b251-cec10497e44b","resolution":{"observed_at":"2026-08-07T04:42:46.052305Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:42:46.041904Z","title":"SynFocal dataset","venue":null,"work_id":"cae10978-c16d-4c86-8a78-4ff8c6689c58","year":null},"citing_paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images","version":4},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-07T04:42:45.859616Z"},"links":{"citing_paper":"/paper/2506.09919"},"observation_digest":"sha256:c10d280766a4e6688a9059a20ecebd81a3823575a37564a40de9ffdf3d98ce19","observation_id":"cd0c1b06-b8eb-41c1-8e97-778cd9fdccfb","resolution":{"observed_at":"2026-08-07T04:42:46.044733Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.09919","last_updated":"2026-06-30T08:44:12Z","latest_version":4,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-16T20:25:11.837827Z","submitted_at":"2025-06-11T16:39:23Z","title":"MetricHMSR:Metric Human Mesh and Scene Recovery from Monocular Images"},"reference_resolution":{"displayed":67,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":20,"verified_exact":1,"verified_fuzzy":46},"total_outbound_references":67},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"thesis":"As of 18 August 2026, this Paper Citation Record lists 67 of 67 outbound references and 1 inbound Pith citation observation for arXiv:2506.09919."}