{"as_of":"2026-08-15T02:27:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:6c5a975961ef47212d58c4a120608136ca475bac5bf5ca7fc47504bd0b5e71da","coverage":[{"denominator":66,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":66,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T22:04:45.141516Z","state":"measured"},{"denominator":66,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":66,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-14T06:32:32.682623+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2501.02966/citation-record","integrity":"/paper/2501.02966/integrity","json":"/paper/2501.02966/citation-record.json","paper":"/paper/2501.02966"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:46.252417Z","title":"Chart demonstrating variations in acuity with retinal position","venue":null,"work_id":"4ed67bd4-1aa1-492f-b5cd-298ffd0d5736","year":1974},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.838671Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:869ae92eb33ee5bdcb972f4707c27bd7f0c4fb943863ed1ec1f3b79d5e21ae6c","observation_id":"cac18bab-4f5c-4041-9734-42888d4345ed","resolution":{"observed_at":"2026-08-10T22:04:46.257162Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2207.13492","last_updated":"2022-12-21T10:55:11Z","snapshot_observed_at":"2026-08-13T14:55:41.132015Z","submitted_at":"2022-07-27T12:27:57Z","title":"Time to augment self-supervised visual representation learning","version":2},"cited_work":{"arxiv_id":"2207.13492","doi":null,"metadata_source":"pith","pith_arxiv_id":"2207.13492","snapshot_observed_at":"2026-08-10T22:04:45.392765Z","title":"Time to augment self-supervised visual representation learning","venue":"cs.LG","work_id":"13854fc2-31f4-444d-90f6-2b00494f9dd4","year":2022},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.844055Z"},"links":{"cited_paper":"/paper/2207.13492","citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:662cdb0d6886f244a4d71c43516c98352ab192eea17ce1f2eb77e20f32ed6ff4","observation_id":"3e2133d3-ac5f-42f4-a24a-ca89ae1200a9","resolution":{"observed_at":"2026-08-10T22:04:45.397978Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:46.237088Z","title":"Toddler- inspired embodied vision for learning object representations","venue":null,"work_id":"6a319c3f-3ef4-42a6-9546-2cf70f303b10","year":2022},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.849396Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:ecc220f70d6fc2d91ddfee7088ebcd234d51b29e378eca9eae1aa582f851ee6a","observation_id":"780c8886-b0c5-476b-acfe-b41e6dca6bd0","resolution":{"observed_at":"2026-08-10T22:04:46.241942Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.05143","last_updated":"2024-04-19T14:08:17Z","snapshot_observed_at":"2026-08-14T02:13:48.369091Z","submitted_at":"2024-04-19T14:08:17Z","title":"Learning Object Semantic Similarity with Self-Supervision","version":1},"cited_work":{"arxiv_id":"2405.05143","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.05143","snapshot_observed_at":"2026-08-10T22:04:45.372187Z","title":"Learning Object Semantic Similarity with Self-Supervision","venue":"cs.CV","work_id":"e95ef8c6-d216-40d1-8af6-d012d0846fc5","year":2024},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.854291Z"},"links":{"cited_paper":"/paper/2405.05143","citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:766766b647892b65840e344310f69e61b6bb5aa9bf504f82e3a8426637ab3e41","observation_id":"7af5028a-c280-4cd1-99ba-22ec4629e687","resolution":{"observed_at":"2026-08-10T22:04:45.376955Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.06704","last_updated":"2024-08-08T09:41:40Z","snapshot_observed_at":"2026-08-12T23:26:01.036853Z","submitted_at":"2024-07-09T09:31:15Z","title":"Self-supervised visual learning from interactions with objects","version":2},"cited_work":{"arxiv_id":"2407.06704","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.06704","snapshot_observed_at":"2026-08-10T22:04:45.352131Z","title":"Self-supervised visual learning from interactions with objects","venue":"cs.CV","work_id":"29126db8-df01-4a7a-9571-b4263cfe1aee","year":2024},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.859723Z"},"links":{"cited_paper":"/paper/2407.06704","citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:049ffd1eb5e4f7cc8bb25b9e0ce0b2be2527feb913d039c50ec73c413ec60f7e","observation_id":"130a6969-c6a2-449f-8351-a47ed7b225ee","resolution":{"observed_at":"2026-08-10T22:04:45.356991Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:44.864685Z","title":"Emerg- ing properties in self-supervised vision transformers","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.864685Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:fc9551f082cca9b20c4a9d4d3c1f5b8f1527104a83d8c28f3184476f32bbef24","observation_id":"e28c4c5e-350f-42ea-ab22-08edfae272a5","resolution":{"observed_at":"2026-08-10T22:04:44.864685Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:46.211222Z","title":"Big self-supervised mod- els are strong semi-supervised learners","venue":null,"work_id":"fc52f465-38ab-44e2-ad5d-abcb164d9266","year":2020},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.870278Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:18852ddb52329a9617f92527e443304ea62f049e1e1daab819e1ce92be5f33bf","observation_id":"b90c1c92-fea6-4bfa-861e-c365470e9552","resolution":{"observed_at":"2026-08-10T22:04:46.216113Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:46.196000Z","title":"An empirical study of training self-supervised vision transformers","venue":null,"work_id":"5abb9dc9-c618-48c0-b6d8-5ed6974912be","year":2021},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.874555Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:910c5611240a7f79dfb74a33ccd97e9bfeda0bf4364d6410c56e81aa8aa3dfb0","observation_id":"b3d0cbcf-5fbd-4779-9d7a-b48e8a793b53","resolution":{"observed_at":"2026-08-10T22:04:46.201315Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:46.179915Z","title":"Describing textures in the wild","venue":null,"work_id":"1074c962-226e-47dd-8ab5-22c02a42a88e","year":2014},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.879010Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:6304260176b84f13b4812f83ff461fd39ac57122cf06e09d6e33d92e754d71ae","observation_id":"5db4808a-bedd-4841-9b7d-912d77d327ef","resolution":{"observed_at":"2026-08-10T22:04:46.184727Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:46.164580Z","title":"An analysis of single-layer networks in unsupervised feature learning","venue":null,"work_id":"91c2d221-4191-4166-8f7c-e20a0de9cd94","year":2011},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.883336Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:09144abc2a58fe05d341bc852605ee0138d2bca8acb86cd179a14bac23541970","observation_id":"1bc6e465-7dac-40db-b83f-2b7ecbc5e788","resolution":{"observed_at":"2026-08-10T22:04:46.169178Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:46.149310Z","title":"solo-learn: A library of self- supervised methods for visual representation learning","venue":null,"work_id":"088cbb65-a33a-4081-8017-f81266d80f99","year":2022},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.887566Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:1ea8f4b3a2402702b72c4e29926ca9765cfa7270a6c9c0168b45f4ec647d0a51","observation_id":"da5c18cf-4109-4fe4-a600-08edb3f4930f","resolution":{"observed_at":"2026-08-10T22:04:46.154671Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:46.135003Z","title":"Hvm-1: Large-scale video models pretrained with nearly 5000 hours of human-like video data","venue":null,"work_id":"72375675-0f4c-4afa-8cd4-8a5f48c19c3b","year":2024},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.892263Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:d916d2ff010e0d52d319a966661634fcff0a6445d7dbd6d3fbaa59c80e161cea","observation_id":"11741c1d-3945-457e-b637-44601fa1b69a","resolution":{"observed_at":"2026-08-10T22:04:46.139592Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:46.120218Z","title":"Learning invariance from transformation se- quences","venue":null,"work_id":"721138ee-d2da-49e2-b09c-3cb1878e483f","year":1991},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.896724Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:0f4992bde33b5e9f2972673b0c88f128ebff8897e44dfeb22bd2724d7e4f881b","observation_id":"1b9b4a59-fae6-40ad-bb49-aa55e30f822b","resolution":{"observed_at":"2026-08-10T22:04:46.124833Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:44.900972Z","title":"Partial success in closing the gap between human and machine vision","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.900972Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:80056689a81b1fdece2f2515d188bc07aee94743626f9b4b4b54ee4962d01631","observation_id":"444597c8-96a8-43c7-9c4a-a3a9a3392bad","resolution":{"observed_at":"2026-08-10T22:04:44.900972Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2003.07990","last_updated":"2020-05-07T17:23:14Z","snapshot_observed_at":"2026-08-14T06:47:34.820346Z","submitted_at":"2020-03-18T00:07:21Z","title":"Watching the World Go By: Representation Learning from Unlabeled Videos","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2003.07990","snapshot_observed_at":"2026-08-10T22:04:44.905361Z","title":"Watching the world go by: Representation learning from un- labeled videos","venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.905361Z"},"links":{"cited_paper":"/paper/2003.07990","citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:bd288e640c98abd33c2f3318a31bf624b5c538bccd4972f8867d85fbf565e0b7","observation_id":"f43079da-7722-4fd2-9673-964553dbc5b1","resolution":{"observed_at":"2026-08-10T22:04:44.905361Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1706.02677","last_updated":"2018-04-30T21:53:41Z","snapshot_observed_at":"2026-08-09T05:23:26.365677Z","submitted_at":"2017-06-08T16:51:53Z","title":"Accurate, Large Minibatch SGD: Training ImageNet in 1 Hour","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1706.02677","snapshot_observed_at":"2026-08-10T22:04:44.909987Z","title":"Girshick, Pieter No- ordhuis, Lukasz Wesolowski, Aapo Kyrola, Andrew Tul- 7 loch, Yangqing Jia, and Kaiming He","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.909987Z"},"links":{"cited_paper":"/paper/1706.02677","citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:3b7bc5bebaab61a5e09e23e75841abce3f504b387fe0685eb6ab317a51a2381e","observation_id":"903c4596-32df-411d-b602-2411742f278e","resolution":{"observed_at":"2026-08-10T22:04:44.909987Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:46.096240Z","title":null,"venue":null,"work_id":"34bdc87d-cb3a-4746-99d7-cbc3e75bceaa","year":2022},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.914692Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:327180407ca0345e7945fcc91e4c5e4a8ce6adf095f38679ced05cb611edf11a","observation_id":"ef61bb56-067f-4a1b-a884-99f0ad8648e1","resolution":{"observed_at":"2026-08-10T22:04:46.100931Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:44.918979Z","title":"Ego4d: Around the world in 3,000 hours of egocentric video","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.918979Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:d55641c22a525ef43c64bf635b653f36411075b0114d0ef7275c3df52b487a3d","observation_id":"29bdaafe-7741-4ba4-9c33-f2eeb1177472","resolution":{"observed_at":"2026-08-10T22:04:44.918979Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18934","last_updated":"2024-08-13T16:01:14Z","snapshot_observed_at":"2026-08-13T04:18:27.004680Z","submitted_at":"2024-02-15T10:34:28Z","title":"The Visual Experience Dataset: Over 200 Recorded Hours of Integrated Eye Movement, Odometry, and Egocentric Video","version":2},"cited_work":{"arxiv_id":"2404.18934","doi":null,"metadata_source":"pith","pith_arxiv_id":"2404.18934","snapshot_observed_at":"2026-08-10T22:04:45.299041Z","title":"The Visual Experience Dataset: Over 200 Recorded Hours of Integrated Eye Movement, Odometry, and Egocentric Video","venue":"cs.CV","work_id":"ba8e1567-a53f-4928-8a2a-d0cdb9b32184","year":2024},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.923205Z"},"links":{"cited_paper":"/paper/2404.18934","citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:2dff2fc3d6bc94b3c49821f264472e97f16b5895a005162d8e6c661cafaad63d","observation_id":"8146818a-7d06-4438-bd43-828f78ea3bc8","resolution":{"observed_at":"2026-08-10T22:04:45.304466Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:46.072643Z","title":"Bootstrap your own latent-a new approach to self-supervised learning","venue":null,"work_id":"b6a66d96-b1d5-4909-a1f3-83abee6b1815","year":2020},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.927973Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:466722eb899fd1dab73b6f9d9ce21239f8c7479e773a628e69f47e67c8ddad62","observation_id":"3b823abd-fecd-4d14-826e-0c13617a4cdb","resolution":{"observed_at":"2026-08-10T22:04:46.077433Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:46.057857Z","title":"Deep residual learning for image recognition","venue":null,"work_id":"647fc346-9b10-494c-99e5-4e13c5a8b199","year":2016},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.932040Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:67b2dfc9fde6820ae56272a0dd243825244e88d921324fcd3bb871209165cb4b","observation_id":"00ba6c55-e974-48b4-aa1f-a2ed12205550","resolution":{"observed_at":"2026-08-10T22:04:46.062471Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:44.936200Z","title":"Masked autoencoders are scalable vision learners","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.936200Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:aef8fd27eb2dfbb09216df63d306317db39f24d96079730ea8dfc0411af2c71b","observation_id":"2c1c665c-bfad-4c21-bb80-22240b0f1fe4","resolution":{"observed_at":"2026-08-10T22:04:44.936200Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:46.033338Z","title":"Space-time correspondence as a contrastive random walk","venue":null,"work_id":"888836ab-369f-4f64-a8b2-b82cf31e0dc2","year":null},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.940447Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:8ca7ed8333c34aaca2f0a4e8b60a5512e36babd1958e66f4d3ce4cc71118ef7b","observation_id":"5099586b-24aa-4e2b-a85b-d807ea05bff4","resolution":{"observed_at":"2026-08-10T22:04:46.038251Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:46.018295Z","title":"Learning image representations tied to ego-motion","venue":null,"work_id":"dd2f83dc-7d71-4631-959d-ddf2a6efeb70","year":2015},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.944973Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:f4aefff3d4b8f31e37b149b94066e44b940f8b600537b21d8e119dbb6b8994ba","observation_id":"e1c1caa2-6aeb-4d64-97df-f49ce1931f29","resolution":{"observed_at":"2026-08-10T22:04:46.023177Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:46.002367Z","title":"Slow and steady feature analysis: higher order temporal coherence in video","venue":null,"work_id":"ce9a485b-d27c-46af-8921-65a3b8fa832a","year":2016},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.949399Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:7d0e4aaad5a4fd315525dfa28b2646fea6fbd8f77456dae4934eb5ffbc7a3e20","observation_id":"b4e07363-86ef-4b67-9785-95bddf0ddad1","resolution":{"observed_at":"2026-08-10T22:04:46.007942Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.986205Z","title":"3d object representations for fine-grained categorization","venue":null,"work_id":"cba9c836-6f3d-4217-aa38-af5ae86abc97","year":2013},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.953823Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:c4bfd9cc68d01f4e93e2e184677ed227175863625c7351db14d0e41cd9a9e1b9","observation_id":"0c19499b-238b-4b20-81bf-3cd66dd80a0d","resolution":{"observed_at":"2026-08-10T22:04:45.991275Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.970503Z","title":"Learning multiple layers of features from tiny images","venue":null,"work_id":"c83ede01-e9f5-44e0-8a4f-46d3fe0f3d99","year":2009},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.958226Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:4e3f5988f82a6b69f7392de016f007c57846b0deef9aeaf0a9ceab2ac8aca595","observation_id":"cdbae478-740b-4fcf-9cb4-3e1d073cfc1d","resolution":{"observed_at":"2026-08-10T22:04:45.975197Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.955130Z","title":"In the eye of transformer: Global-local correlation for egocentric gaze estimation","venue":null,"work_id":"2d1b3792-a088-4fc6-9f2f-3b4cc7154d4d","year":2022},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.962804Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:ef9db06d4b9ae4a8a4e612329c28940b84c9fb4a70fff111d0f1e94b2f362656","observation_id":"2a81c838-7008-41eb-878f-3e5ddb2a08d3","resolution":{"observed_at":"2026-08-10T22:04:45.959976Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.940208Z","title":"Core50: a new dataset and benchmark for continuous object recognition","venue":null,"work_id":"67a32a30-32d2-4dbd-995e-9353f638e2ff","year":2017},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.966938Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:8ea5a70b9ee65e2a307e0234900592063a06d62751c95ec8bd427c524509c9bc","observation_id":"1930a14b-64dd-405e-bf02-251a4d5c3aed","resolution":{"observed_at":"2026-08-10T22:04:45.945234Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10447","last_updated":"2025-07-22T22:34:57Z","snapshot_observed_at":"2026-08-12T23:42:19.272655Z","submitted_at":"2024-06-14T23:52:27Z","title":"The BabyView dataset: High-resolution egocentric videos of infants' and young children's everyday experiences","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10447","snapshot_observed_at":"2026-08-10T22:04:44.971497Z","title":"The babyview dataset: High-resolution egocentric videos of in- fants’ and young children’s everyday experiences","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.971497Z"},"links":{"cited_paper":"/paper/2406.10447","citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:31447af6330b11d2c557495254c33e36466a87e8bc792427785e0fc1de8f8662","observation_id":"762ce1c1-4cff-4f91-aac2-0e1066aa7473","resolution":{"observed_at":"2026-08-10T22:04:44.971497Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.09905","last_updated":"2024-09-20T01:18:11Z","snapshot_observed_at":"2026-08-12T23:42:49.268895Z","submitted_at":"2024-06-14T10:23:53Z","title":"Nymeria: A Massive Collection of Multimodal Egocentric Daily Motion in the Wild","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.09905","snapshot_observed_at":"2026-08-10T22:04:44.976071Z","title":"Nymeria: A massive collection of multimodal egocentric daily motion in the wild","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.976071Z"},"links":{"cited_paper":"/paper/2406.09905","citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:6173b135081448c471f89ce88c475702e813657b52eb3523e6ef5df97e328948","observation_id":"45756f10-4192-4ef2-a74b-59ccf945eaf1","resolution":{"observed_at":"2026-08-10T22:04:44.976071Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.925290Z","title":"Vip: Towards universal visual reward and representation via value-implicit pre-training","venue":null,"work_id":"89914400-a597-4373-b84c-a8092c129fb1","year":null},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.980501Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:44455a27f85613c29400b342eb75d9730fbc39fa0a9c2b46e8ab0b6ef6dd74c6","observation_id":"d897e606-0e65-4a3a-b1b5-0339a8483fbc","resolution":{"observed_at":"2026-08-10T22:04:45.930216Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1306.5151","last_updated":"2013-06-21T14:31:57Z","snapshot_observed_at":"2026-08-12T17:35:23.022229Z","submitted_at":"2013-06-21T14:31:57Z","title":"Fine-Grained Visual Classification of Aircraft","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1306.5151","snapshot_observed_at":"2026-08-10T22:04:44.985110Z","title":"Blaschko, and Andrea Vedaldi","venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.985110Z"},"links":{"cited_paper":"/paper/1306.5151","citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:4a08b8d387b6a2ad609bf5dfe6133e082cb8611afd04ab3a4254c63894e7ce45","observation_id":"fac9742a-c793-4019-a98b-e15896ff52e4","resolution":{"observed_at":"2026-08-10T22:04:44.985110Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.810857Z","title":"Where are we in the search for an artificial visual cortex for embodied intelli- gence? Advances in Neural Information Processing Systems, 36:655–677, 2023","venue":null,"work_id":"f4facc7b-ed78-471a-b02f-81b319047ea2","year":2023},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.989944Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:2e938bf864aeca82a08eea7d4ae583f763bb3a8923493f1a4c05654a328d0a2d","observation_id":"c41427a7-d32f-44d2-9a67-a9cd545b14a8","resolution":{"observed_at":"2026-08-10T22:04:45.914844Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:44.994360Z","title":"R3m: A universal visual repre- sentation for robot manipulation","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.994360Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:005c8989b310ca2a9ca2115cb0ccbe5888d3dd9941bc190972cc50e36a2857c5","observation_id":"44a15815-32f4-4522-8878-e0acce51c4fe","resolution":{"observed_at":"2026-08-10T22:04:44.994360Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.786046Z","title":"Columbia object image library (coil-20)","venue":null,"work_id":"fa2d0607-fdfa-481b-8b99-0d7206e30541","year":1996},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:44.998793Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:693676a87da0564333a5e57fa49eee5b96be4a864ce1be9bc5b3d14a26cd2961","observation_id":"26ca85d3-23ca-4367-994c-104e236f9b32","resolution":{"observed_at":"2026-08-10T22:04:45.790821Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.770936Z","title":"Automated flower classification over a large number of classes","venue":null,"work_id":"b49d780b-85d6-487d-aafd-0a12327841ab","year":2008},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.003015Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:3ca6b846b6f3f1e4e2be359dc7628589e713d806b7c6778cd8515bc390557cb0","observation_id":"691696eb-b946-4303-a96a-249910bce304","resolution":{"observed_at":"2026-08-10T22:04:45.775536Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.03712","last_updated":"2023-08-10T16:23:03Z","snapshot_observed_at":"2026-08-13T10:39:13.544920Z","submitted_at":"2023-08-07T16:31:38Z","title":"Scaling may be all you need for achieving human-level object recognition capacity with human-like visual experience","version":2},"cited_work":{"arxiv_id":"2308.03712","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.03712","snapshot_observed_at":"2026-08-10T22:04:45.228969Z","title":"Scaling may be all you need for achieving human-level object recognition capacity with human-like visual experience","venue":"cs.CV","work_id":"1a8a1305-3ea7-4b0a-9b8a-9d0872b1cde3","year":2023},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.007792Z"},"links":{"cited_paper":"/paper/2308.03712","citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:80bc21c67cefa81a9e4f9d0f7d6c1e149466141c2586e97de005c593c661d217","observation_id":"15f598a4-1cbe-4072-9b39-f6f59da519ca","resolution":{"observed_at":"2026-08-10T22:04:45.236356Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.756111Z","title":"Learning high-level vi- sual representations from a child’s perspective without strong inductive biases","venue":null,"work_id":"c487d71f-e22e-4a52-a147-4554fd30d43c","year":2024},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.012715Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:41dd71816ef6f477bc668e54cb1b5374ff9545ab95d1bc3bd921ee230c8367b5","observation_id":"3d1ce256-548d-49cf-a0cf-fd53c2f9d3f7","resolution":{"observed_at":"2026-08-10T22:04:45.760993Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.00300","last_updated":"2024-10-16T19:10:03Z","snapshot_observed_at":"2026-08-13T04:29:35.750017Z","submitted_at":"2024-02-01T03:27:26Z","title":"Self-supervised learning of video representations from a child's perspective","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.00300","snapshot_observed_at":"2026-08-10T22:04:45.017469Z","title":"Self-supervised learning of video representations from a child’s perspective","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.017469Z"},"links":{"cited_paper":"/paper/2402.00300","citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:f4494b42cdeb53216831f203aa0df301845a24e449d753f8d96f670a1a8aa279","observation_id":"74c1a2fb-8fb1-45ff-8bb0-d2533efaa7c3","resolution":{"observed_at":"2026-08-10T22:04:45.017469Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.741610Z","title":"Self- supervised learning through the eyes of a child","venue":null,"work_id":"8f39fbf8-54e4-44c2-b51e-8e2017278c0c","year":null},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.022545Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:c8dd16c653d865ccb479d67630b0bac81edf34c1396f7d3770d56d5ae0950e8c","observation_id":"3eab0ae2-719a-4eb5-9f8f-c9a6f3073929","resolution":{"observed_at":"2026-08-10T22:04:45.746481Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.726108Z","title":"Are vi- sion transformers more data hungry than newborn visual sys- tems? Advances in Neural Information Processing Systems, 36, 2024","venue":null,"work_id":"458260ad-1c8b-446d-9577-7488cc05fdf0","year":2024},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.027607Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:744519c80f41d6939ceff3ec195b2402e5a7e0f38d2d00ca4845742cb1c1031b","observation_id":"36bf38e8-80dc-4af9-ae4d-d3bdd99f858c","resolution":{"observed_at":"2026-08-10T22:04:45.731151Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.711666Z","title":null,"venue":null,"work_id":"29e41b26-e408-4ed5-94c5-2032afe5b54d","year":2012},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.032426Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:b08ca16241958940f91115965a5baf35a997b2bad21e4b0de2388ea0e94344bd","observation_id":"7739606c-fd0d-4140-be32-8c905dc0ac4e","resolution":{"observed_at":"2026-08-10T22:04:45.716257Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.696540Z","title":"Self-supervised video pretraining yields robust and more human-aligned visual representations","venue":null,"work_id":"11ef1c46-d08e-4a2c-ac32-b5a29c59ac7b","year":null},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.037676Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:286a8e532ffe349b826283a1901f1ff61c1c252ced4bfdaca448f77ec10c367b","observation_id":"2aa19d88-7f6f-46e1-b17e-52fa937a7cdb","resolution":{"observed_at":"2026-08-10T22:04:45.701646Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.682007Z","title":"Object recognition in primates: What can early visual areas contribute? Fron- tiers in Behavioral Neuroscience, 18:1425496, 2024","venue":null,"work_id":"1190e99c-f204-4a4a-99fc-3d0a0661a5d6","year":2024},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.043018Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:32ca0b735c6b0a2ccac754f2a1c5110d84d50aa329077cc409eb38b388a626d5","observation_id":"38aa210b-761e-4372-9b4d-9025d25e93fd","resolution":{"observed_at":"2026-08-10T22:04:45.686871Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.667698Z","title":"Bottom-up saliency mod- els for still images: A practical review","venue":null,"work_id":"3396e9c8-d945-4f94-8905-c65d146c1279","year":2016},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.047505Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:7befaea829d97fef1c0cbb69c638ae0e449f43d46b540be91900a8278ee409aa","observation_id":"acd7b81f-b9c0-4115-ba9a-a82685ab67be","resolution":{"observed_at":"2026-08-10T22:04:45.672310Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.653113Z","title":"Imagenet large scale visual recognition challenge","venue":null,"work_id":"ed2ea438-7609-4881-84bb-0dbf9733e2e9","year":2015},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.052110Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:78df1f998ca8fb164920032661b8105624541df936839196b36a8efd418bfe37","observation_id":"d0d531d5-b04a-43b7-bca1-a12bb3ec2c08","resolution":{"observed_at":"2026-08-10T22:04:45.657823Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.637775Z","title":"Time does tell: Self-supervised time- tuning of dense image representations","venue":null,"work_id":"5c343f0f-53fd-48e3-a8df-15b8680ae1f5","year":2023},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.056682Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:54d8289d74b0e3d62b99eb3e06a0a6f83eb478268d45dad83675003ee3aa4176","observation_id":"cdb3576e-9bc8-4832-a6dc-3a346111b8a9","resolution":{"observed_at":"2026-08-10T22:04:45.642931Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.622380Z","title":"A computational account of self-supervised visual learning from egocentric object play","venue":null,"work_id":"c693c4e1-e0ab-4c3c-a984-87e890f13aaf","year":2023},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.061626Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:67033326462c0be919d3a627769ce20f6f703e37ae48574d44a8b44439c03565","observation_id":"1b3950bf-d791-4e6e-a912-54dea90688a3","resolution":{"observed_at":"2026-08-10T22:04:45.627050Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.606666Z","title":"Caregiver talk shapes toddler vision: A computational study of dyadic play","venue":null,"work_id":"507544c3-6cf9-4cd3-ae53-a7719385ac3c","year":2023},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.066095Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:d5867eb27cf9f9137a9956a19c098e7b20cbfa2460e28fdfb8baa422676bfbd9","observation_id":"0112cbcc-a875-4a01-b29f-c6db3bb56e55","resolution":{"observed_at":"2026-08-10T22:04:45.611402Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.591816Z","title":"Time-contrastive networks: Self-supervised learn- ing from video","venue":null,"work_id":"8810a5ed-b459-44b3-8314-ede1a93bed4f","year":2018},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.070644Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:b514eca68a6cec36a63272ee6d75b5141831e36706da84580ad807d8f66e5524","observation_id":"7fef5805-4560-4307-87c1-aeee7d82bec4","resolution":{"observed_at":"2026-08-10T22:04:45.596650Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.576880Z","title":"Curriculum learning with infant egocentric videos","venue":null,"work_id":"867df86b-3511-418c-9ef6-42fe167c11d2","year":2024},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.075457Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:1db10c7eeae86e600c35909d305f7e5060a14cd106c9fd1350813ac4a4132ece","observation_id":"d92bde83-47e7-4b86-9aff-741ad5d676c8","resolution":{"observed_at":"2026-08-10T22:04:45.581892Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.562504Z","title":"Saycam: A large, longitudinal audiovisual dataset recorded from the infant’s perspective","venue":null,"work_id":"f182fba3-2be4-4744-bb37-e0b57cb2d8cf","year":2021},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.080109Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:36d3474a388dd2f15b8af96c2073181bdc3968f9a0d8568018fd92c6329842bb","observation_id":"ce4e2245-2d97-423d-8068-f7dba36a03b4","resolution":{"observed_at":"2026-08-10T22:04:45.567384Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.084550Z","title":"Con- trastive multiview coding","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.084550Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:adab2e73a2195dd41a08cf5305535cc87ff83aee6ddcc64e0888a1ee124a7491","observation_id":"03f5751e-f32b-4a02-a57f-bb13b38f9517","resolution":{"observed_at":"2026-08-10T22:04:45.084550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.089501Z","title":"Repre- sentation learning with contrastive predictive coding, 2019","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.089501Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:683f5d1987cfdb9958cc34c5a4e66081068104b8f9daebaee5a6a615a931c087","observation_id":"4ac6b865-ad06-47a8-8bb5-b934d4517960","resolution":{"observed_at":"2026-08-10T22:04:45.089501Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.528635Z","title":"Is imagenet worth 1 video? learning strong image encoders from 1 long unlabelled video","venue":null,"work_id":"312a7715-7125-4306-97cb-bae32fe97fd3","year":2024},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.094468Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:d7b80ecf2f2c0666447a482bd8ed0fcc6776e672d1bef93f325aeef9967c2d9f","observation_id":"0b883b7c-d65f-405d-8445-bfb1091a6679","resolution":{"observed_at":"2026-08-10T22:04:45.533308Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.514225Z","title":"On the use of cortical magnification and saccades as biological proxies for data augmentation","venue":null,"work_id":"19fa563a-5150-44d4-92dd-2cac5a34efc8","year":2021},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.099160Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:3178754b87d2d2f7aa23ef92c75f2a06b09c5a09d040a85b24a527fc4a5b4774","observation_id":"8349f479-3ee6-4b2e-a82a-d850d6ab1dec","resolution":{"observed_at":"2026-08-10T22:04:45.519055Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.499679Z","title":"Eliott, James Ainooson, Joshua H","venue":null,"work_id":"24f155a8-314b-49d5-af95-8cbc663a7c77","year":2017},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.103500Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:8e71694277f5179dfc570e5cad192e1150c541d8c4295c6bfca8913591e54719","observation_id":"36b2bc3c-d6cb-443d-b06b-fc468a06a508","resolution":{"observed_at":"2026-08-10T22:04:45.504430Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.485397Z","title":"Cortical magnification factor and the gan- glion cell density of the primate retina","venue":null,"work_id":"2cf8f960-b01d-4f9e-8714-8a7b4a7942b6","year":1989},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.108168Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:8c122dc983c7e4f957291e3a4366a370fcfb7a227dec72d26d9f8b39001b2d63","observation_id":"01ddb606-0c5d-4ef0-a19f-8ac6b8cfc0c7","resolution":{"observed_at":"2026-08-10T22:04:45.490205Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.469629Z","title":"Slow feature analysis: Unsupervised learning of invariances","venue":null,"work_id":"97e9d107-ed8d-4721-a533-274635f04e0b","year":2002},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.112973Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:f93302055954f3aa29c14ace838e0e57bc11047c150e491998b9355fbc8e5144","observation_id":"29718189-a083-4c28-ab7d-babd46611227","resolution":{"observed_at":"2026-08-10T22:04:45.474740Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.454996Z","title":"Noise or signal: The role of image backgrounds in object recognition","venue":null,"work_id":"a52543f7-59b1-4e16-9876-929cffc954e7","year":2021},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.117673Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:04992d5ae387e835f12ed1c54e1079f310e9ece15c81a5911797240c4811b181","observation_id":"11b410af-7864-4291-8220-26cb065c2705","resolution":{"observed_at":"2026-08-10T22:04:45.459951Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.438982Z","title":"Rethinking self-supervised correspondence learning: A video frame-level similarity per- spective","venue":null,"work_id":"f09dc97e-01e7-4821-b370-f555e3462cfc","year":2021},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.122447Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:9f7f9f5cfed59f98f85aac4ce2efc02a6de919025396ab35e9d7ff0081bd537e","observation_id":"cd89ce47-6827-4586-97bc-f9588e8abfb7","resolution":{"observed_at":"2026-08-10T22:04:45.444485Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1708.03888","last_updated":"2017-09-13T23:25:07Z","snapshot_observed_at":"2026-08-14T20:41:41.185318Z","submitted_at":"2017-08-13T11:01:57Z","title":"Large Batch Training of Convolutional Networks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1708.03888","snapshot_observed_at":"2026-08-10T22:04:45.127222Z","title":"Scaling SGD batch size to 32k for imagenet training","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.127222Z"},"links":{"cited_paper":"/paper/1708.03888","citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:b98e9fb3cefa18374ace31a7154a66834912c4e53d0f88322d3d65c0a24abe3c","observation_id":"604c4d06-0d18-43ee-9758-8439f3dc3b8b","resolution":{"observed_at":"2026-08-10T22:04:45.127222Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.423658Z","title":"Representation of central and peripheral vision in the primate cerebral cortex: Insights from studies of the marmoset brain","venue":null,"work_id":"bc8d5ab5-3404-476f-abd0-b4b7018b34e3","year":2015},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.132116Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:f0d2b61b14b2e17d300688432835ded6071321a498458abc0a7283dd15d547d1","observation_id":"e5041db1-5551-46b6-8a26-9d1d81b949c4","resolution":{"observed_at":"2026-08-10T22:04:45.428771Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:04:45.408521Z","title":"Places: A 10 million image database for scene recognition","venue":null,"work_id":"03403ee1-081d-445a-8888-1ac4d941f006","year":2017},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.136641Z"},"links":{"citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:c2c7fc4b92b3d39abc392764e075104218ba63b006b4f812e7cb777b552fc5e1","observation_id":"787012ed-2983-42a3-9dc4-714c5d6d9293","resolution":{"observed_at":"2026-08-10T22:04:45.413324Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2203.14415","last_updated":"2022-03-27T23:42:05Z","snapshot_observed_at":"2026-08-13T16:14:30.894916Z","submitted_at":"2022-03-27T23:42:05Z","title":"Mugs: A Multi-Granular Self-Supervised Learning Framework","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2203.14415","snapshot_observed_at":"2026-08-10T22:04:45.141516Z","title":"Mugs: A multi- granular self-supervised learning framework","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-10T22:04:45.141516Z"},"links":{"cited_paper":"/paper/2203.14415","citing_paper":"/paper/2501.02966"},"observation_digest":"sha256:928ba85d5f20994133dca68a92f43443cb2f4047e681a9b4af19c5423433a414","observation_id":"96b9502f-7659-4deb-90a6-3c146f3a19bc","resolution":{"observed_at":"2026-08-10T22:04:45.141516Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2501.02966","last_updated":"2025-01-06T12:21:40Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-11T17:24:36.682889Z","submitted_at":"2025-01-06T12:21:40Z","title":"Human Gaze Boosts Object-Centered Representation Learning"},"reference_resolution":{"displayed":66,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":17,"verified_exact":5,"verified_fuzzy":44},"total_outbound_references":66},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"thesis":"As of 15 August 2026, this Paper Citation Record lists 66 of 66 outbound references and 0 inbound Pith citation observations for arXiv:2501.02966."}