{"as_of":"2026-08-05T20:45:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:5d37c096c8d7d59c867bd46f8edb65b597513c5e5d044f6e38ee113fc404583f","coverage":[{"denominator":30,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":30,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-29T17:58:58.353336Z","state":"measured"},{"denominator":30,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":30,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2605.27737/citation-record","integrity":"/paper/2605.27737/integrity","json":"/paper/2605.27737/citation-record.json","paper":"/paper/2605.27737"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Flamingo: a visual language model for few-shot learning","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:a97f5fc1a28522abb520ed2664cad97b4d3d326bd6da9b82b2c0733682a9f2e6","observation_id":"954bfb8d-5d20-430f-82c5-237ba8409e95","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Qwen2 technical report.arXiv (Cornell University), 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:6edba0b200d25281de4870a074218d749d103c1c297d8c05c2bc546a21001ad0","observation_id":"a9b49539-1be0-4f3e-bd93-490228c14ecd","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Qwen-vl: A versatile vision-language model for un- derstanding, localization, text reading, and beyond, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:6ecc6ce5a9fbb0e048f9d2b1941d2fc9af6d24fbe4618e464bd1956c2e5dc135","observation_id":"b2b2b2a2-04c8-4354-9a6e-d9a9aa96b987","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.13923","last_updated":"2025-02-19T18:00:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-19T18:00:14Z","title":"Qwen2.5-VL Technical Report","version":1},"cited_work":{"arxiv_id":"2502.13923","doi":"10.48550/arxiv.2502.13923","metadata_source":"pith","pith_arxiv_id":"2502.13923","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Qwen2.5-VL Technical Report","venue":"cs.CV","work_id":"69dffacb-bfe8-442d-be86-48624c60426f","year":2025},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"cited_paper":"/paper/2502.13923","citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:17904256f848754fe7bad67ef3817a23010b762adb1460a1052d4fc42a9a8d0e","observation_id":"7b4afcd5-708f-4bde-aee7-8a2735c2a104","resolution":{"observed_at":"2026-06-29T18:03:48.006747Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:19:13.082554+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:19:13.082554+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Mobilevlm : A fast, strong and open vision language assistant for mobile devices.arXiv (Cornell University), 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:aad1106ce70ad3cdc4ffa45b12ddb71d11622e48bad992658693a67e71f543f7","observation_id":"d0cf9473-c173-4fe1-8ece-03fbf1a58786","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Mobilevlm v2: Faster and stronger baseline for vision language model.arXiv (Cornell University), 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:ba71e3a322ccf15088e9f6aa13575fef17d2ae610b0b220af039fc2b34be6abb","observation_id":"bfa0a8ba-0e56-4d11-8c1b-e98111859d84","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Lovif @ cvpr 2026: Challenge on efficient vlm for multimodal creative qual- ity scoring.https : / / www","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:33e98d5bd2240480200dc67f0dd04faa4dae29e2e176e74114e956750c1bfa8c","observation_id":"7ebb378b-0292-4203-81ab-fa97758028ea","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.08691","last_updated":"2023-07-17T17:50:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-07-17T17:50:36Z","title":"FlashAttention-2: Faster Attention with Better Parallelism and Work Partitioning","version":1},"cited_work":{"arxiv_id":"2307.08691","doi":"10.48550/arxiv.2307.08691","metadata_source":"pith","pith_arxiv_id":"2307.08691","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"FlashAttention-2: Faster Attention with Better Parallelism and Work Partitioning","venue":"cs.LG","work_id":"fff3953b-5efb-4753-bee4-002f59995810","year":2023},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"cited_paper":"/paper/2307.08691","citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:80dc99f385fd425f962416cc6ecebc01aff3271a3e76ab2c70e121764cb85e44","observation_id":"bbe114f4-f3d4-436e-a264-f09c6175d342","resolution":{"observed_at":"2026-06-29T18:03:48.015433Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T15:50:18.243793+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T15:50:18.243793+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2110.02861","last_updated":"2022-06-20T16:05:15Z","snapshot_observed_at":"2026-07-06T11:55:04.054344Z","submitted_at":"2021-10-06T15:43:20Z","title":"8-bit Optimizers via Block-wise Quantization","version":2},"cited_work":{"arxiv_id":"2110.02861","doi":"10.48550/arxiv.2110.02861","metadata_source":"arxiv_reference","pith_arxiv_id":"2110.02861","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Dettmers, M","venue":"arXiv (Cornell University)","work_id":"dc9d8ece-f716-493d-89a8-2a30dedcceb4","year":2021},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"cited_paper":"/paper/2110.02861","citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:45e8da11b45704e486850a775e28de72ee7fe93d8e27a9ab80a82d14d4d6ffde","observation_id":"4fcba71d-9a04-47dc-b5a5-b08b92e64a1c","resolution":{"observed_at":"2026-06-29T18:03:48.009935Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2209.05040","last_updated":"2022-10-05T04:13:47Z","snapshot_observed_at":"2026-07-06T13:51:05.005337Z","submitted_at":"2022-09-12T06:31:13Z","title":"SANCL: Multimodal Review Helpfulness Prediction with Selective Attention and Natural Contrastive Learning","version":5},"cited_work":{"arxiv_id":"2209.05040","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2209.05040","snapshot_observed_at":"2026-06-29T18:03:48.018397Z","title":"Sancl: Multimodal review helpfulness prediction with selective attention and natural contrastive learning.arXiv preprint arXiv:2209.05040, 2022","venue":null,"work_id":"6dd36d8d-478c-4fc2-9f5e-019deda7a0c0","year":2022},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"cited_paper":"/paper/2209.05040","citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:2a4877861eab0b9f9f76806eadad03d140e02c970c31419e981d21c520f97dc6","observation_id":"10077c86-1609-40ea-bf08-406aebb71df7","resolution":{"observed_at":"2026-06-29T18:03:48.020088Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Vbpr: Visual bayesian per- sonalized ranking from implicit feedback.Proceedings of the AAAI Conference on Artificial Intelligence, 2016","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:cd0b16999aca4370ca2dda7b253953587f0859589f750ba81efcebc71eae6aef","observation_id":"84184723-c0fd-49e7-a73f-e6b6b66e0729","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.03952","last_updated":"2026-04-20T06:05:54Z","snapshot_observed_at":"2026-08-03T00:56:47.399756Z","submitted_at":"2024-03-06T18:56:36Z","title":"Bridging Language and Items for Retrieval and Recommendation: Benchmarking LLMs as Semantic Encoders","version":2},"cited_work":{"arxiv_id":"2403.03952","doi":"10.1007/978-3-031-56060-6","metadata_source":"pith","pith_arxiv_id":"2403.03952","snapshot_observed_at":"2026-08-05T02:49:54.815029Z","title":"Bridging Language and Items for Retrieval and Recommendation: Benchmarking LLMs as Semantic Encoders","venue":"cs.IR","work_id":"357c2b19-af52-48bc-a55a-e34d2d84b8ff","year":2024},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"cited_paper":"/paper/2403.03952","citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:7c8f54ff40e24513c8acca585fce49e0ee97f3f0d8bf99db2fb29395c50e7626","observation_id":"57c09eaa-6e4d-4979-b6b4-f8b91cedd08e","resolution":{"observed_at":"2026-06-29T18:03:48.022658Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Aesbench: An expert benchmark for multimodal large language models on image aesthetics perception.arXiv (Cornell University), 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:245fc8347f2773fe3f47d566a4741659693a8bb450f24cc858f2fa3e2bcf31e6","observation_id":"18a5e3d6-354e-4cc4-9a60-2a7349429d0e","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Huggingfacetb/smolvlm2-256m-video- instruct model card.https : / / huggingface","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:8823e54247cf9b4f1c0931c484fbbc4dd118fc6cbc3af9fc61da6c9e6a5eb79f","observation_id":"025742fd-6fdd-4b5a-82ee-edb68e0741ff","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Mcauley-lab/amazon-reviews- 2023 dataset card.https : / / huggingface","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:9a5a4d21c8f12e30c90913b91d078241a2e5adc7f0723575eb95f2b0e29f59aa","observation_id":"5d32b651-92d1-4632-bac3-33dc07cbadde","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Building and better understanding vision- language models: insights and future directions.arXiv (Cor- nell University), 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:f08083e3ba0829affb1fc0c885601e7d9c75a2fce1f1f38bb4705c571534743e","observation_id":"d93a29b1-517e-4fdd-9b98-607927bc2c5f","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"What matters when building vision-language models? arXiv (Cornell University), 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:2bc6ef7149a357dd0dca996ba3cedb66fe0984a914cd5471e072ea43b96cc1bf","observation_id":"4c45dc12-8fa9-4112-9fde-59b3ed2025b0","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Delving deep into engagement prediction of short videos","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:e929a15b64774e93d503480d4314b9e60fa7f1937543232bce8a3b1fdcd8b065","observation_id":"e508fcc9-4790-46bf-a130-1220796ddabc","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Vquala 2025 challenge on engagement prediction for short videos: Methods and results","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:7021e34ee33359d136aa2eeaa0cae7743aa8d395fb84b9c7516e1a5d23a64cc0","observation_id":"816e0c38-b428-4bec-b381-f1c97b8ddc97","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Blip-2: Bootstrapping language-image pre-training with frozen image encoders and large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:4afdb4e4eff55f401a8b9503fb40e388960eb619a1d16734fd65eb6290a6167a","observation_id":"a03191ad-4f39-4158-8751-94a5f2cfe938","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Improved baselines with visual instruction tuning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:a8826110549ca4a249bf15e1ac67d59def0f6afa9985a52b22bd8e2e0b5f27ce","observation_id":"c947d237-d884-46cf-b49a-3fa14bdc858f","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Multi- perspective coherent reasoning for helpfulness prediction of multimodal reviews","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:64bd04f73df3c9506d45e513d97b015688f8b9decd107c4f88e3bee5dfde1daa","observation_id":"d4652587-4e2d-4e6e-8471-2d598e55d395","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Multimodal recommender systems: A survey.ACM Computing Surveys,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:4d4795213b3495a492b3a88f191ead0f293945f1ae94efb7cc8663a8dc2921da","observation_id":"e7abd74f-2170-4e66-aede-41a989e29fdb","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05299","last_updated":"2025-04-07T17:58:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-04-07T17:58:57Z","title":"SmolVLM: Redefining small and efficient multimodal models","version":1},"cited_work":{"arxiv_id":"2504.05299","doi":"10.18653/v1/2025.acl-long.718","metadata_source":"pith","pith_arxiv_id":"2504.05299","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SmolVLM: Redefining small and efficient multimodal models","venue":"cs.AI","work_id":"810cc87f-f6bd-4443-9673-5e30982426b9","year":2025},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"cited_paper":"/paper/2504.05299","citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:089edc3b7e1af962ad79faf9ed194559d6b509a538d36609999bd74fbcd5948f","observation_id":"475d82ca-453a-4d5b-99b1-bcdf107e12e2","resolution":{"observed_at":"2026-06-29T18:03:48.012477Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:22c8fdb7be2a85707d2487c2edcde54cea7849f52ae4bf75d531c995c5c9dca2","observation_id":"3edd74a3-259b-47fa-8522-68a1a3737f95","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Engagement prediction of short videos with large multimodal models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:604c99f94d186eb56e4d5e70724b2d0db1759b8850cfa84ced097306ea22c4d4","observation_id":"235cad66-f27c-482b-a54c-0fc57027e775","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Deepseek-vl2: Mixture-of-experts vision-language models for advanced multimodal understanding.arXiv (Cornell Uni- versity), 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:3f31fe712cd85cedc86a1d2dca11a935eb7a84f5b82beeccd01f7342dd62a276","observation_id":"5bc35159-5c18-4a4e-a417-b30de02e9d02","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:2fdf6b39d80a25090a7474ec84425ffb9ca36c4d03df041c8cf42d7aea192625","observation_id":"254a3e62-b3c6-44ec-a3bd-e415de2b87ee","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Minicpm-v: A gpt-4v level mllm on your phone.arXiv (Cor- nell University), 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:2f22553bd50dfba6e9fb2534c5088392a364c034b888e2ad45b6a76d22886b7f","observation_id":"5f527483-ccbd-4fa5-9c48-86e751d13962","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T17:58:58.353336Z","title":"Depicting beyond scores: Advanc- ing image quality assessment through multi-modal language models.Lecture notes in computer science, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-06-29T17:58:58.353336Z"},"links":{"citing_paper":"/paper/2605.27737"},"observation_digest":"sha256:d33ebcb4021503ce787e2f22d2c63f8482165d42c741f85c56ff6977c6b2876d","observation_id":"346d74f4-b520-4b0b-94fa-7265c178d044","resolution":{"observed_at":"2026-06-29T17:58:58.353336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2605.27737","last_updated":"2026-05-26T22:27:58Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-01T07:53:04.594255Z","submitted_at":"2026-05-26T22:27:58Z","title":"Bounded-Compute Multimodal Regression for Product-Rating Prediction"},"reference_resolution":{"displayed":30,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":24,"verified_exact":5,"verified_fuzzy":0},"total_outbound_references":30},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 30 of 30 outbound references and 0 inbound Pith citation observations for arXiv:2605.27737."}