{"as_of":"2026-08-08T03:39:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:7265ca38e2985a91d082411ae76192eef414e6b26c62ba398babbaab8e314fe3","coverage":[{"denominator":55,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":55,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T15:30:03.702022Z","state":"measured"},{"denominator":56,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":56,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-25T04:52:41.981616Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-25T04:55:23.628724Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"cited_work":{"arxiv_id":"2505.17090","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.17090","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Emosign: A multimodal dataset for understanding emotions in american sign language","venue":null,"work_id":"58576d30-e39d-43d3-b8b6-75c131ddd0f7","year":2025},"citing_paper":{"arxiv_id":"2605.23328","last_updated":"2026-07-13T13:17:59Z","snapshot_observed_at":"2026-07-16T23:19:36.229434Z","submitted_at":"2026-05-22T07:44:20Z","title":"Emotion Recognition in Sign Language Conversation","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-25T04:52:41.981616Z"},"links":{"cited_paper":"/paper/2505.17090","citing_paper":"/paper/2605.23328"},"observation_digest":"sha256:946f82e416efb44d7eb68f2f9db9437e5bb8e35ebf8a2d5500ee90f076b2f6b1","observation_id":"42e2ddd3-c450-4355-bf55-9aeee5b1be02","resolution":{"observed_at":"2026-05-25T04:55:23.631848Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2505.17090/citation-record","integrity":"/paper/2505.17090/integrity","json":"/paper/2505.17090/citation-record.json","paper":"/paper/2505.17090"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2404.03413","last_updated":"2024-04-04T12:46:01Z","snapshot_observed_at":"2026-07-06T17:55:36.748672Z","submitted_at":"2024-04-04T12:46:01Z","title":"MiniGPT4-Video: Advancing Multimodal LLMs for Video Understanding with Interleaved Visual-Textual Tokens","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.03413","snapshot_observed_at":"2026-08-07T15:29:59.428434Z","title":"Minigpt4-video: Advancing multimodal llms for video understanding with interleaved visual-textual tokens.arXiv preprint arXiv:2404.03413, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T15:29:59.428434Z"},"links":{"cited_paper":"/paper/2404.03413","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:f001c2d2f026c8dceb28e82fdda867da85243a546d8b8076df06d8ab486e0421","observation_id":"d2f6c73f-2372-44a3-b160-0537fb465047","resolution":{"observed_at":"2026-08-07T15:29:59.428434Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.13923","last_updated":"2025-02-19T18:00:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-19T18:00:14Z","title":"Qwen2.5-VL Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.13923","snapshot_observed_at":"2026-08-07T15:29:59.519673Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T15:29:59.519673Z"},"links":{"cited_paper":"/paper/2502.13923","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:8eafec716b20ccfaf3ee076f9acac25549caa6c77ad801837df85a55d9aab577","observation_id":"041d0a2e-8201-4acc-b67b-8dc9c6137635","resolution":{"observed_at":"2026-08-07T15:29:59.519673Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2204.05862","last_updated":"2022-04-12T15:02:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-04-12T15:02:38Z","title":"Training a Helpful and Harmless Assistant with Reinforcement Learning from Human Feedback","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2204.05862","snapshot_observed_at":"2026-08-07T15:29:59.608050Z","title":"Training a helpful and harmless assistant with reinforcement learning from human feedback.arXiv preprint arXiv:2204.05862, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T15:29:59.608050Z"},"links":{"cited_paper":"/paper/2204.05862","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:6196a350336521d6de389c8d4271eed0735e0345d950e8a6e841d114694dc225","observation_id":"601aef98-c094-4021-9bb8-cca8376267ca","resolution":{"observed_at":"2026-08-07T15:29:59.608050Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:08.268958Z","title":"Openface: an open source facial behavior analysis toolkit","venue":null,"work_id":"67b2abb3-7fda-4eef-8626-a610d6975941","year":2016},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T15:29:59.736285Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:f9742700fc4907f814d2b98bde9200d96ac47ec7060b0d165fb24cee19aa2d2f","observation_id":"f1a67b60-2e22-4092-92fb-d589e6100502","resolution":{"observed_at":"2026-08-07T15:30:08.337388Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:08.140855Z","title":"Sign pose-based transformer for word-level sign language recognition","venue":null,"work_id":"6d2f1df6-f1a0-458b-b2e7-d030d166b015","year":2022},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T15:29:59.831909Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:2af362e52c032a9e7555093225f7475fcb880d9e517fced9202daa461cf1066c","observation_id":"6eaaf614-fb10-47af-b08e-e7bca10d4cfc","resolution":{"observed_at":"2026-08-07T15:30:08.194062Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:29:59.965311Z","title":"Iemocap: Interactive emotional dyadic motion capture database.Language resources and evaluation, 42:335–359, 2008","venue":null,"work_id":null,"year":2008},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T15:29:59.965311Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:d839982c5eccad2b22244824ddb308abac2e71a37efaf76a701168fda139d715","observation_id":"fe0fde76-831e-4ca8-b861-30fac383167d","resolution":{"observed_at":"2026-08-07T15:29:59.965311Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:07.994877Z","title":"Multimodal sentiment analysis with word-level fusion and reinforcement learning","venue":null,"work_id":"39388dc3-d658-44aa-876b-64c85e8f2031","year":2017},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:00.041054Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:2042fe1e31ac0819ae586bba08db29d0e71f18d1c1d1df4611472470e514427f","observation_id":"ccedc296-ee90-4a99-820a-9f4360b4fad4","resolution":{"observed_at":"2026-08-07T15:30:08.068564Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11161","last_updated":"2024-11-02T02:30:50Z","snapshot_observed_at":"2026-07-06T18:31:53.591578Z","submitted_at":"2024-06-17T03:01:22Z","title":"Emotion-LLaMA: Multimodal Emotion Recognition and Reasoning with Instruction Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.11161","snapshot_observed_at":"2026-08-07T15:30:00.144663Z","title":"Emotion-llama: Multimodal emotion recognition and reasoning with instruction tuning.arXiv preprint arXiv:2406.11161, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:00.144663Z"},"links":{"cited_paper":"/paper/2406.11161","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:24aea86de10b4a18fb927849bc181d8053130e61c4f9867285665d8a9ed40bb5","observation_id":"54f4f0a3-be67-4492-a5a9-ff14ac46d9dc","resolution":{"observed_at":"2026-08-07T15:30:00.144663Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.08072","last_updated":"2025-05-12T21:17:51Z","snapshot_observed_at":"2026-08-07T15:47:43.323600Z","submitted_at":"2025-05-12T21:17:51Z","title":"Perspectives on Capturing Emotional Expressiveness in Sign Language","version":1},"cited_work":{"arxiv_id":"2505.08072","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.08072","snapshot_observed_at":"2026-08-07T15:30:04.531881Z","title":"Perspectives on Capturing Emotional Expressiveness in Sign Language","venue":"cs.HC","work_id":"ed931735-e77d-4a1b-aaf6-f5fe60a17208","year":2025},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:00.237723Z"},"links":{"cited_paper":"/paper/2505.08072","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:b5ebc41389b3bf401953225a6971552cec46e369a152f0ba8c8513489561fd01","observation_id":"6c6c1fbb-6b15-458a-86e2-76a85175823d","resolution":{"observed_at":"2026-08-07T15:30:04.646878Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:07.880566Z","title":"University of California, Berkeley, 2014","venue":null,"work_id":"e4dcd663-d604-4166-80e5-02ddb7b3cb48","year":2014},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:00.307728Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:41f2ae99654061f9b0a1ba3079b2d48d6d6f25e1abd7f5378060dbd67af60f88","observation_id":"377474c0-b06a-492e-a5b5-836ea679a9ad","resolution":{"observed_at":"2026-08-07T15:30:07.935497Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:07.760078Z","title":"Self-report captures 27 distinct categories of emotion bridged by continuous gradients.Proceedings of the national academy of sciences, 114(38): E7900–E7909, 2017","venue":null,"work_id":"059bec04-d590-43ff-9131-e06cd6956eb9","year":2017},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:00.416461Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:67e0bb8a941d0bb9b4d018c6012ad94c939a04f9aa82cde3f1c4f9b9b8496e30","observation_id":"71669260-6444-4a7d-8d38-402ca8dc5281","resolution":{"observed_at":"2026-08-07T15:30:07.798937Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:07.641931Z","title":"Prosody in the comprehension of spoken language: A literature review.Language and speech, 40(2):141–201, 1997","venue":null,"work_id":"46e52669-3276-40af-9b6d-1dbdbf972557","year":1997},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:00.550900Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:92a93cb66e7974c7f135c2636d5196ddee73ae2da38da4609f256d198efee257","observation_id":"07420c47-11aa-4405-9ba1-49d2671ca540","resolution":{"observed_at":"2026-08-07T15:30:07.681886Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:00.641079Z","title":"How2sign: a large-scale multimodal dataset for continuous american sign language","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:00.641079Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:a9c1412295464fb3b59107d29e1ab7165f7ed6aebddb6962a4299caf0b364ffc","observation_id":"ff99d8de-c4d3-4962-890e-63deac64474a","resolution":{"observed_at":"2026-08-07T15:30:00.641079Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:00.705511Z","title":"An argument for basic emotions.Cognition & emotion, 6(3-4):169–200, 1992","venue":null,"work_id":null,"year":1992},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:00.705511Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:93806bfa8e8ce868cb3aae8ed959c02234e6cc657d603f7b4fc034b3d58579a3","observation_id":"dd19aadf-6a6e-490a-908f-9a216504656d","resolution":{"observed_at":"2026-08-07T15:30:00.705511Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:07.478607Z","title":"Facial expressions, emotions, and sign languages.Frontiers in psychology, 4:115, 2013","venue":null,"work_id":"6f98145e-7554-4c03-a7a2-89b48ebfe58f","year":2013},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:00.765099Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:8fbe6944cc60f679c5124007e6af8cc412a0c88c1612e5719a24c8f715282e8e","observation_id":"f33ef3c2-5a59-4ef1-9da6-78714407f394","resolution":{"observed_at":"2026-08-07T15:30:07.520265Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.10718","last_updated":"2025-04-30T02:19:25Z","snapshot_observed_at":"2026-07-06T18:15:46.130919Z","submitted_at":"2024-05-17T12:01:43Z","title":"SignLLM: Sign Language Production Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.10718","snapshot_observed_at":"2026-08-07T15:30:00.820749Z","title":"Signllm: Sign languages production large language models.arXiv preprint arXiv:2405.10718, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:00.820749Z"},"links":{"cited_paper":"/paper/2405.10718","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:d215c7345400b488537516352376a50e1cbef49e8b2ec04e69592c3b75f936d9","observation_id":"30d5a90d-052b-473e-8b97-bbcd169836ea","resolution":{"observed_at":"2026-08-07T15:30:00.820749Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:07.346911Z","title":"Perception of emotions in the hand movement quality of finnish sign language.Journal of nonverbal behavior, 28:53–64, 2004","venue":null,"work_id":"2edf7558-9ba4-4816-84e0-af570420f514","year":2004},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:00.887459Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:cca18f6f6166c57f997d76d3dcb526b259b8fd1f5eb76102e800361eb3ecf4e0","observation_id":"cc7dcd43-edfa-438b-b94f-e4b542d30b4c","resolution":{"observed_at":"2026-08-07T15:30:07.390454Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.07987","last_updated":"2024-07-25T09:33:50Z","snapshot_observed_at":"2026-08-06T11:02:06.607024Z","submitted_at":"2024-05-13T17:58:30Z","title":"The Platonic Representation Hypothesis","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.07987","snapshot_observed_at":"2026-08-07T15:30:00.958293Z","title":"The platonic representation hypothesis.arXiv preprint arXiv:2405.07987, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:00.958293Z"},"links":{"cited_paper":"/paper/2405.07987","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:b028df3767ec0a7823253e4d525d17d4ef4510f4d58f4879bcfcc26d13369231","observation_id":"97f8cf7f-6634-4325-b234-c3c55bc05a5f","resolution":{"observed_at":"2026-08-07T15:30:00.958293Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:01.039674Z","title":"Vader: A parsimonious rule-based model for sentiment analysis of social media text","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:01.039674Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:0d6018af29a7de31619e497bb7effd4e53fd535e062a3f8eb7c56c5dfb512d4f","observation_id":"8089c1ce-b1cd-4b13-a64f-e8d614cbd4d7","resolution":{"observed_at":"2026-08-07T15:30:01.039674Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:07.161820Z","title":"they’re not willing to accommodate deaf patients","venue":null,"work_id":"2495bbda-11c3-487a-aa1d-2c83634614eb","year":2022},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:01.096430Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:e800bf3eb522983bf4d0e520f3e285c55b3fba80f21cfaf0834840a832523a6d","observation_id":"01ce1b56-b59c-49af-80cf-297431ff3c2d","resolution":{"observed_at":"2026-08-07T15:30:07.224575Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:01.155901Z","title":"Dfew: A large-scale database for recognizing dynamic facial expressions in the wild","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:01.155901Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:5d964314d5429fe88d2573d66f4a99331bb654bda21fe7dc74550b1fd92a40d4","observation_id":"7beef794-9a64-4328-abac-a153142fbc0a","resolution":{"observed_at":"2026-08-07T15:30:01.155901Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1812.01053","last_updated":"2019-11-20T22:42:52Z","snapshot_observed_at":"2026-07-06T07:18:49.521593Z","submitted_at":"2018-12-03T19:41:16Z","title":"MS-ASL: A Large-Scale Data Set and Benchmark for Understanding American Sign Language","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1812.01053","snapshot_observed_at":"2026-08-07T15:30:01.229565Z","title":"Ms-asl: A large-scale data set and benchmark for understanding american sign language.arXiv preprint arXiv:1812.01053, 2018","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:01.229565Z"},"links":{"cited_paper":"/paper/1812.01053","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:5d5ef58e9d203531f1a2e281d4e15bf703abf8148775f43f34df6c84b8578bb6","observation_id":"a02217cc-b16e-4af1-a124-4d5703e17078","resolution":{"observed_at":"2026-08-07T15:30:01.229565Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:06.992458Z","title":"Context-aware emotion recognition networks","venue":null,"work_id":"d9f5ba80-6c55-4b05-9e41-a3eddc35dc88","year":2019},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:01.304545Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:ee2494df144483a81c4958e9bd52e715fe11df1f6186e53f99a9c4ee4d6f91da","observation_id":"fc13ea44-bbc9-4b3c-a45d-9f056ee320a5","resolution":{"observed_at":"2026-08-07T15:30:07.042020Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.00574","last_updated":"2025-02-04T07:57:52Z","snapshot_observed_at":"2026-07-06T18:08:13.456457Z","submitted_at":"2024-05-01T15:25:54Z","title":"EALD-MLLM: Emotion Analysis in Long-sequential and De-identity videos with Multi-modal Large Language Model","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.00574","snapshot_observed_at":"2026-08-07T15:30:01.374418Z","title":"Eald-mllm: Emotion analysis in long-sequential and de-identity videos with multi-modal large language model.arXiv preprint arXiv:2405.00574, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:01.374418Z"},"links":{"cited_paper":"/paper/2405.00574","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:4c5b79b12838e524d430077dc20c441aa0cf4e410ce9afbec1dad7404f5fcbd7","observation_id":"91d16b9c-6aa1-4f01-86d1-298dc03f1cac","resolution":{"observed_at":"2026-08-07T15:30:01.374418Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:01.439837Z","title":"Mimeqa: Towards socially-intelligent nonverbal foundation models.arXiv preprint arXiv:2502.16671, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:01.439837Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:4946cc38688e2ae3c5044a96cc5248e7b0dc652b2fde4146dc0d81fc3b52c020","observation_id":"a91e881e-6c55-4c56-8666-2eac6267dcfb","resolution":{"observed_at":"2026-08-07T15:30:01.439837Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:06.864812Z","title":"Mer 2023: Multi-label learning, modality robustness, and semi- supervised learning","venue":null,"work_id":"9ecba13d-a379-4c91-8fda-dcaf78677e2a","year":2023},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:01.511547Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:58c36e75888eec2d4aa30ab9c6b1b96cb272d1c0dd35c055f2a9b451dded8f33","observation_id":"cf3b454f-89da-4970-bcf6-77b1ade2461b","resolution":{"observed_at":"2026-08-07T15:30:06.932678Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.16566","last_updated":"2025-05-07T13:20:08Z","snapshot_observed_at":"2026-07-06T20:27:04.664873Z","submitted_at":"2025-01-27T23:18:39Z","title":"AffectGPT: A New Dataset, Model, and Benchmark for Emotion Understanding with Multimodal Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.16566","snapshot_observed_at":"2026-08-07T15:30:01.599045Z","title":"Affectgpt: A new dataset, model, and benchmark for emotion understanding with multimodal large language models.arXiv preprint arXiv:2501.16566, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:01.599045Z"},"links":{"cited_paper":"/paper/2501.16566","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:e6a8a2fde6c8a354546d26344cd0d0bff3f3ad1b471f584c1b0c0df1439ff405","observation_id":"2709bacd-1e47-4e65-b0df-341064a66a43","resolution":{"observed_at":"2026-08-07T15:30:01.599045Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.16524","last_updated":"2024-12-21T08:01:08Z","snapshot_observed_at":"2026-07-06T20:11:24.538172Z","submitted_at":"2024-12-21T08:01:08Z","title":"LLaVA-SLT: Visual Language Tuning for Sign Language Translation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.16524","snapshot_observed_at":"2026-08-07T15:30:01.669077Z","title":"Llava-slt: Visual language tuning for sign language translation.arXiv preprint arXiv:2412.16524, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:01.669077Z"},"links":{"cited_paper":"/paper/2412.16524","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:a16ca7c9a4f1f657853e7f9f1b3d6ed1419d78a1b7b5cc3e4ff852d3a35ec330","observation_id":"14810b1d-e5e9-46e4-b548-82fe42ff0b8e","resolution":{"observed_at":"2026-08-07T15:30:01.669077Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.03418","last_updated":"2024-07-03T18:00:48Z","snapshot_observed_at":"2026-08-03T01:16:48.562339Z","submitted_at":"2024-07-03T18:00:48Z","title":"HEMM: Holistic Evaluation of Multimodal Foundation Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.03418","snapshot_observed_at":"2026-08-07T15:30:01.743465Z","title":"Hemm: Holistic evaluation of multimodal foundation models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:01.743465Z"},"links":{"cited_paper":"/paper/2407.03418","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:550ed60a9391b841a52c138fd543070c29d0b326a98cc77b13b288f400dde45b","observation_id":"c7274a0d-2748-43f9-8c94-f50fc60536af","resolution":{"observed_at":"2026-08-07T15:30:01.743465Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.05738","last_updated":"2024-12-07T20:18:46Z","snapshot_observed_at":"2026-08-06T08:18:35.672785Z","submitted_at":"2024-12-07T20:18:46Z","title":"Exploring the Impact of Emotional Voice Integration in Sign-to-Speech Translators for Deaf-to-Hearing Communication","version":1},"cited_work":{"arxiv_id":"2412.05738","doi":null,"metadata_source":"pith","pith_arxiv_id":"2412.05738","snapshot_observed_at":"2026-08-07T15:30:04.103942Z","title":"Exploring the Impact of Emotional Voice Integration in Sign-to-Speech Translators for Deaf-to-Hearing Communication","venue":"cs.HC","work_id":"70a7ea27-6062-437c-8e6b-502d752694be","year":2024},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:01.822069Z"},"links":{"cited_paper":"/paper/2412.05738","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:a09add2a63ae24af36729475499b3d8f4569081f7ce811ccc2efccf8a21c9649","observation_id":"9724b48a-a784-4ca8-8cad-5c3bb7dc2ffe","resolution":{"observed_at":"2026-08-07T15:30:04.170294Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03744","last_updated":"2024-05-15T19:22:44Z","snapshot_observed_at":"2026-07-06T16:28:22.350574Z","submitted_at":"2023-10-05T17:59:56Z","title":"Improved Baselines with Visual Instruction Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03744","snapshot_observed_at":"2026-08-07T15:30:01.881209Z","title":"Improved baselines with visual instruction tuning, 2024.URL https://arxiv","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:01.881209Z"},"links":{"cited_paper":"/paper/2310.03744","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:06003c429d1b25dbe6ff1e04b5330fe91291cf2bddc4eba47de5dbd0132bdaeb","observation_id":"fb466e14-4687-4d61-910a-667e4b16b3df","resolution":{"observed_at":"2026-08-07T15:30:01.881209Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:06.722659Z","title":"Mafw: A large-scale, multi-modal, compound affective database for dynamic facial expression recognition in the wild","venue":null,"work_id":"1681ddbb-e87d-4b86-8e3a-95fbdf81588f","year":2022},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:01.978099Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:da5597f5512da7fe7b72eb634cc513ba4371cdee508ffdf997667f22abec452e","observation_id":"fd826690-b581-4ec8-b550-f3b7a22e6ccb","resolution":{"observed_at":"2026-08-07T15:30:06.771715Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15185","last_updated":"2023-12-23T07:46:55Z","snapshot_observed_at":"2026-08-08T00:25:44.100410Z","submitted_at":"2023-12-23T07:46:55Z","title":"emotion2vec: Self-Supervised Pre-Training for Speech Emotion Representation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.15185","snapshot_observed_at":"2026-08-07T15:30:02.058069Z","title":"emotion2vec: Self-supervised pre-training for speech emotion representation.arXiv preprint arXiv:2312.15185, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:02.058069Z"},"links":{"cited_paper":"/paper/2312.15185","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:ffe1870baa31b717cc4fd7d6e9a31ded83058ff515e293a60d8a17d43c63bfd5","observation_id":"fa03a2b8-92fd-4ae8-afab-9bee14f50ef5","resolution":{"observed_at":"2026-08-07T15:30:02.058069Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:02.120622Z","title":"Morevqa: Exploring modular reasoning models for video question answering","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:02.120622Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:480e10e31f6b6627e49539604596cdd2fb84bf0f574e23fe51e704df71cf71af","observation_id":"c80be57d-f1e7-4c28-be73-de26a3900eed","resolution":{"observed_at":"2026-08-07T15:30:02.120622Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:06.576548Z","title":"New shared & interconnected asl resources: Signstream® 3 software; dai 2 for web access to linguistically annotated video corpora; and a sign bank","venue":null,"work_id":"78a9e3e1-85c0-48c4-a660-2e80e2cb81e3","year":2018},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:02.200799Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:1dd019cab46f2dd6c8923c1f0f84ae6c25eccd8114aae8b3aafa9bef91b0eb5f","observation_id":"d442b20e-5101-4606-a493-6bc554c89cc1","resolution":{"observed_at":"2026-08-07T15:30:06.629685Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2201.07899","last_updated":"2022-01-19T22:48:36Z","snapshot_observed_at":"2026-07-06T12:28:57.325286Z","submitted_at":"2022-01-19T22:48:36Z","title":"ASL Video Corpora & Sign Bank: Resources Available through the American Sign Language Linguistic Research Project (ASLLRP)","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2201.07899","snapshot_observed_at":"2026-08-07T15:30:02.267537Z","title":"Asl video corpora & sign bank: Resources available through the american sign language linguistic research project (asllrp)","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:02.267537Z"},"links":{"cited_paper":"/paper/2201.07899","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:675e870404a2169006633d2d0c3b8538a7a312d0088916e9f8adde80c1e30963","observation_id":"bc3d9048-db66-420f-81a3-f5c205e2f08a","resolution":{"observed_at":"2026-08-07T15:30:02.267537Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:06.443946Z","title":"De Gruyter Mouton, 2012","venue":null,"work_id":"0256475e-a085-4509-96c6-3d436f4fc927","year":2012},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:02.338980Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:768eff2872bfa431061dcb62a142daa143048281396e4d976e6c3e22ade9684e","observation_id":"b328c3f2-233d-4d5a-9ed2-5207ac5577d2","resolution":{"observed_at":"2026-08-07T15:30:06.505195Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1810.02508","last_updated":"2019-06-04T12:33:49Z","snapshot_observed_at":"2026-07-06T07:06:12.726165Z","submitted_at":"2018-10-05T03:50:24Z","title":"MELD: A Multimodal Multi-Party Dataset for Emotion Recognition in Conversations","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1810.02508","snapshot_observed_at":"2026-08-07T15:30:02.422904Z","title":"Meld: A multimodal multi-party dataset for emotion recognition in conversa- tions.arXiv preprint arXiv:1810.02508, 2018","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:02.422904Z"},"links":{"cited_paper":"/paper/1810.02508","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:3eef7c25b208808837064a300ed183636957baa237c9ff224ca4f8748475a23a","observation_id":"18a472da-0eb3-43ad-a503-d4395e809cec","resolution":{"observed_at":"2026-08-07T15:30:02.422904Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:06.285720Z","title":"Affective prosody in american sign language.Sign Language Studies, 75(1):113–128, 1992","venue":null,"work_id":"2dff36d1-8204-44c5-a7bb-e4342eeac3b4","year":1992},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:02.465923Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:6c436b003a8084823c3ac3e9dd2e9e1c4e1450c9ea60e5be751e6941bfbeb908","observation_id":"32925ab3-63df-4710-8751-cb8c804e811c","resolution":{"observed_at":"2026-08-07T15:30:06.350857Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:02.550950Z","title":"A circumplex model of affect.Journal of personality and social psychology, 39(6):1161, 1980","venue":null,"work_id":null,"year":1980},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:02.550950Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:78f6d003f865b83af5a200a489a0d3808f30c97076e3efa1c454ce4d16965c82","observation_id":"b0ec4c9a-b9bf-45ed-8a48-86caae2f8873","resolution":{"observed_at":"2026-08-07T15:30:02.550950Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12870","last_updated":"2022-11-19T16:06:02Z","snapshot_observed_at":"2026-07-06T13:13:55.421435Z","submitted_at":"2022-05-25T15:43:31Z","title":"Open-Domain Sign Language Translation Learned from Online Video","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12870","snapshot_observed_at":"2026-08-07T15:30:02.619890Z","title":"Open-domain sign language translation learned from online video.arXiv preprint arXiv:2205.12870, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:02.619890Z"},"links":{"cited_paper":"/paper/2205.12870","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:0b7cce0d87a244287c43a3e7708745d393a036dc8c8e9f58639e2cb229cf2de8","observation_id":"38282ae4-b796-4fe1-92e9-f058be1d36c0","resolution":{"observed_at":"2026-08-07T15:30:02.619890Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:06.141556Z","title":null,"venue":null,"work_id":"47cc12db-5da7-43eb-a21e-1c99b2f0d35a","year":2019},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:02.709111Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:d860ef10c0b4fe4d3ab5396188da63cb55d09da87ae196f8e6a8d21c0c7d6fc5","observation_id":"91aa7793-6297-4183-b10c-f07d6a0c48ca","resolution":{"observed_at":"2026-08-07T15:30:06.212504Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:05.991437Z","title":"Roberta-lstm: a hybrid model for sentiment analysis with transformer and recurrent neural network.IEEE Access, 10:21517–21525, 2022","venue":null,"work_id":"5abd57cb-bf92-482c-a5d4-a7cd927fe6c8","year":2022},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:02.783706Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:1b1dcd8d080884c15f7b189242450335e2589a7bfc54cf386e22d117f337a604","observation_id":"d2930ab4-1496-42c8-b0f3-345032b283f7","resolution":{"observed_at":"2026-08-07T15:30:06.075855Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:02.820602Z","title":"Youtube-asl: A large-scale, open-domain american sign language-english parallel corpus.Advances in Neural Information Processing Systems, 36:29029–29047, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:02.820602Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:66c5fc6ed2d74796b335e11dbb4558f6c8da4defa80bac417718237258c43f5f","observation_id":"03a1d710-1288-4ae3-b174-7e61b536bcb6","resolution":{"observed_at":"2026-08-07T15:30:02.820602Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:05.822684Z","title":"Gallaudet University Press, 2000","venue":null,"work_id":"1ec1c7de-1724-445e-b29e-484fb7922379","year":2000},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:02.908405Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:2341822eccbd593e895c64e68d55536dc6a5040db428c7934b8208900dccec6f","observation_id":"8facebb7-cef4-4275-95ae-d1d0a269ceaf","resolution":{"observed_at":"2026-08-07T15:30:05.876738Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:05.695497Z","title":"The future of emotion in human- computer interaction","venue":null,"work_id":"741b1cab-7595-4749-81c1-331b8122c50f","year":2022},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:02.983924Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:9dad9473d6b6beb61359ee0ef0542d202ab829ee5c0501404e23cde86a392a61","observation_id":"c0f19a77-92f7-41b4-80c2-50dedd5c7b13","resolution":{"observed_at":"2026-08-07T15:30:05.741736Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:03.052507Z","title":"Can i trust your answer? visually grounded video question answering","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:03.052507Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:d5d101ef891d3ecf6b031b641258d42466d68855eff621a97a1695d843fb97c7","observation_id":"b1fa4f4d-26ae-4695-95eb-3caecd5dd4f8","resolution":{"observed_at":"2026-08-07T15:30:03.052507Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:05.524986Z","title":"Detecting depression severity from vocal prosody.IEEE transactions on affective computing, 4(2):142–150, 2012","venue":null,"work_id":"acf17dc3-9128-481f-a0bf-00f3c3205316","year":2012},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:03.138996Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:2d90691d4a4c72b5191d7b527d7a8b423953d407882a9cfdb39fddbc8af41ea6","observation_id":"42b5fcdc-d589-4e96-8077-ef34bd93caab","resolution":{"observed_at":"2026-08-07T15:30:05.617887Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.05783","last_updated":"2024-11-08T18:50:37Z","snapshot_observed_at":"2026-08-06T18:49:38.482239Z","submitted_at":"2024-11-08T18:50:37Z","title":"ASL STEM Wiki: Dataset and Benchmark for Interpreting STEM Articles","version":1},"cited_work":{"arxiv_id":"2411.05783","doi":null,"metadata_source":"pith","pith_arxiv_id":"2411.05783","snapshot_observed_at":"2026-08-07T15:30:03.843088Z","title":"ASL STEM Wiki: Dataset and Benchmark for Interpreting STEM Articles","venue":"cs.CL","work_id":"b6b9e0a1-baea-4dfa-a936-b4e1c3491897","year":2024},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:03.220688Z"},"links":{"cited_paper":"/paper/2411.05783","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:7579299cf1c3ac93f43fcd51a6c7df20a69c3066a055c7df1976d3deaa5c8321","observation_id":"124c3081-0d87-4db5-96c9-a67dd5d6269d","resolution":{"observed_at":"2026-08-07T15:30:03.898781Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1707.07250","last_updated":"2017-07-23T05:54:20Z","snapshot_observed_at":"2026-07-06T05:52:17.418793Z","submitted_at":"2017-07-23T05:54:20Z","title":"Tensor Fusion Network for Multimodal Sentiment Analysis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.07250","snapshot_observed_at":"2026-08-07T15:30:03.301698Z","title":"Tensor fusion network for multimodal sentiment analysis.arXiv preprint arXiv:1707.07250, 2017","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:03.301698Z"},"links":{"cited_paper":"/paper/1707.07250","citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:17ef7471e9974d50194a9146beae2b2c4aec98028b0a6a6a5bdda4e35f0577dd","observation_id":"8eb1f1a7-95f8-4b1d-859e-18e2aa976da1","resolution":{"observed_at":"2026-08-07T15:30:03.301698Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:05.371358Z","title":"Multimodal language analysis in the wild: Cmu-mosei dataset and interpretable dynamic fusion graph","venue":null,"work_id":"5cb6a6b8-f1fa-4482-a574-0aad418da510","year":2018},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:03.378293Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:48c693b2ac66a578cca81b745f3f6cddbe9a4f5e41e0cc9333df541f44a93757","observation_id":"e1d5c81f-7cff-4838-84c3-73b809968df6","resolution":{"observed_at":"2026-08-07T15:30:05.437269Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:05.207999Z","title":"filename","venue":null,"work_id":"3139a90a-d44a-48e4-b0a1-08ef1e1218ca","year":null},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:03.457512Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:58ce9da7c7385e68eaac0d63a70a5cda42e2096b9c686885aeac7ef2659076d8","observation_id":"25ba6666-8621-48dc-8f36-80a24780120d","resolution":{"observed_at":"2026-08-07T15:30:05.274566Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:05.091256Z","title":null,"venue":null,"work_id":"e23ff47b-2b3d-43f0-a005-4b7cd92bd8bf","year":null},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:03.532471Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:31bac3874a6766d6a744a1abba2d8e4934631fd59fc685c2ba0fa4fb5b109d04","observation_id":"904be561-1b41-45a6-9b2d-b180f7581823","resolution":{"observed_at":"2026-08-07T15:30:05.141605Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:04.956013Z","title":null,"venue":null,"work_id":"544c9232-2bbd-477b-949a-b074b52a6618","year":null},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:03.625333Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:1aaa9d799717aa038ca6fbdc480db124d19637b11c9fd79d5db87af988430561","observation_id":"eabe1079-bb44-4f81-9475-9eaf47f832f8","resolution":{"observed_at":"2026-08-07T15:30:05.006225Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:30:04.805026Z","title":"If Mary gets home late, John will probably be upset","venue":null,"work_id":"1e315543-56a9-48b1-8498-2f84ef5933e9","year":null},"citing_paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:03.702022Z"},"links":{"citing_paper":"/paper/2505.17090"},"observation_digest":"sha256:c1291d5b8f4a7f0a0c6850818229ec05b5fdcc5b7c1411d9b729bf54f4715917","observation_id":"f168e2cb-7a86-466d-b70b-d6a7a168413d","resolution":{"observed_at":"2026-08-07T15:30:04.873498Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2505.17090","last_updated":"2025-05-20T22:14:38Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-08T00:26:41.300341Z","submitted_at":"2025-05-20T22:14:38Z","title":"EmoSign: A Multimodal Dataset for Understanding Emotions in American Sign Language"},"reference_resolution":{"displayed":55,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":30,"verified_exact":3,"verified_fuzzy":21},"total_outbound_references":55},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 55 of 55 outbound references and 1 inbound Pith citation observation for arXiv:2505.17090."}