{"as_of":"2026-08-07T16:46:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:f0eb49b0be381346ff02971532694e990ea78442a80654d68b3779b35d343623","coverage":[{"denominator":27,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":27,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T12:48:09.853658Z","state":"measured"},{"denominator":30,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":30,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":3,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":3,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-02T17:05:10.690456Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-13T23:08:25.106938Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.21476","snapshot_observed_at":"2026-07-14T20:22:12.729190Z","title":"Malkin, N., Jain, M., Bengio, E., Sun, C., and Bengio, Y","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2603.18375","last_updated":"2026-07-10T23:41:27Z","snapshot_observed_at":"2026-08-04T02:39:29.405773Z","submitted_at":"2026-03-19T00:23:57Z","title":"Relationship-Centered Care: Relatedness and Responsible Design for Human Connections in Mental-Health Care","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-07-14T20:22:12.729190Z"},"links":{"cited_paper":"/paper/2507.21476","citing_paper":"/paper/2603.18375"},"observation_digest":"sha256:73d729173473a13404518155d8953ff5051ef4213e0c709157bb3530cc839550","observation_id":"7da47090-9728-4baa-8a07-054361c8b7ce","resolution":{"observed_at":"2026-07-14T20:22:12.729190Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"cited_work":{"arxiv_id":"2507.21476","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.21476","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Which llms get the joke? probing non-stem reasoning abilities with humorbench","venue":null,"work_id":"29dda195-5572-4dec-9b0e-66e41f2bca51","year":2025},"citing_paper":{"arxiv_id":"2604.19786","last_updated":"2026-07-29T21:26:11Z","snapshot_observed_at":"2026-08-02T23:16:53.127257Z","submitted_at":"2026-03-31T18:54:15Z","title":"HumorRank: A Tournament-Based Leaderboard for Evaluating Humor Generation in Large Language Models","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-05-13T23:07:38.352198Z"},"links":{"cited_paper":"/paper/2507.21476","citing_paper":"/paper/2604.19786"},"observation_digest":"sha256:c8872c13744e69e17d8a062856a16ce986b13c5741c2fc71475ff57d0606c60a","observation_id":"0dc073df-7454-4b5b-beb0-9f15dbc037eb","resolution":{"observed_at":"2026-05-13T23:08:25.110234Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.21476","snapshot_observed_at":"2026-08-02T17:05:10.690456Z","title":"Which llms get the joke? probing non-stem reasoning abilities with humorbench","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.19786","last_updated":"2026-07-29T21:26:11Z","snapshot_observed_at":"2026-08-02T23:16:53.127257Z","submitted_at":"2026-03-31T18:54:15Z","title":"HumorRank: A Tournament-Based Leaderboard for Evaluating Humor Generation in Large Language Models","version":3},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-02T17:05:10.690456Z"},"links":{"cited_paper":"/paper/2507.21476","citing_paper":"/paper/2604.19786"},"observation_digest":"sha256:5b9bebec4ad9bc9ee7bb9a509952221fcfb8ed342b332be733701ad54bf2b0b4","observation_id":"6af44ac6-42c1-4c0b-87bb-15bc9f946ce6","resolution":{"observed_at":"2026-08-02T17:05:10.690456Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2507.21476/citation-record","integrity":"/paper/2507.21476/integrity","json":"/paper/2507.21476/citation-record.json","paper":"/paper/2507.21476"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2504.21318","last_updated":"2025-04-30T05:05:09Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-04-30T05:05:09Z","title":"Phi-4-reasoning Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.21318","snapshot_observed_at":"2026-08-06T12:48:07.539696Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:07.539696Z"},"links":{"cited_paper":"/paper/2504.21318","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:c4d68c4b079c86d7fb2785d1cfb07468c3bdf5ad264a72630e2c44f874849d3b","observation_id":"a953727c-fb1c-4391-b553-3239a5d0b3bf","resolution":{"observed_at":"2026-08-06T12:48:07.539696Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:48:12.200603Z","title":"https: //blog.google/technology/google-deepmind/ gemini-model-thinking-updates-march-2025/","venue":null,"work_id":"a4e88b40-07a7-4e87-a0f8-a981fcf0de48","year":2025},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:07.793387Z"},"links":{"citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:634e7af06eb82b59643c8ff099268bc016504ff7d65265146d5fc36068810475","observation_id":"12470ab8-9ba2-4b65-8aff-82ea03bf0790","resolution":{"observed_at":"2026-08-06T12:48:12.204493Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:48:11.972592Z","title":"In Working Notes of CLEF 2025 - Conference and Labs of the Evaluation F orum","venue":null,"work_id":"34d8da6f-92b1-4039-b81e-731220d35f53","year":2025},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:07.887911Z"},"links":{"citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:c65af65cf4288ebfdd61d1c88110717c7b7d19039b020a8bf0025f0548ce288e","observation_id":"683f9eb8-8304-4b22-9cc8-ac4ddc7001b2","resolution":{"observed_at":"2026-08-06T12:48:12.130059Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.19187","last_updated":"2025-05-06T15:11:32Z","snapshot_observed_at":"2026-08-01T06:17:18.773056Z","submitted_at":"2025-02-26T14:50:50Z","title":"BIG-Bench Extra Hard","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.19187","snapshot_observed_at":"2026-08-06T12:48:07.971609Z","title":"arXiv preprint arXiv:2502.19187","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:07.971609Z"},"links":{"cited_paper":"/paper/2502.19187","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:adaa39bc56c4064fc6c243046f19e4fca1c18a138b604db32d33bcb56b8b1420","observation_id":"62bbde9c-ac16-4bc1-92d8-1c737c8671eb","resolution":{"observed_at":"2026-08-06T12:48:07.971609Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.23137","last_updated":"2026-04-15T02:38:45Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-03-29T16:08:51Z","title":"When 'YES' Meets 'BUT': Can Large Models Comprehend Contradictory Humor Through Comparative Reasoning?","version":2},"cited_work":{"arxiv_id":"2503.23137","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.23137","snapshot_observed_at":"2026-08-06T12:48:11.028577Z","title":"When 'YES' Meets 'BUT': Can Large Models Comprehend Contradictory Humor Through Comparative Reasoning?","venue":"cs.CV","work_id":"d7dd18c4-b11e-4220-96d8-31e3866bc0b4","year":2025},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.109209Z"},"links":{"cited_paper":"/paper/2503.23137","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:cdfe5428baf14ae0dde7ac05581b5886feec29f6454502afcc2ea8f6fcdc9b7b","observation_id":"80ac73db-40d3-45a8-86ed-31cbec686418","resolution":{"observed_at":"2026-08-06T12:48:11.099671Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.20050","last_updated":"2023-05-31T17:24:00Z","snapshot_observed_at":"2026-08-05T13:11:04.104454Z","submitted_at":"2023-05-31T17:24:00Z","title":"Let's Verify Step by Step","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.20050","snapshot_observed_at":"2026-08-06T12:48:08.158360Z","title":"arXiv preprint arXiv:2305.20050","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.158360Z"},"links":{"cited_paper":"/paper/2305.20050","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:b99e000e9de3233b8419fc89771db1bbf3d6e54fc487ac625b937eb9ea4cff72","observation_id":"f28ad304-3e3b-4fcb-a08e-165d092eff2d","resolution":{"observed_at":"2026-08-06T12:48:08.158360Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:48:11.842208Z","title":"In Proceed- ings of the 2023 Conference on Empirical Methods in Natural Language Processing, pages 11069–11081","venue":null,"work_id":"ea9fa2ed-29c1-4c33-a82b-3a22502f0cd8","year":2023},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.199060Z"},"links":{"citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:9ac0f6928fd2776e662dc3e1f7a8fd0b9f6f77cd0819f5a194623e8885fc8f8c","observation_id":"7cc49c60-75f7-49d8-b82f-f6234cad64d4","resolution":{"observed_at":"2026-08-06T12:48:11.956878Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.09583","last_updated":"2025-06-04T08:58:56Z","snapshot_observed_at":"2026-08-02T06:48:43.121988Z","submitted_at":"2023-08-18T14:23:21Z","title":"WizardMath: Empowering Mathematical Reasoning for Large Language Models via Reinforced Evol-Instruct","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.09583","snapshot_observed_at":"2026-08-06T12:48:08.248926Z","title":"arXiv preprint arXiv:2308.09583","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.248926Z"},"links":{"cited_paper":"/paper/2308.09583","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:3edf81e905245b3e18957c6e03ac24e579cf863a69a5ef16bd1af013298956f6","observation_id":"ad9be2f5-65c2-4393-bae4-9f82ba291152","resolution":{"observed_at":"2026-08-06T12:48:08.248926Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.09479","last_updated":"2024-05-13T01:25:12Z","snapshot_observed_at":"2026-08-06T14:28:32.497451Z","submitted_at":"2023-06-15T20:11:23Z","title":"Inverse Scaling: When Bigger Isn't Better","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.09479","snapshot_observed_at":"2026-08-06T12:48:08.300856Z","title":"arXiv preprint arXiv:2306.09479","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.300856Z"},"links":{"cited_paper":"/paper/2306.09479","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:4750684b1d92d12eb97eac97ffe5d7d530259419e693a1bfeca5cf45e787acf8","observation_id":"b084aec8-6fee-4087-a879-fb1e266d68de","resolution":{"observed_at":"2026-08-06T12:48:08.300856Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.01257","last_updated":"2025-01-03T16:36:12Z","snapshot_observed_at":"2026-07-31T19:08:41.925451Z","submitted_at":"2025-01-02T13:49:00Z","title":"CodeElo: Benchmarking Competition-level Code Generation of LLMs with Human-comparable Elo Ratings","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.01257","snapshot_observed_at":"2026-08-06T12:48:08.423730Z","title":"arXiv preprint arXiv:2501.01257","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.423730Z"},"links":{"cited_paper":"/paper/2501.01257","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:e435cf659d981f338aca2e0da589c1a5abd5ef0f73537b6623f9c907a8a85181","observation_id":"75acb1a9-c225-43a2-9ff5-3e318ecf04b2","resolution":{"observed_at":"2026-08-06T12:48:08.423730Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.12022","last_updated":"2023-11-20T18:57:34Z","snapshot_observed_at":"2026-08-04T22:55:15.345443Z","submitted_at":"2023-11-20T18:57:34Z","title":"GPQA: A Graduate-Level Google-Proof Q&A Benchmark","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.12022","snapshot_observed_at":"2026-08-06T12:48:08.472560Z","title":"arXiv preprint arXiv:2311.12022","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.472560Z"},"links":{"cited_paper":"/paper/2311.12022","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:9621eab38490b5c6f52bca9b9b2fcbf53e4cea1c8ca05b7870f5ee60e24fcd8b","observation_id":"1c712a67-ee2b-4547-b5fb-db7e402e4fd9","resolution":{"observed_at":"2026-08-06T12:48:08.472560Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-06T12:48:08.489572Z","title":"arXiv preprint arXiv:2402.03300","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.489572Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:63f8fde3b9f3d1f0e57833ca78d1024a95c3c1d2b88b2ec60c79727503adb2b7","observation_id":"06d8d6d1-29d0-44c1-8461-11d8fd80da06","resolution":{"observed_at":"2026-08-06T12:48:08.489572Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12077","last_updated":"2025-01-21T12:05:47Z","snapshot_observed_at":"2026-08-03T07:28:43.639236Z","submitted_at":"2025-01-21T12:05:47Z","title":"Phishing Awareness via Game-Based Learning","version":1},"cited_work":{"arxiv_id":"2501.12077","doi":null,"metadata_source":"pith","pith_arxiv_id":"2501.12077","snapshot_observed_at":"2026-08-06T12:48:10.667611Z","title":"Phishing Awareness via Game-Based Learning","venue":"cs.CR","work_id":"493aac99-b87f-4586-92be-d9e582a7d199","year":2025},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.645463Z"},"links":{"cited_paper":"/paper/2501.12077","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:ab502075bcb01cc5bfbc1e7484acdaf98b73d5fa003b3ad0eef33e243c985fae","observation_id":"c7ae82a2-ee3b-4401-854d-6be08571409b","resolution":{"observed_at":"2026-08-06T12:48:10.816118Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.00127","last_updated":"2025-04-30T18:48:06Z","snapshot_observed_at":"2026-08-07T15:57:34.744926Z","submitted_at":"2025-04-30T18:48:06Z","title":"Between Underthinking and Overthinking: An Empirical Study of Reasoning Length and correctness in LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.00127","snapshot_observed_at":"2026-08-06T12:48:08.790009Z","title":"arXiv preprint arXiv:2505.00127","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.790009Z"},"links":{"cited_paper":"/paper/2505.00127","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:748481c636e673dbeb7bda1b5281e0db7b6c48cf5d0f6e5bf7dd8649b9454145","observation_id":"f61fd422-528b-483d-81a9-585fc134065a","resolution":{"observed_at":"2026-08-06T12:48:08.790009Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.21380","last_updated":"2026-04-12T10:37:12Z","snapshot_observed_at":"2026-08-02T16:04:50.458535Z","submitted_at":"2025-03-27T11:20:17Z","title":"Challenging the Boundaries of Reasoning: An Olympiad-Level Math Benchmark for Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.21380","snapshot_observed_at":"2026-08-06T12:48:08.957774Z","title":"arXiv preprint arXiv:2503.21380","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.957774Z"},"links":{"cited_paper":"/paper/2503.21380","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:1bbe400d469be7432613b0d22bceceeacc81acbd77b162eecb19f94c5beb9458","observation_id":"a26467da-b9b2-4995-ac92-11b90ac777cb","resolution":{"observed_at":"2026-08-06T12:48:08.957774Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:48:11.549249Z","title":"In Proceedings of the 2022 Confer- ence on Empirical Methods in Natural Language Processing, pages 2866–2879","venue":null,"work_id":"77f364ed-d150-4c0b-bfd7-e012e66814ec","year":2022},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:09.095642Z"},"links":{"citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:4bcba1cf7b422f121ec9441d3adb7eaf1eb1a7d0f8c4edf9b611a85756c03d8e","observation_id":"96465f1a-7000-4334-ae1f-ad044ac3a944","resolution":{"observed_at":"2026-08-06T12:48:11.719843Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.02642","last_updated":"2024-01-05T05:26:25Z","snapshot_observed_at":"2026-08-05T08:38:24.137256Z","submitted_at":"2024-01-05T05:26:25Z","title":"Signatures of room-temperature superconductivity emerging in two-dimensional domains within the new Bi/Pb-based ceramic cuprate superconductors at ambient pressure","version":1},"cited_work":{"arxiv_id":"2401.02642","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.02642","snapshot_observed_at":"2026-08-06T12:48:10.430487Z","title":"Signatures of room-temperature superconductivity emerging in two-dimensional domains within the new Bi/Pb-based ceramic cuprate superconductors at ambient pressure","venue":"cond-mat.supr-con","work_id":"c6206d2b-c29c-410a-97ba-d09089894de5","year":2024},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:09.185449Z"},"links":{"cited_paper":"/paper/2401.02642","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:348b2e387dd3306f0282f3331f15d323bf77c82b397a709ceea8c7c3268ebdcf","observation_id":"a0c34740-672c-4533-a901-2407262fa64c","resolution":{"observed_at":"2026-08-06T12:48:10.523027Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:48:11.326899Z","title":"In Proceedings of the 2024 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Tech- nologies, pages 4213–4228","venue":null,"work_id":"5c5952c1-4ddd-46d9-b30a-25b712da0499","year":2024},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:09.311008Z"},"links":{"citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:095e21b76b3529590146d7e3132d8df95a9ac434051f90e9686d0c34b8d8c2d6","observation_id":"b61d379f-95df-47da-bca1-d3318f3f7d37","resolution":{"observed_at":"2026-08-06T12:48:11.432296Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.14273","last_updated":"2024-01-25T16:08:27Z","snapshot_observed_at":"2026-07-31T01:40:38.269489Z","submitted_at":"2024-01-25T16:08:27Z","title":"Uniformly rotating vortices for the lake equation","version":1},"cited_work":{"arxiv_id":"2401.14273","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.14273","snapshot_observed_at":"2026-08-06T12:48:10.132245Z","title":"Uniformly rotating vortices for the lake equation","venue":"math.AP","work_id":"189f5c88-e813-42f3-ba8c-b5c272cfff63","year":2024},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:09.407105Z"},"links":{"cited_paper":"/paper/2401.14273","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:effacfc9225eab0b5a5e53be70f409709ec190a043c917c22f9e14ccdaaa1407","observation_id":"656c4b5c-aa05-4ebe-8311-f2e126843037","resolution":{"observed_at":"2026-08-06T12:48:10.276133Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:48:09.550871Z","title":"arXiv preprint arXiv:2502.18080","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:09.550871Z"},"links":{"citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:3d5e2132d792c97a4b534c374731b0fe7842a07a5ac34e7e744281d75e73e443","observation_id":"9c75b147-50f5-4294-a4eb-4c281786a722","resolution":{"observed_at":"2026-08-06T12:48:09.550871Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10522","last_updated":"2024-12-18T05:21:24Z","snapshot_observed_at":"2026-08-03T17:34:17.504501Z","submitted_at":"2024-06-15T06:26:25Z","title":"Humor in AI: Massive Scale Crowd-Sourced Preferences and Benchmarks for Cartoon Captioning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10522","snapshot_observed_at":"2026-08-06T12:48:09.697591Z","title":"arXiv preprint arXiv:2406.10522","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:09.697591Z"},"links":{"cited_paper":"/paper/2406.10522","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:98ca85ef03eff65b8fbc02f74808a1c95853127680cd5885ba4ab5115487b43c","observation_id":"b202ac87-475a-49a2-a998-8a4f9e111b8b","resolution":{"observed_at":"2026-08-06T12:48:09.697591Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.20356","last_updated":"2025-02-27T18:29:09Z","snapshot_observed_at":"2026-08-06T10:25:12.513740Z","submitted_at":"2025-02-27T18:29:09Z","title":"Bridging the Creativity Understanding Gap: Small-Scale Human Alignment Enables Expert-Level Humor Ranking in LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.20356","snapshot_observed_at":"2026-08-06T12:48:09.853658Z","title":"arXiv preprint arXiv:2502.20356","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:09.853658Z"},"links":{"cited_paper":"/paper/2502.20356","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:1ee339229e690bc0d7ea0138b912e6e1f35a131670ad53ce23ecabbd1a64bd95","observation_id":"2720ae89-6683-4c2a-a731-9dbeb5c2d960","resolution":{"observed_at":"2026-08-06T12:48:09.853658Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2004.10645","last_updated":"2020-10-05T03:28:21Z","snapshot_observed_at":"2026-08-04T08:38:55.832006Z","submitted_at":"2020-04-22T15:42:13Z","title":"AmbigQA: Answering Ambiguous Open-domain Questions","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2004.10645","snapshot_observed_at":"2026-08-06T12:48:08.379481Z","title":"arXiv preprint arXiv:2004.10645","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.379481Z"},"links":{"cited_paper":"/paper/2004.10645","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:aed6fc98a2a48466806e09903fb3ef6297fdbedc74488b6ba7fb00f7d7168c00","observation_id":"66882ed1-a3c6-4913-8e54-06cd911d2f7f","resolution":{"observed_at":"2026-08-06T12:48:08.379481Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.14858","last_updated":"2022-07-01T02:15:12Z","snapshot_observed_at":"2026-08-05T15:41:22.691461Z","submitted_at":"2022-06-29T18:54:49Z","title":"Solving Quantitative Reasoning Problems with Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.14858","snapshot_observed_at":"2026-08-06T12:48:08.040390Z","title":"arXiv preprint arXiv:2206.14858","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.040390Z"},"links":{"cited_paper":"/paper/2206.14858","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:e7bd6a18572697c834514285ed8e262afc30a71799bbf6f838830454642309dc","observation_id":"5a21e266-f982-45af-b54a-c35634c4fe36","resolution":{"observed_at":"2026-08-06T12:48:08.040390Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.17452","last_updated":"2024-02-21T12:59:22Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-09-29T17:59:38Z","title":"ToRA: A Tool-Integrated Reasoning Agent for Mathematical Problem Solving","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.17452","snapshot_observed_at":"2026-08-06T12:48:07.929176Z","title":"arXiv preprint arXiv:2309.17452","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:07.929176Z"},"links":{"cited_paper":"/paper/2309.17452","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:39b699da8e93b1da96b3b407220fc8d07f62d1ddf826ef3eaf0a8ce4f15166ba","observation_id":"1c50d8ef-ea44-4d1b-a529-5caa53466861","resolution":{"observed_at":"2026-08-06T12:48:07.929176Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.04132","last_updated":"2024-03-07T01:22:38Z","snapshot_observed_at":"2026-08-02T17:55:33.750637Z","submitted_at":"2024-03-07T01:22:38Z","title":"Chatbot Arena: An Open Platform for Evaluating LLMs by Human Preference","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.04132","snapshot_observed_at":"2026-08-06T12:48:07.648384Z","title":"arXiv preprint arXiv:2403.04132","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:07.648384Z"},"links":{"cited_paper":"/paper/2403.04132","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:dd65f9f41d5989721b43b1d602320adcd23857bec7e35d8fe98368f76d89afb4","observation_id":"64403be8-ce55-4562-a051-38c6ef8c9e4b","resolution":{"observed_at":"2026-08-06T12:48:07.648384Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.04604","last_updated":"2025-01-08T05:24:50Z","snapshot_observed_at":"2026-08-04T14:31:04.378756Z","submitted_at":"2024-12-05T20:40:28Z","title":"ARC Prize 2024: Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.04604","snapshot_observed_at":"2026-08-06T12:48:07.697172Z","title":"arXiv preprint arXiv:2412.04604","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:07.697172Z"},"links":{"cited_paper":"/paper/2412.04604","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:345464821b1ab401b2fd70c79d698760ebe41ec5d23746157c8da9c64f7d7398","observation_id":"3118d44f-b9c8-4fe4-8ee5-0a339bc9a225","resolution":{"observed_at":"2026-08-06T12:48:07.697172Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-06T12:48:07.112830Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench"},"reference_resolution":{"displayed":27,"state_counts":{"malformed_identifier":0,"metadata_mismatch":4,"parse_uncertain":0,"unresolved":18,"verified_exact":0,"verified_fuzzy":5},"total_outbound_references":27},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 27 of 27 outbound references and 3 inbound Pith citation observations for arXiv:2507.21476."}