{"as_of":"2026-08-10T07:39:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:3af335447670e2546b653c5b99fd02d466e89d3f665bd8d2b3a642823d069792","coverage":[{"denominator":27,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":27,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T12:48:09.853658Z","state":"measured"},{"denominator":30,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":30,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":3,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":3,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-02T17:05:10.690456Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-13T23:08:25.106938Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.21476","snapshot_observed_at":"2026-07-14T20:22:12.729190Z","title":"Malkin, N., Jain, M., Bengio, E., Sun, C., and Bengio, Y","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2603.18375","last_updated":"2026-07-10T23:41:27Z","snapshot_observed_at":"2026-08-07T22:46:43.471693Z","submitted_at":"2026-03-19T00:23:57Z","title":"Relationship-Centered Care: Relatedness and Responsible Design for Human Connections in Mental-Health Care","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-07-14T20:22:12.729190Z"},"links":{"cited_paper":"/paper/2507.21476","citing_paper":"/paper/2603.18375"},"observation_digest":"sha256:c3c1a560804e74d49b4150f3b486088d667fcb8844d81c41bae1f89057bb02c8","observation_id":"7da47090-9728-4baa-8a07-054361c8b7ce","resolution":{"observed_at":"2026-07-14T20:22:12.729190Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"cited_work":{"arxiv_id":"2507.21476","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.21476","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Which llms get the joke? probing non-stem reasoning abilities with humorbench","venue":null,"work_id":"29dda195-5572-4dec-9b0e-66e41f2bca51","year":2025},"citing_paper":{"arxiv_id":"2604.19786","last_updated":"2026-07-29T21:26:11Z","snapshot_observed_at":"2026-08-02T23:16:53.127257Z","submitted_at":"2026-03-31T18:54:15Z","title":"HumorRank: A Tournament-Based Leaderboard for Evaluating Humor Generation in Large Language Models","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-05-13T23:07:38.352198Z"},"links":{"cited_paper":"/paper/2507.21476","citing_paper":"/paper/2604.19786"},"observation_digest":"sha256:f0a0c300c55aab9574ffb8c9e6dbc214d04b2e2e79687d85fb3d961d607f348b","observation_id":"0dc073df-7454-4b5b-beb0-9f15dbc037eb","resolution":{"observed_at":"2026-05-13T23:08:25.110234Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.21476","snapshot_observed_at":"2026-08-02T17:05:10.690456Z","title":"Which llms get the joke? probing non-stem reasoning abilities with humorbench","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.19786","last_updated":"2026-07-29T21:26:11Z","snapshot_observed_at":"2026-08-02T23:16:53.127257Z","submitted_at":"2026-03-31T18:54:15Z","title":"HumorRank: A Tournament-Based Leaderboard for Evaluating Humor Generation in Large Language Models","version":3},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-02T17:05:10.690456Z"},"links":{"cited_paper":"/paper/2507.21476","citing_paper":"/paper/2604.19786"},"observation_digest":"sha256:55f83a960c65c039d7e48a5ba9c79b6b7a1c40ddea937936a07bef5a7cde4dae","observation_id":"6af44ac6-42c1-4c0b-87bb-15bc9f946ce6","resolution":{"observed_at":"2026-08-02T17:05:10.690456Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2507.21476/citation-record","integrity":"/paper/2507.21476/integrity","json":"/paper/2507.21476/citation-record.json","paper":"/paper/2507.21476"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2504.21318","last_updated":"2025-04-30T05:05:09Z","snapshot_observed_at":"2026-08-08T08:43:53.657438Z","submitted_at":"2025-04-30T05:05:09Z","title":"Phi-4-reasoning Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.21318","snapshot_observed_at":"2026-08-06T12:48:07.539696Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:07.539696Z"},"links":{"cited_paper":"/paper/2504.21318","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:5581fb477214c5e17f1c76e5776c793f942e4ae8ada1a909652a3e8a5e0de4a6","observation_id":"a953727c-fb1c-4391-b553-3239a5d0b3bf","resolution":{"observed_at":"2026-08-06T12:48:07.539696Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:48:12.200603Z","title":"https: //blog.google/technology/google-deepmind/ gemini-model-thinking-updates-march-2025/","venue":null,"work_id":"a4e88b40-07a7-4e87-a0f8-a981fcf0de48","year":2025},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:07.793387Z"},"links":{"citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:84159cf300b3945353403aa720f3ee1c9d9321926b004dba21b595fcb4fa0774","observation_id":"12470ab8-9ba2-4b65-8aff-82ea03bf0790","resolution":{"observed_at":"2026-08-06T12:48:12.204493Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:48:11.972592Z","title":"In Working Notes of CLEF 2025 - Conference and Labs of the Evaluation F orum","venue":null,"work_id":"34d8da6f-92b1-4039-b81e-731220d35f53","year":2025},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:07.887911Z"},"links":{"citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:0f148969b07ea716e54f2b44861f9da91ed3aa3f23d6b09c5cb7c78dc68c63f7","observation_id":"683f9eb8-8304-4b22-9cc8-ac4ddc7001b2","resolution":{"observed_at":"2026-08-06T12:48:12.130059Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.19187","last_updated":"2025-05-06T15:11:32Z","snapshot_observed_at":"2026-08-07T17:46:00.833139Z","submitted_at":"2025-02-26T14:50:50Z","title":"BIG-Bench Extra Hard","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.19187","snapshot_observed_at":"2026-08-06T12:48:07.971609Z","title":"arXiv preprint arXiv:2502.19187","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:07.971609Z"},"links":{"cited_paper":"/paper/2502.19187","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:99a577178bc8bb7d4e67f1f54c17472ea44bf1a3bf71f9a348f179c2389f672e","observation_id":"62bbde9c-ac16-4bc1-92d8-1c737c8671eb","resolution":{"observed_at":"2026-08-06T12:48:07.971609Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.23137","last_updated":"2026-04-15T02:38:45Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-03-29T16:08:51Z","title":"When 'YES' Meets 'BUT': Can Large Models Comprehend Contradictory Humor Through Comparative Reasoning?","version":2},"cited_work":{"arxiv_id":"2503.23137","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.23137","snapshot_observed_at":"2026-08-06T12:48:11.028577Z","title":"When 'YES' Meets 'BUT': Can Large Models Comprehend Contradictory Humor Through Comparative Reasoning?","venue":"cs.CV","work_id":"d7dd18c4-b11e-4220-96d8-31e3866bc0b4","year":2025},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.109209Z"},"links":{"cited_paper":"/paper/2503.23137","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:3c413d13cddf08500d21896d54386acb73c8d8e5af12e9310a911eb790400b38","observation_id":"80ac73db-40d3-45a8-86ed-31cbec686418","resolution":{"observed_at":"2026-08-06T12:48:11.099671Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.20050","last_updated":"2023-05-31T17:24:00Z","snapshot_observed_at":"2026-08-05T13:11:04.104454Z","submitted_at":"2023-05-31T17:24:00Z","title":"Let's Verify Step by Step","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.20050","snapshot_observed_at":"2026-08-06T12:48:08.158360Z","title":"arXiv preprint arXiv:2305.20050","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.158360Z"},"links":{"cited_paper":"/paper/2305.20050","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:59483aa8b3d49fb26e9ca96eaf738d43f9e31010dcce0b4b8c57220bc646864d","observation_id":"f28ad304-3e3b-4fcb-a08e-165d092eff2d","resolution":{"observed_at":"2026-08-06T12:48:08.158360Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:48:11.842208Z","title":"In Proceed- ings of the 2023 Conference on Empirical Methods in Natural Language Processing, pages 11069–11081","venue":null,"work_id":"ea9fa2ed-29c1-4c33-a82b-3a22502f0cd8","year":2023},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.199060Z"},"links":{"citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:7994f08310d0d085f9466986dbad31aabc66424cd297aa1b4f6fe7b9ac759741","observation_id":"7cc49c60-75f7-49d8-b82f-f6234cad64d4","resolution":{"observed_at":"2026-08-06T12:48:11.956878Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.09583","last_updated":"2025-06-04T08:58:56Z","snapshot_observed_at":"2026-08-02T06:48:43.121988Z","submitted_at":"2023-08-18T14:23:21Z","title":"WizardMath: Empowering Mathematical Reasoning for Large Language Models via Reinforced Evol-Instruct","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.09583","snapshot_observed_at":"2026-08-06T12:48:08.248926Z","title":"arXiv preprint arXiv:2308.09583","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.248926Z"},"links":{"cited_paper":"/paper/2308.09583","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:7a8ef00dc9c116fe9b0467b5a434349ff0dc430c684612dc45ebcb1ded5b6b34","observation_id":"ad9be2f5-65c2-4393-bae4-9f82ba291152","resolution":{"observed_at":"2026-08-06T12:48:08.248926Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.09479","last_updated":"2024-05-13T01:25:12Z","snapshot_observed_at":"2026-08-06T14:28:32.497451Z","submitted_at":"2023-06-15T20:11:23Z","title":"Inverse Scaling: When Bigger Isn't Better","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.09479","snapshot_observed_at":"2026-08-06T12:48:08.300856Z","title":"arXiv preprint arXiv:2306.09479","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.300856Z"},"links":{"cited_paper":"/paper/2306.09479","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:3e063c3b191597885ce1340dcf54c2326a140cefb91816564b8d04629b4a6615","observation_id":"b084aec8-6fee-4087-a879-fb1e266d68de","resolution":{"observed_at":"2026-08-06T12:48:08.300856Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.01257","last_updated":"2025-01-03T16:36:12Z","snapshot_observed_at":"2026-07-31T19:08:41.925451Z","submitted_at":"2025-01-02T13:49:00Z","title":"CodeElo: Benchmarking Competition-level Code Generation of LLMs with Human-comparable Elo Ratings","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.01257","snapshot_observed_at":"2026-08-06T12:48:08.423730Z","title":"arXiv preprint arXiv:2501.01257","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.423730Z"},"links":{"cited_paper":"/paper/2501.01257","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:922aeb17e264cc487ef11b843e25ee0f5f1141e4c23f68b8278214c10adc34e4","observation_id":"75acb1a9-c225-43a2-9ff5-3e318ecf04b2","resolution":{"observed_at":"2026-08-06T12:48:08.423730Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.12022","last_updated":"2023-11-20T18:57:34Z","snapshot_observed_at":"2026-08-04T22:55:15.345443Z","submitted_at":"2023-11-20T18:57:34Z","title":"GPQA: A Graduate-Level Google-Proof Q&A Benchmark","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.12022","snapshot_observed_at":"2026-08-06T12:48:08.472560Z","title":"arXiv preprint arXiv:2311.12022","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.472560Z"},"links":{"cited_paper":"/paper/2311.12022","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:4b34f0191bff697f7b6a62f4ad5f3b4738c0e119e58622dcb225d479a82547a2","observation_id":"1c712a67-ee2b-4547-b5fb-db7e402e4fd9","resolution":{"observed_at":"2026-08-06T12:48:08.472560Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-06T12:48:08.489572Z","title":"arXiv preprint arXiv:2402.03300","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.489572Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:e9131dec6a59d38375a22c43c6131f2ad69988fe1875e80f5abbc2ca4add0e59","observation_id":"06d8d6d1-29d0-44c1-8461-11d8fd80da06","resolution":{"observed_at":"2026-08-06T12:48:08.489572Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12077","last_updated":"2025-01-21T12:05:47Z","snapshot_observed_at":"2026-08-03T07:28:43.639236Z","submitted_at":"2025-01-21T12:05:47Z","title":"Phishing Awareness via Game-Based Learning","version":1},"cited_work":{"arxiv_id":"2501.12077","doi":null,"metadata_source":"pith","pith_arxiv_id":"2501.12077","snapshot_observed_at":"2026-08-06T12:48:10.667611Z","title":"Phishing Awareness via Game-Based Learning","venue":"cs.CR","work_id":"493aac99-b87f-4586-92be-d9e582a7d199","year":2025},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.645463Z"},"links":{"cited_paper":"/paper/2501.12077","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:4efb3571df22125ee42476ab94bccb1a0fe180942971cc050ba98fa41c75b078","observation_id":"c7ae82a2-ee3b-4401-854d-6be08571409b","resolution":{"observed_at":"2026-08-06T12:48:10.816118Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.00127","last_updated":"2025-04-30T18:48:06Z","snapshot_observed_at":"2026-08-07T17:34:13.816265Z","submitted_at":"2025-04-30T18:48:06Z","title":"Between Underthinking and Overthinking: An Empirical Study of Reasoning Length and correctness in LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.00127","snapshot_observed_at":"2026-08-06T12:48:08.790009Z","title":"arXiv preprint arXiv:2505.00127","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.790009Z"},"links":{"cited_paper":"/paper/2505.00127","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:7b5926e632b1bdc00ba5f1b7e86e58124ffc43437c322fc12b0dcc6333889ab3","observation_id":"f61fd422-528b-483d-81a9-585fc134065a","resolution":{"observed_at":"2026-08-06T12:48:08.790009Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.21380","last_updated":"2026-04-12T10:37:12Z","snapshot_observed_at":"2026-08-02T16:04:50.458535Z","submitted_at":"2025-03-27T11:20:17Z","title":"Challenging the Boundaries of Reasoning: An Olympiad-Level Math Benchmark for Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.21380","snapshot_observed_at":"2026-08-06T12:48:08.957774Z","title":"arXiv preprint arXiv:2503.21380","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.957774Z"},"links":{"cited_paper":"/paper/2503.21380","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:bc3afe69b6e13c02e9ba63e12cc0474af1a6ffada363d884a4681131cd33397e","observation_id":"a26467da-b9b2-4995-ac92-11b90ac777cb","resolution":{"observed_at":"2026-08-06T12:48:08.957774Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:48:11.549249Z","title":"In Proceedings of the 2022 Confer- ence on Empirical Methods in Natural Language Processing, pages 2866–2879","venue":null,"work_id":"77f364ed-d150-4c0b-bfd7-e012e66814ec","year":2022},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:09.095642Z"},"links":{"citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:74c6e2bde7653f0b3ae7d01179725dee88af92a90246254f55e64c85e12eda8a","observation_id":"96465f1a-7000-4334-ae1f-ad044ac3a944","resolution":{"observed_at":"2026-08-06T12:48:11.719843Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.02642","last_updated":"2024-01-05T05:26:25Z","snapshot_observed_at":"2026-08-05T08:38:24.137256Z","submitted_at":"2024-01-05T05:26:25Z","title":"Signatures of room-temperature superconductivity emerging in two-dimensional domains within the new Bi/Pb-based ceramic cuprate superconductors at ambient pressure","version":1},"cited_work":{"arxiv_id":"2401.02642","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.02642","snapshot_observed_at":"2026-08-06T12:48:10.430487Z","title":"Signatures of room-temperature superconductivity emerging in two-dimensional domains within the new Bi/Pb-based ceramic cuprate superconductors at ambient pressure","venue":"cond-mat.supr-con","work_id":"c6206d2b-c29c-410a-97ba-d09089894de5","year":2024},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:09.185449Z"},"links":{"cited_paper":"/paper/2401.02642","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:4aa6301adb98d6c879210c117faa82a61247cefb10d9e6ebae3467a2021824fc","observation_id":"a0c34740-672c-4533-a901-2407262fa64c","resolution":{"observed_at":"2026-08-06T12:48:10.523027Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:48:11.326899Z","title":"In Proceedings of the 2024 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Tech- nologies, pages 4213–4228","venue":null,"work_id":"5c5952c1-4ddd-46d9-b30a-25b712da0499","year":2024},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:09.311008Z"},"links":{"citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:528ac3a551437f859b2be863dfde7bba3320476b371f444b4a33b1fd1fe3a107","observation_id":"b61d379f-95df-47da-bca1-d3318f3f7d37","resolution":{"observed_at":"2026-08-06T12:48:11.432296Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.14273","last_updated":"2024-01-25T16:08:27Z","snapshot_observed_at":"2026-07-31T01:40:38.269489Z","submitted_at":"2024-01-25T16:08:27Z","title":"Uniformly rotating vortices for the lake equation","version":1},"cited_work":{"arxiv_id":"2401.14273","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.14273","snapshot_observed_at":"2026-08-06T12:48:10.132245Z","title":"Uniformly rotating vortices for the lake equation","venue":"math.AP","work_id":"189f5c88-e813-42f3-ba8c-b5c272cfff63","year":2024},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:09.407105Z"},"links":{"cited_paper":"/paper/2401.14273","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:01ba1aabae89d328ea6264a402b1d47de3ccd8b79c3f318137777967d3c441a0","observation_id":"656c4b5c-aa05-4ebe-8311-f2e126843037","resolution":{"observed_at":"2026-08-06T12:48:10.276133Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:48:09.550871Z","title":"arXiv preprint arXiv:2502.18080","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:09.550871Z"},"links":{"citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:8ebb557a2ac4acdbe588276faffe291b3cacb9c3b952a6672d50a750b6f6540e","observation_id":"9c75b147-50f5-4294-a4eb-4c281786a722","resolution":{"observed_at":"2026-08-06T12:48:09.550871Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10522","last_updated":"2024-12-18T05:21:24Z","snapshot_observed_at":"2026-08-03T17:34:17.504501Z","submitted_at":"2024-06-15T06:26:25Z","title":"Humor in AI: Massive Scale Crowd-Sourced Preferences and Benchmarks for Cartoon Captioning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10522","snapshot_observed_at":"2026-08-06T12:48:09.697591Z","title":"arXiv preprint arXiv:2406.10522","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:09.697591Z"},"links":{"cited_paper":"/paper/2406.10522","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:79b2201b03a5a18e46b25cd838b3d39f94d608abaaccea5b77805b559bfea59e","observation_id":"b202ac87-475a-49a2-a998-8a4f9e111b8b","resolution":{"observed_at":"2026-08-06T12:48:09.697591Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.20356","last_updated":"2025-02-27T18:29:09Z","snapshot_observed_at":"2026-08-09T14:28:31.973792Z","submitted_at":"2025-02-27T18:29:09Z","title":"Bridging the Creativity Understanding Gap: Small-Scale Human Alignment Enables Expert-Level Humor Ranking in LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.20356","snapshot_observed_at":"2026-08-06T12:48:09.853658Z","title":"arXiv preprint arXiv:2502.20356","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:09.853658Z"},"links":{"cited_paper":"/paper/2502.20356","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:abcfb907aa6e1e119753c24097a79c9c44f15d2bcfa53c34f0902373aa19037f","observation_id":"2720ae89-6683-4c2a-a731-9dbeb5c2d960","resolution":{"observed_at":"2026-08-06T12:48:09.853658Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2004.10645","last_updated":"2020-10-05T03:28:21Z","snapshot_observed_at":"2026-08-10T00:47:33.544216Z","submitted_at":"2020-04-22T15:42:13Z","title":"AmbigQA: Answering Ambiguous Open-domain Questions","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2004.10645","snapshot_observed_at":"2026-08-06T12:48:08.379481Z","title":"arXiv preprint arXiv:2004.10645","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.379481Z"},"links":{"cited_paper":"/paper/2004.10645","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:d1caadf7d732f3bd0f7b4409eb5cd7cac3d242013a6eacea37915bc4a423ce68","observation_id":"66882ed1-a3c6-4913-8e54-06cd911d2f7f","resolution":{"observed_at":"2026-08-06T12:48:08.379481Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.14858","last_updated":"2022-07-01T02:15:12Z","snapshot_observed_at":"2026-08-05T15:41:22.691461Z","submitted_at":"2022-06-29T18:54:49Z","title":"Solving Quantitative Reasoning Problems with Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.14858","snapshot_observed_at":"2026-08-06T12:48:08.040390Z","title":"arXiv preprint arXiv:2206.14858","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:08.040390Z"},"links":{"cited_paper":"/paper/2206.14858","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:1ddf6140bd98610e4d8bf40f2e072a3cecfd26c7c3740da5cc70131ba6aeb31c","observation_id":"5a21e266-f982-45af-b54a-c35634c4fe36","resolution":{"observed_at":"2026-08-06T12:48:08.040390Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.17452","last_updated":"2024-02-21T12:59:22Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-09-29T17:59:38Z","title":"ToRA: A Tool-Integrated Reasoning Agent for Mathematical Problem Solving","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.17452","snapshot_observed_at":"2026-08-06T12:48:07.929176Z","title":"arXiv preprint arXiv:2309.17452","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:07.929176Z"},"links":{"cited_paper":"/paper/2309.17452","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:3ac638c06b3fcb2affc7f5af26282a67b3da2955a113752b95730e13725ab38b","observation_id":"1c50d8ef-ea44-4d1b-a529-5caa53466861","resolution":{"observed_at":"2026-08-06T12:48:07.929176Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.04132","last_updated":"2024-03-07T01:22:38Z","snapshot_observed_at":"2026-08-02T17:55:33.750637Z","submitted_at":"2024-03-07T01:22:38Z","title":"Chatbot Arena: An Open Platform for Evaluating LLMs by Human Preference","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.04132","snapshot_observed_at":"2026-08-06T12:48:07.648384Z","title":"arXiv preprint arXiv:2403.04132","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:07.648384Z"},"links":{"cited_paper":"/paper/2403.04132","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:fe485cab93fa1e40f686dbb4ea39b982bc09936095fedc59904fca3da2500018","observation_id":"64403be8-ce55-4562-a051-38c6ef8c9e4b","resolution":{"observed_at":"2026-08-06T12:48:07.648384Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.04604","last_updated":"2025-01-08T05:24:50Z","snapshot_observed_at":"2026-08-04T14:31:04.378756Z","submitted_at":"2024-12-05T20:40:28Z","title":"ARC Prize 2024: Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.04604","snapshot_observed_at":"2026-08-06T12:48:07.697172Z","title":"arXiv preprint arXiv:2412.04604","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-06T12:48:07.697172Z"},"links":{"cited_paper":"/paper/2412.04604","citing_paper":"/paper/2507.21476"},"observation_digest":"sha256:483475c3c8b6a422c979891de6262ade28ea7038a25a57f0f923449d143910fe","observation_id":"3118d44f-b9c8-4fe4-8ee5-0a339bc9a225","resolution":{"observed_at":"2026-08-06T12:48:07.697172Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2507.21476","last_updated":"2025-07-29T03:44:43Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-09T05:51:49.231984Z","submitted_at":"2025-07-29T03:44:43Z","title":"Which LLMs Get the Joke? Probing Non-STEM Reasoning Abilities with HumorBench"},"reference_resolution":{"displayed":27,"state_counts":{"malformed_identifier":0,"metadata_mismatch":4,"parse_uncertain":0,"unresolved":18,"verified_exact":0,"verified_fuzzy":5},"total_outbound_references":27},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 27 of 27 outbound references and 3 inbound Pith citation observations for arXiv:2507.21476."}