{"as_of":"2026-08-19T04:46:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:5947e447f10fd2b1bf57aea1b3232666a3a6e2cc87c4ce7e56a155fa9e5be97d","coverage":[{"denominator":100,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":100,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T11:03:29.370944Z","state":"measured"},{"denominator":101,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":101,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-18T06:34:40.430872+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T01:02:51.934157Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-06T01:02:53.997511Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"cited_work":{"arxiv_id":"2501.16750","doi":null,"metadata_source":"pith","pith_arxiv_id":"2501.16750","snapshot_observed_at":"2026-08-06T01:02:53.997511Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","venue":"cs.CR","work_id":"6aedf8e0-e5e2-4bc4-9481-56ce30815143","year":2025},"citing_paper":{"arxiv_id":"2508.03990","last_updated":"2025-08-06T00:45:02Z","snapshot_observed_at":"2026-08-15T20:05:25.723773Z","submitted_at":"2025-08-06T00:45:02Z","title":"Are Today's LLMs Ready to Explain Well-Being Concepts?","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-06T01:02:51.934157Z"},"links":{"cited_paper":"/paper/2501.16750","citing_paper":"/paper/2508.03990"},"observation_digest":"sha256:08444986bf9a9e13bca9020b52ab4a60691b4b36fec8ad7fb8acd0694e1d7ffe","observation_id":"16bb6919-7b9a-4b14-bf1b-9cb3e6c48e0e","resolution":{"observed_at":"2026-08-06T01:02:54.096747Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2501.16750/citation-record","integrity":"/paper/2501.16750/integrity","json":"/paper/2501.16750/citation-record.json","paper":"/paper/2501.16750"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:27.902823Z","title":"https://en.wikipedia.org/wiki/ Coleman-Liau_index","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:27.902823Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:6b11086470ebb73d360e3c11db060683cb05da087ae380fa04a3df65cb2fd7ef","observation_id":"3ad7f745-cbff-421d-9ec5-fc59bfb5fbbb","resolution":{"observed_at":"2026-08-10T11:03:27.902823Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:27.907255Z","title":"https://github.com/unitaryai/detoxify","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:27.907255Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:01e6a58f98071a31c9890430e4057de22f4ea2a927d1ac8503e79a2553be180a","observation_id":"03b1fd37-1304-4c6b-ad09-a10ebfb26c4c","resolution":{"observed_at":"2026-08-10T11:03:27.907255Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:27.911411Z","title":"https://gdpr-info.eu/","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:27.911411Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:37a74b21f53bfc284fd669266b6e445dba25be13cdb4902cf7b5ac661b1cc061","observation_id":"d4f50355-d933-4a08-88d7-26d878d1b1b5","resolution":{"observed_at":"2026-08-10T11:03:27.911411Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:27.915052Z","title":"https://chatgpt.com/g/g-w0y3CvDM 9-freddy-griffin","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:27.915052Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:d1f50cc7ac6f35601e00e66c899cb05d5d7c55e8dd81d65ff2a56cf28facb9eb","observation_id":"6dcf60a8-89d3-48d1-939a-e30de2bde689","resolution":{"observed_at":"2026-08-10T11:03:27.915052Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:27.919260Z","title":"https://chatgpt.com/g/g-JlQ9WBdHB-hate /","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:27.919260Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:17d8f80a023cb789fe2165c9aa31d4633cc55b7d3c9eb6aabe272fa9eed26a98","observation_id":"a63bdbfd-cad7-412f-9a4b-8f93bc1429cd","resolution":{"observed_at":"2026-08-10T11:03:27.919260Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:27.937731Z","title":"https://chatgpt.com/g/g-87uTmBE65-ru de-gpt","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:27.937731Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:fb55b636236603116441f0d0c74ecc612fd41fcadc5f2715f4cb8ca962ba2048","observation_id":"fd1fe658-5211-4357-b5f7-f4cc5a9b7645","resolution":{"observed_at":"2026-08-10T11:03:27.937731Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:27.967561Z","title":"https://www.perspectiveapi.com","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:27.967561Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:2f031039e89fde2ea2effe00c191010fec09915b769ea0be1ef0de7687c242a7","observation_id":"0bbde68a-1663-43f4-a4ce-e354f24e7abe","resolution":{"observed_at":"2026-08-10T11:03:27.967561Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:27.987546Z","title":"https://osf.io/edua3/","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:27.987546Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:ab6b258ad411bdf80b98c61ed6f37f42feea3fb7ca8c7e027007fe7685f765e3","observation_id":"3ac119e8-d12a-4ca5-a0b5-a24c13c232be","resolution":{"observed_at":"2026-08-10T11:03:27.987546Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:28.018050Z","title":"https://lmsys.org/blog/2023-03-30-vicuna/","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.018050Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:31ffea6f4aba51f7c45977527328450cd7118fa373f23b46e5094504c045a00e","observation_id":"8edefd31-f421-45fc-b7bc-f66f21388a73","resolution":{"observed_at":"2026-08-10T11:03:28.018050Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:28.048010Z","title":"ADL Task Force Issues Report Detailing Widespread Anti-Semitic Harassment of Journalists on Twitter During 2016 Campaign","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.048010Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:a7e5b06fb79c04d27a82d829ac896309cc3a718f32bab280319600f12a811340","observation_id":"ebaeb6ad-bb19-4144-8aed-535777dd7068","resolution":{"observed_at":"2026-08-10T11:03:28.048010Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:28.056730Z","title":"Online Hate and Harassment: The American Experi- ence 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.056730Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:84e759f97782b247985cba9097d0eaf9699e713d57e6aeec7e335c5bebe4d2dc","observation_id":"e659f29b-ae11-4dec-bf41-e49ed0e947be","resolution":{"observed_at":"2026-08-10T11:03:28.056730Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:28.094242Z","title":"Aunties, Strangers, and the FBI: Online Privacy Concerns and Experiences of Muslim- American Women","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.094242Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:5922adfa2839766cfd20706c5df406926018273256b706f5809454ab0e2c09a1","observation_id":"350e5bd4-1ec6-429b-b67f-34892d7950fc","resolution":{"observed_at":"2026-08-10T11:03:28.094242Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:28.127008Z","title":"Google’s Jigsaw was trying to fight toxic speech with AI","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.127008Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:95abdb91816b4347800af6614c81b339e4d75bcb8a8ca1560be7a4a435a86c28","observation_id":"ef305288-9568-414e-bd57-75a81094b2ab","resolution":{"observed_at":"2026-08-10T11:03:28.127008Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.01680","last_updated":"2023-07-04T12:22:40Z","snapshot_observed_at":"2026-08-18T18:09:10.093843Z","submitted_at":"2023-07-04T12:22:40Z","title":"Robust Hate Speech Detection in Social Media: A Cross-Dataset Empirical Evaluation","version":1},"cited_work":{"arxiv_id":"2307.01680","doi":null,"metadata_source":"pith","pith_arxiv_id":"2307.01680","snapshot_observed_at":"2026-08-10T11:03:29.652452Z","title":"Robust Hate Speech Detection in Social Media: A Cross-Dataset Empirical Evaluation","venue":"cs.CL","work_id":"ea04de47-8b3e-435a-a38b-d75c2f1fb458","year":2023},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.161185Z"},"links":{"cited_paper":"/paper/2307.01680","citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:dc9549b3c4cf885440d73eacd0dea762b38bc7d22c2277e319253acc70f4461f","observation_id":"8a0f5ff9-abec-4355-a5b8-1cdd73640df6","resolution":{"observed_at":"2026-08-10T11:03:29.690850Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:28.189110Z","title":"METEOR: An Auto- matic Metric for MT Evaluation with Improved Correlation with Human Judgments","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.189110Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:6d49ef2f9c14bb3c3c2745d57b65fe02b9905db08671cacb42dea2cb245a9440","observation_id":"085a2934-d5cf-43a2-977e-d98ce24af45b","resolution":{"observed_at":"2026-08-10T11:03:28.189110Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:28.199948Z","title":"The Pushshift Reddit Dataset","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.199948Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:b990abc126a52491f90b799a98110e44761c66d5dfd22c121a5f21a18613f6fb","observation_id":"20622d68-2f80-4d1b-bf9d-8f72e7490f70","resolution":{"observed_at":"2026-08-10T11:03:28.199948Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:28.218992Z","title":"Nuanced Metrics for Measuring Unin- tended Bias with Real Data for Text Classification","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.218992Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:ec535ed0297961b628fcfa85a80049592a969c4a3b950e74bd2189a63d54fd39","observation_id":"b2799bf0-0fce-4029-a73a-5d097b6c6e68","resolution":{"observed_at":"2026-08-10T11:03:28.218992Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:28.232006Z","title":"HateGAN: Adversarial Generative-Based Data Augmentation for Hate Speech Detec- tion","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.232006Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:438625888be09a982a08e84f833d46c44f4cc70a50a2caa9c2309ae181fa9b42","observation_id":"7793d8ca-21ff-48cf-997d-30f68d8a965d","resolution":{"observed_at":"2026-08-10T11:03:28.232006Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:28.241154Z","title":"John, Noah Constant, Mario Guajardo- Cespedes, Steve Yuan, Chris Tar, Brian Strope, and Ray Kurzweil","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.241154Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:4cc801b1757e8f00b60c534d3e776ae53fabc7476f6ef9a10d6bdecd244899b4","observation_id":"e8e5201c-059b-43ca-bd76-21bab6bfc547","resolution":{"observed_at":"2026-08-10T11:03:28.241154Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:28.254179Z","title":"Hate is not Binary: Studying Abusive Behavior of #GamerGate on Twitter","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.254179Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:444cee2a93ee3719a96a45a0882022896f591f199f4ab1a13d734f244c921174","observation_id":"9d5aa832-e0aa-453f-b517-7e8ed0cccfec","resolution":{"observed_at":"2026-08-10T11:03:28.254179Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:28.269724Z","title":"Christiano, Jan Leike, Tom B","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.269724Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:0fd5104ac21aaa7f3aa7616c6cb3897025453927a17a46071f7b7a04903e0eb0","observation_id":"a93cd6a6-24b0-4a66-8eb4-3a6261c457f8","resolution":{"observed_at":"2026-08-10T11:03:28.269724Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.376974Z","title":"Hate Campaign","venue":null,"work_id":"671deb33-d81b-44b1-83c2-87d53db39c0b","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.284063Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:c9ca22c7cf1e5a0c48c7b9f42eda7de0faf4cb2dcba26a6729bd71e1db357230","observation_id":"2df97a1e-c6ba-4887-bd58-49efbe1fe76c","resolution":{"observed_at":"2026-08-10T11:03:31.380797Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.364866Z","title":"Free dolly: Introducing the world’s first truly open instruction-tuned llm, 2023","venue":null,"work_id":"31716884-4072-4b15-bea1-0d21238941e3","year":2023},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.296688Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:ea814cbaf8fb323be4caa1773c1c7386086fe9de374f2652a2fdc78a7a40fbdc","observation_id":"99b942b8-788a-41f9-a1d3-8136f5181f24","resolution":{"observed_at":"2026-08-10T11:03:31.368765Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.05335","last_updated":"2023-04-11T16:53:54Z","snapshot_observed_at":"2026-08-16T15:41:14.308965Z","submitted_at":"2023-04-11T16:53:54Z","title":"Toxicity in ChatGPT: Analyzing Persona-assigned Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.05335","snapshot_observed_at":"2026-08-10T11:03:28.304717Z","title":"Toxicity in Chat- GPT: Analyzing Persona-assigned Language Models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.304717Z"},"links":{"cited_paper":"/paper/2304.05335","citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:593769092bceb1ce8feebcb21f44f5d41a3fd7324a79d3134cda5417fc686fb8","observation_id":"6bc594ad-b278-462e-b48c-96687bb8aec7","resolution":{"observed_at":"2026-08-10T11:03:28.304717Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.353815Z","title":"BERT: Pre-training of Deep Bidirectional Trans- formers for Language Understanding","venue":null,"work_id":"dc0041b7-c6eb-4a3f-8e25-efe95111b2b6","year":2019},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.309122Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:8dd9d5bd1e3f6fbd2716cfaade3d86c0524b5d7a4239d196c58d7f4bd2d485d0","observation_id":"28de4e90-701c-4e55-b3dc-17e14de7a540","resolution":{"observed_at":"2026-08-10T11:03:31.357619Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.343995Z","title":null,"venue":null,"work_id":"1d204b78-352d-49f2-a3c1-53b2ff4fd577","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.315782Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:4f4f5a74a09360547f06b2c3af4de89e505799c5c149cf223bb1478d1eb0649e","observation_id":"eea97b84-2711-4c36-8827-81c04e587bf6","resolution":{"observed_at":"2026-08-10T11:03:31.347120Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.334197Z","title":"Paraphrase a text","venue":null,"work_id":"b716ac48-cfaf-424b-bca8-88121a616901","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.327174Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:5b86102130b5f03eb96fe0f169bd3c1593866549f0a042d29e389244b7895ddb","observation_id":"613a99d5-4c1b-4084-b042-e22588cb4e20","resolution":{"observed_at":"2026-08-10T11:03:31.337317Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.322875Z","title":"Guide: Large Language Models-Generated Fraud, Malware, and Vulnerabilities","venue":null,"work_id":"689a5a1d-5730-4b99-91b2-b197055a9660","year":2023},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.333912Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:c79ca9b7bb9535e93339395086b21648e714025c4865371df3671bc93d7f4063","observation_id":"76ba4c38-f8d0-461e-b25b-b4f8cff0cc8f","resolution":{"observed_at":"2026-08-10T11:03:31.326970Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.312055Z","title":"Black-Box Generation of Adversarial Text Sequences to Evade Deep Learning Classifiers","venue":null,"work_id":"927b827f-c1f3-474a-a8af-68a48068a27e","year":2018},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.350922Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:205be0cfac5fe1055839356bc891a53ff5164de703730c043f761ed750996aab","observation_id":"04699b60-c2af-4c3b-b9f2-8c5f1663c1fe","resolution":{"observed_at":"2026-08-10T11:03:31.315813Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2101.00027","last_updated":"2020-12-31T19:00:10Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-12-31T19:00:10Z","title":"The Pile: An 800GB Dataset of Diverse Text for Language Modeling","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2101.00027","snapshot_observed_at":"2026-08-10T11:03:28.360191Z","title":"The Pile: An 800GB Dataset of Diverse Text for Language Modeling","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.360191Z"},"links":{"cited_paper":"/paper/2101.00027","citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:6721eada002a4702936c42f8b96002be786edbd6e87de216e03d30d847836273","observation_id":"c7948ace-3c0f-4252-a6a9-4621348310dc","resolution":{"observed_at":"2026-08-10T11:03:28.360191Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.301034Z","title":null,"venue":null,"work_id":"4cff11b2-fb11-48a5-9f55-a0e3bf543364","year":2018},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.364370Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:46776402fe7e7e965e561d9e80523f9fea4bb6bc7594c6452088cfd31d3c5beb","observation_id":"7352fa66-8c9f-435c-a322-6dd973b7e909","resolution":{"observed_at":"2026-08-10T11:03:31.304561Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.286079Z","title":"Hancock, and Zakir Durumeric","venue":null,"work_id":"96811aa3-d43c-461b-a238-65968d1480f6","year":2023},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.368006Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:d11ec5e0f743de3b4f0f78148dee8f09f821cdec94f0e008814cad87728787cf","observation_id":"8a497de1-fa71-4e2f-a6e4-cbf5a18d9def","resolution":{"observed_at":"2026-08-10T11:03:31.289569Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.238669Z","title":"ToxiGen: A Large-Scale Machine-Generated Dataset for Adversarial and Implicit Hate Speech Detection","venue":null,"work_id":"40774e9a-1b1f-4453-8944-2ef6b724d41d","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.372041Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:2ac43c4dec1a73f8737d85420c6a318850f5b6d954416eec46a07b8489c24d73","observation_id":"3fe4fc75-27d8-45af-ab7c-d413cf709c09","resolution":{"observed_at":"2026-08-10T11:03:31.267866Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.184376Z","title":"Hess, Kelley P","venue":null,"work_id":"a126cbb7-6d0f-437e-8109-13e0268a5e73","year":1984},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.378471Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:4e96ded7d6d8cc9f83cc719399dc8490ed246aa1da43bb7a5f1e0877c7205c02","observation_id":"c02e3e85-f651-41d7-aad1-4cec723a7953","resolution":{"observed_at":"2026-08-10T11:03:31.213406Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1702.08138","last_updated":"2017-02-27T04:07:36Z","snapshot_observed_at":"2026-08-14T21:14:40.365911Z","submitted_at":"2017-02-27T04:07:36Z","title":"Deceiving Google's Perspective API Built for Detecting Toxic Comments","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1702.08138","snapshot_observed_at":"2026-08-10T11:03:28.386994Z","title":"Deceiving Google’s Perspective API Built for Detecting Toxic Comments","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.386994Z"},"links":{"cited_paper":"/paper/1702.08138","citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:22bcc133f0e79fddbecd9d08ccad1362937192af78e54e5e85b6ba6406d7f182","observation_id":"cdf58a07-63fb-49cd-8294-29708f51959c","resolution":{"observed_at":"2026-08-10T11:03:28.386994Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.114152Z","title":"Adversarial Example Generation with Syntactically Controlled Paraphrase Networks","venue":null,"work_id":"b783af63-ea9d-4bb0-bd66-5844eb6672a7","year":2018},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.390871Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:bf8ce6c00b8058c453370b7d11da76d761a25d56cd51e8029ad2e09c6cbb0727","observation_id":"c7bd1e3f-72df-40d0-b70b-8a873d84cdfb","resolution":{"observed_at":"2026-08-10T11:03:31.149871Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.070015Z","title":"High Accuracy and High Fi- delity Extraction of Neural Networks","venue":null,"work_id":"2d8a829b-2c28-4860-a3a8-49b4c3cb13c5","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.416115Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:3cd6e3e8a39357f6c4959eb06830026b8add7de0d02fe7a5c9d3f84f2b64b143","observation_id":"867775fb-fa87-4bb4-ab05-25d298584f0d","resolution":{"observed_at":"2026-08-10T11:03:31.082157Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:28.447493Z","title":"Is BERT Really Robust? A Strong Baseline for Natural Lan- guage Attack on Text Classification and Entailment","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.447493Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:86597dd580a60694a54b8b0d4e1215542f9669c2f85bd4d8c783024c43ba9f2e","observation_id":"18b20293-f9a8-4252-9dcc-d9f66218911d","resolution":{"observed_at":"2026-08-10T11:03:28.447493Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.051215Z","title":"Toxic Comment Classification Challenge, 2017","venue":null,"work_id":"bf63a6e1-49c7-49e5-8bbd-0c3f73941520","year":2017},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.484402Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:701e0ae353d9e9b7cdc9b22aa6251730d1142f20aab5e7da039305ff3295e5ec","observation_id":"68f67a60-d589-40c7-854f-bb23bb8592e4","resolution":{"observed_at":"2026-08-10T11:03:31.055167Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.038583Z","title":"Jigsaw Unintended Bias in Toxicity Classification,","venue":null,"work_id":"1bd31836-5cf3-4b39-b574-df3f6ca9090b","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.491766Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:7aba160ce145d20a1ef21fc1ba359021f17c97b4a4e2337409856f5ea90e517e","observation_id":"850c24bd-9013-43db-a4cb-0eb5d49bbb12","resolution":{"observed_at":"2026-08-10T11:03:31.042736Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.027463Z","title":null,"venue":null,"work_id":"6722a8a2-67bc-407e-bb73-8e3943f2e537","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.537604Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:31d7fb3f2c7e75b82b73fc8b6f4c47362c53dcb76ff3cdc94115f51f5c72f43c","observation_id":"3560f2f0-ae31-4139-b05c-a2b323d259dc","resolution":{"observed_at":"2026-08-10T11:03:31.030756Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.015884Z","title":"Content Analysis: An Introduction to Its Methodology","venue":null,"work_id":"25cd515b-9b6e-4ac3-993e-2f616b4d7431","year":2018},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.567684Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:9cebb3467ee15b84e6ed2531f74ce3e8c51eca0103a881dd60f977073cf9f98e","observation_id":"1119b65d-bc8f-44bd-80f4-5b0d50529dbe","resolution":{"observed_at":"2026-08-10T11:03:31.020120Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:31.002456Z","title":"Parikh, Nico- las Papernot, and Mohit Iyyer","venue":null,"work_id":"dcd0597f-0353-41aa-bff5-b3206b895751","year":2020},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.590059Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:b86fda9b99a4bc77b9f16b14832e84098d00bb68a6e8dbe3445e6ca55d9132bd","observation_id":"fd7a6e8a-f462-44dc-b6ce-68e12e98ba1b","resolution":{"observed_at":"2026-08-10T11:03:31.006299Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2108.12521","last_updated":"2021-08-27T22:47:02Z","snapshot_observed_at":"2026-08-16T18:00:26.332660Z","submitted_at":"2021-08-27T22:47:02Z","title":"TweetBLM: A Hate Speech Dataset and Analysis of Black Lives Matter-related Microblogs on Twitter","version":1},"cited_work":{"arxiv_id":"2108.12521","doi":null,"metadata_source":"pith","pith_arxiv_id":"2108.12521","snapshot_observed_at":"2026-08-10T11:03:29.534349Z","title":"TweetBLM: A Hate Speech Dataset and Analysis of Black Lives Matter-related Microblogs on Twitter","venue":"cs.CL","work_id":"8c914584-3bf6-4428-8d30-c2c8ba90bb4e","year":2021},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.620351Z"},"links":{"cited_paper":"/paper/2108.12521","citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:1c11096a18d4da7bfe9028505870c5193eca5269e11f890b3b98c98231f0ec1b","observation_id":"af4b1e44-f981-4f32-92e8-7b956d279b36","resolution":{"observed_at":"2026-08-10T11:03:29.570718Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.991376Z","title":"TextBugger: Generating Adversarial Text Against Real-world Applications","venue":null,"work_id":"24bba327-8a5b-464a-a674-fae0d09f50dd","year":2019},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.652938Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:e6fefbc66144186e573e2cf102cdd01e9e99d70e366b3b3b8df9a296a4a03e7d","observation_id":"5ba05853-35ad-47aa-80bb-f9575da20453","resolution":{"observed_at":"2026-08-10T11:03:30.995160Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1907.11692","last_updated":"2019-07-26T17:48:29Z","snapshot_observed_at":"2026-08-16T14:33:50.657682Z","submitted_at":"2019-07-26T17:48:29Z","title":"RoBERTa: A Robustly Optimized BERT Pretraining Approach","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1907.11692","snapshot_observed_at":"2026-08-10T11:03:28.689441Z","title":"RoBERTa: A Robustly Opti- mized BERT Pretraining Approach","venue":null,"work_id":null,"year":1907},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.689441Z"},"links":{"cited_paper":"/paper/1907.11692","citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:59267758317e494df0beeb73981ddb17401ac6afb52268fec68aaf9707a0356b","observation_id":"4d76440f-0804-46ca-aff2-72fd0895ebf8","resolution":{"observed_at":"2026-08-10T11:03:28.689441Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.972880Z","title":"A Holistic Approach to Undesired Content Detection in the Real World","venue":null,"work_id":"0dbfab23-674e-4364-b230-ec7711d23236","year":2022},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.710325Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:421a4d780b1905bf7d766ba032f70aa4348ef004498668127b725e8eafbd00d2","observation_id":"cd03df65-d1a3-4b55-95ec-c71b069e6fb4","resolution":{"observed_at":"2026-08-10T11:03:30.983943Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.946615Z","title":"HateX- plain: A Benchmark Dataset for Explainable Hate Speech De- 14 tection","venue":null,"work_id":"39a84c4b-9575-4c05-b9b2-840d52aeb5d5","year":2021},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.727341Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:0b2be7dd9238e14637f437438d7ee4ffeb7d960cdfedb466a9c229230fdc019e","observation_id":"0a6cefe5-9337-45eb-991d-8652c1d33b41","resolution":{"observed_at":"2026-08-10T11:03:30.958534Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.921032Z","title":"AI Trained on 4Chan Becomes ‘Hate Speech Machine’","venue":null,"work_id":"dc2483c6-5597-450f-989b-3ac55db0dd6b","year":2022},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.739293Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:e93dac2342247c9594f20f4109ba14727d7c8a914d9d14d0376e5e518bfc85a5","observation_id":"515c39ea-1658-4802-a2eb-f09dec649daa","resolution":{"observed_at":"2026-08-10T11:03:30.930790Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.860002Z","title":"Mazurek, Florian Schaub, and Elissa M","venue":null,"work_id":"bcf8025c-ac07-4a0c-bb06-08feaabc11d1","year":2021},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.749638Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:7f669c8ff998318d5d35727dc95a56e337daef7a994d5b8a6c85116adb2d9100","observation_id":"fbfdfdfb-1768-4df7-a7e6-fe61556966a1","resolution":{"observed_at":"2026-08-10T11:03:30.892022Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.765169Z","title":"The Challenge of Detecting Hate Speech","venue":null,"work_id":"2030586f-e065-42da-b43b-b6b4e9e757f6","year":2022},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.766621Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:5a18f9adfe9ac038f82ea62c22016d61d4178c9e849549d6d04d9d058cdad059","observation_id":"76e81ca1-999c-42c7-93f7-be9c5dde4df2","resolution":{"observed_at":"2026-08-10T11:03:30.796736Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.743238Z","title":null,"venue":null,"work_id":"df616e15-07f3-46e7-a0b3-6e56f40dea34","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.783465Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:6c7d2e2cb2e1241022055ed94d696c5711474991c4cb495603e51a294109f404","observation_id":"df5f1a68-3d31-4125-893d-ab986c018863","resolution":{"observed_at":"2026-08-10T11:03:30.746708Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.731276Z","title":"What is hate speech","venue":null,"work_id":"e709b59a-9531-4f14-a975-bed3d9d0f7c5","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.800210Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:3f52c19f563f28b9c5a1aab429ca5270d560c56b5e114a4d233dcb7f46ca3ef1","observation_id":"54135c04-8df8-43c4-97a4-01d6cffc24ee","resolution":{"observed_at":"2026-08-10T11:03:30.735138Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.720279Z","title":"Handling Disagreement in Hate Speech Modelling","venue":null,"work_id":"febe86fc-9729-4b55-9213-22be172dbb6a","year":2022},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.805787Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:25a6f1dfda0a1f7e2b679dd191234fc266c2bf7a2900ac6e7dd2414e867e96e3","observation_id":"4c2f3670-5515-45eb-8bde-5f11cd7172ee","resolution":{"observed_at":"2026-08-10T11:03:30.723835Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.708711Z","title":"I Know What You Trained Last Summer: A Survey on Stealing Ma- chine Learning Models and Defences","venue":null,"work_id":"5760bff4-1291-4bc2-8d6c-43d2527c1c3a","year":2023},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.810279Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:2370b578003984f527452878a1da5a2bcf2d8a663d63fb5fee87c7abb545f3d3","observation_id":"e5466900-68dc-497c-b282-cf1c31e194ff","resolution":{"observed_at":"2026-08-10T11:03:30.712642Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.696508Z","title":null,"venue":null,"work_id":"f21d7ea0-ae5d-4dbb-b3e3-f3300c6ce599","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.816793Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:7a577117ca02895146e9a5ee101d958dfc4f4a62e9503997859c71360a95606e","observation_id":"dd69ecf4-ebe1-418a-9767-1f520db815f3","resolution":{"observed_at":"2026-08-10T11:03:30.700183Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.685578Z","title":"Introducing GPTs","venue":null,"work_id":"f78b0d7f-8d95-444b-a8f5-b7db89257823","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.830570Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:85b2bc8ef97ef0efc3c501c6e2c2871659aa47f56394954e996ca47ebee4a855","observation_id":"566f0571-34f3-4fe8-8d60-856a0dd5e825","resolution":{"observed_at":"2026-08-10T11:03:30.688919Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-17T09:58:46.058102Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-10T11:03:28.836125Z","title":"GPT-4 Technical Report","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.836125Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:41eefb8f5ac859d20feb627e275a591701f195c0389cb9b11ccc743f964dbb93","observation_id":"be5b8a04-8b95-4a9e-b688-049aef60b248","resolution":{"observed_at":"2026-08-10T11:03:28.836125Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.675291Z","title":"Offensive and Hateful Text Multiclassification","venue":null,"work_id":"59dd61a1-91da-467e-9290-6ee437786f87","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.847019Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:037ef0fa6083d505cc4bd2dda72834661c7993e2e594c3892dc0e2fdc5ed4aba","observation_id":"055fc79f-3b4f-4522-bd68-cddc1c2c34b6","resolution":{"observed_at":"2026-08-10T11:03:30.678911Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.647879Z","title":"Facebook’s race-blind practices around hate speech came at the expense of Black users, new docu- ments show","venue":null,"work_id":"7e2db217-975c-44e6-8acc-22bcb77a6372","year":2021},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.857094Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:330eb44ed07af8ef582a88694889f19cb75fb2b267fe64f364e163eb03cba8b0","observation_id":"4512db10-1eff-4c35-b2cd-94e30131da9d","resolution":{"observed_at":"2026-08-10T11:03:30.658933Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.03486","last_updated":"2025-09-11T08:50:08Z","snapshot_observed_at":"2026-08-16T13:55:02.344221Z","submitted_at":"2024-05-06T13:57:03Z","title":"UnsafeBench: Benchmarking Image Safety Classifiers on Real-World and AI-Generated Images","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.03486","snapshot_observed_at":"2026-08-10T11:03:28.872472Z","title":"UnsafeBench: Benchmarking Image Safety Classifiers on Real-World and AI-Generated Im- ages","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.872472Z"},"links":{"cited_paper":"/paper/2405.03486","citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:ffd5b72ac36d05a596abcc5b79fd6210f7e25f1e012f51c1f2591cc94746457a","observation_id":"cc9834e9-ce0b-4570-978c-f1483673bc5a","resolution":{"observed_at":"2026-08-10T11:03:28.872472Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.606635Z","title":"Gener- ating Natural Language Adversarial Examples through Proba- bility Weighted Word Saliency","venue":null,"work_id":"20e95aa3-77ad-4d70-b8c0-b8015d2232ca","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.884560Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:6076d81409a6bdc58ba9e27d3c959695872a0cd48af70901e5ede440daf2a198","observation_id":"bc3954aa-717b-4681-abc0-a62a2463d96d","resolution":{"observed_at":"2026-08-10T11:03:30.626013Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.539195Z","title":"Schuller","venue":null,"work_id":"92771d61-2671-485e-99a7-2684ae3f8037","year":2019},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.891454Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:6618d8bb857c46c66be5eeaf279ed3f2af2918f401913511e7b3c819e78f7931","observation_id":"5d432944-2249-4239-95ba-b70f0ef1aa70","resolution":{"observed_at":"2026-08-10T11:03:30.571183Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.477929Z","title":"Margetts, and Janet B","venue":null,"work_id":"17f1c01d-bd40-4797-a3ca-24ff02307db0","year":2021},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.902216Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:ebf0a63ef737ab0ff695bf0fee4623ef6da244107b3f4af438e05f24cd7ab2f1","observation_id":"10176289-a493-4eb0-9720-25beae53e1c8","resolution":{"observed_at":"2026-08-10T11:03:30.505330Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.419244Z","title":"Sachdeva, Renata Barreto, Geoff Bacon, Alexander Sahn, Claudia von Vacano, and Chris J","venue":null,"work_id":"b6978333-3b88-4dcd-aeef-16d7eabbfb6c","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.906112Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:ece164ef873b825b5dbb5c9627afd4264ed8b29348555395be49ff04e46f960a","observation_id":"5c772455-2cf9-45d1-aa32-3115a565f95c","resolution":{"observed_at":"2026-08-10T11:03:30.443534Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.369484Z","title":"TUBERAIDER: Attributing Coordinated Hate Attacks on YouTube Videos to their Source Communities","venue":null,"work_id":"c77ee518-658a-4a1b-9a32-cac3dead2ea1","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.914055Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:01f76f90e6383f2488764e8e92c4ebab7ae4d714cd5c3186bd58cf50cf80bd1e","observation_id":"574f0623-0853-4351-997b-924c24e0f4f9","resolution":{"observed_at":"2026-08-10T11:03:30.373777Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.357569Z","title":"Generative AI as a Vector for Harassment and Harm","venue":null,"work_id":"0f9694f9-7dfa-49c4-bccc-d1a022f8fd72","year":2023},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.917853Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:9c62eda332c8176ced67e4e2197e32fd626bdad6c8a742136e892455e6426be9","observation_id":"3dd02bca-aeeb-470d-9dc9-b95fc2e2d14b","resolution":{"observed_at":"2026-08-10T11:03:30.361183Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.346668Z","title":"Man who harassed black student online must deliver ‘sincere’ apology, renounce white supremacy","venue":null,"work_id":"65d73905-a361-43f2-9de4-43161dc92b22","year":2018},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.922110Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:17f510ea4334a6b4bd030f83a034f4b6cec28233425bdf0fff32e00dc11e5ce5","observation_id":"b43cb975-e63a-4db4-af32-ce6610a0568f","resolution":{"observed_at":"2026-08-10T11:03:30.350491Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.334310Z","title":"Do Anything Now: Characterizing and Evaluat- ing In-The-Wild Jailbreak Prompts on Large Language Mod- els","venue":null,"work_id":"08a7703d-1568-46c4-baaa-3b410ee41777","year":2024},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.926135Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:3883623ccbb35f465e2a2f5a2eaeab6a288b6ead9b884600d3378ef458b450a7","observation_id":"c945fcff-def5-4e11-b494-4b0f1d8c8bb8","resolution":{"observed_at":"2026-08-10T11:03:30.338709Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.323662Z","title":"On Xing Tian and the Perseverance of Anti-China Sentiment Online","venue":null,"work_id":"69a0241c-d7b8-4986-a6b0-448140ff5df4","year":2022},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.930210Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:6cead73f34e68d545077516d1f0ea2a1195b1b04097256bfcd0ade26199b3c21","observation_id":"8cd9120a-fe98-49f3-845c-8e79e8e9c624","resolution":{"observed_at":"2026-08-10T11:03:30.327376Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.312238Z","title":"Model Stealing Attacks Against Inductive Graph Neural Networks","venue":null,"work_id":"1b732ee5-52e1-41be-af47-01964b783df0","year":2022},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.933692Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:ed458a9dc282388bb05077f734bddedc26d715116d8074f3dd07211ba1b30ee9","observation_id":"50ec7bac-64e6-4f6e-8608-02e913eced0b","resolution":{"observed_at":"2026-08-10T11:03:30.316013Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.299888Z","title":"Analyzing the Targets of Hate in Online Social Media","venue":null,"work_id":"9ff06ad4-e9c7-4ff2-9790-2e7959f5c917","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.937591Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:f7e42b46e250e10513eeaa0f3fac0a64b2072ae77f66976a4a8d3f2be810fbd8","observation_id":"efd0ea7a-3324-4732-bbf3-3ddaf119f503","resolution":{"observed_at":"2026-08-10T11:03:30.304658Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2009.01325","last_updated":"2022-02-15T19:09:36Z","snapshot_observed_at":"2026-08-17T15:45:31.564132Z","submitted_at":"2020-09-02T19:54:41Z","title":"Learning to summarize from human feedback","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.01325","snapshot_observed_at":"2026-08-10T11:03:28.955630Z","title":"Ziegler, Ryan Lowe, Chelsea V oss, Alec Radford, Dario Amodei, and Paul F","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.955630Z"},"links":{"cited_paper":"/paper/2009.01325","citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:7a407eb4a8a928318f9d71fecfe1988f46bf56cbd66586ffe207d2783fd93557","observation_id":"4f3f5cb9-3a3b-48a8-877b-806ceca6bd09","resolution":{"observed_at":"2026-08-10T11:03:28.955630Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.239519Z","title":null,"venue":null,"work_id":"d10b8e53-e2b5-4555-b489-9cdd333a37b6","year":2021},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.980546Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:03fc90f5db410c272d1a51b69f8c1d7e9806952db908457dce925336ff9072cb","observation_id":"2ef07ce3-c05c-4fd6-ac2a-f360ec628db1","resolution":{"observed_at":"2026-08-10T11:03:30.290822Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.201867Z","title":"Large-Scale Hate Speech Detection with Cross-Domain Transfer","venue":null,"work_id":"93edf017-d6af-4ea4-9d71-69c74c6f6e33","year":2022},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.015186Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:217d43cf4d458744c5c6071e771441a325e2eb3ccf98e5d7592561c376d42677","observation_id":"0c967b31-044a-49d2-a191-b7678e5d6db1","resolution":{"observed_at":"2026-08-10T11:03:30.210478Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2302.13971","last_updated":"2023-02-27T17:11:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-02-27T17:11:15Z","title":"LLaMA: Open and Efficient Foundation Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.13971","snapshot_observed_at":"2026-08-10T11:03:29.053661Z","title":"LLaMA: Open and Efficient Foundation Language Mod- els","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.053661Z"},"links":{"cited_paper":"/paper/2302.13971","citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:5dc2cc76102bc59bbb6d0a545b1ec466e283d189fcfc797b598410e2a0ba7af4","observation_id":"e1d50bf6-b164-47f5-9bf6-7978ac6476dc","resolution":{"observed_at":"2026-08-10T11:03:29.053661Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.190061Z","title":"Reiter, and Thomas Ristenpart","venue":null,"work_id":"77ff6bfa-a52a-4772-829f-fa6646a3dcb5","year":2016},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.071955Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:e86836d9fbd8923bc937ba24eeb0e2e9d9342310fb65fd99bf6a197c7d79695e","observation_id":"4913f0b0-d5dc-4d5d-8ad7-f434389b6bda","resolution":{"observed_at":"2026-08-10T11:03:30.194371Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:29.094145Z","title":"Visualizing Data using t-SNE","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.094145Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:97d90f8e51c5b599a499ceace7766eb7423dfddbc40d885fcf7f5732e1be20a9","observation_id":"a76556e8-7611-4441-91cb-99eca8b791e6","resolution":{"observed_at":"2026-08-10T11:03:29.094145Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.171553Z","title":"Learning from the Worst: Dynamically Generated Datasets to Improve Online Hate Detection","venue":null,"work_id":"97328453-c0f8-4453-8678-f15f27cb6feb","year":2021},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.123862Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:c059fe3ecd59411d67f0acaa7e87a0f0cc976441d683bcb5c20696e531605c3a","observation_id":"62f85124-fd18-4fc4-b752-6cbff76e2090","resolution":{"observed_at":"2026-08-10T11:03:30.175553Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.159986Z","title":"Moderating New Waves of Online Hate with Chain-of- Thought Reasoning in Large Language Models","venue":null,"work_id":"7f8d71fe-d9bb-4805-848d-0eb46096db24","year":2024},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.159151Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:ce1f74f2581549dcc120fd7d32fc7e58011ff8a5f918c1ef4966bc68554ec03f","observation_id":"d3ad8727-8c8a-46d9-8470-645e194b289e","resolution":{"observed_at":"2026-08-10T11:03:30.164132Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.138087Z","title":"Vu, Alice Hutchings, and Ross J","venue":null,"work_id":"a6dd49af-33f7-4271-8f6f-5fa779b1cc3d","year":2024},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.193217Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:0293aa767e64ba4967df1739699bfc2a54dd300416ed0ba67836d9b221cd2bcc","observation_id":"a7abee6c-1f43-47e3-9d9c-7ef49454d838","resolution":{"observed_at":"2026-08-10T11:03:30.152770Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.116973Z","title":"There’s so much responsibility on users right now: Expert Advice for Staying Safer From Hate and Harassment","venue":null,"work_id":"1c355870-7dc1-4095-bbde-c46df842bdfb","year":2023},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.206641Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:08e898e65453b3378a7a959535651fb48b47eb9a43cf4dcf6a091551ba77b929","observation_id":"3ffcd2df-e5bf-4d0c-ba60-94956a9dc078","resolution":{"observed_at":"2026-08-10T11:03:30.125214Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2109.07445","last_updated":"2021-09-15T17:27:06Z","snapshot_observed_at":"2026-08-16T17:56:21.563256Z","submitted_at":"2021-09-15T17:27:06Z","title":"Challenges in Detoxifying Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2109.07445","snapshot_observed_at":"2026-08-10T11:03:29.226885Z","title":"Challenges in Detoxifying Language Models","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.226885Z"},"links":{"cited_paper":"/paper/2109.07445","citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:acaca7ae37f9562926b429f4b9528097c359307208a7512cf5637268c7c3da38","observation_id":"ee32c18a-1f36-4b04-b4aa-bc5d59673f41","resolution":{"observed_at":"2026-08-10T11:03:29.226885Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.083916Z","title":"Not All Asians are the Same: A Disaggregated Approach to Identify- ing Anti-Asian Racism in Social Media","venue":null,"work_id":"6c0f4b74-f126-4101-a03b-f69b4b93168b","year":2024},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.232894Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:7fefdc535aaeb32f7dbd48ba9b7044cee71e899e8511c9a2920e1dbe08ded457","observation_id":"10dcf881-889b-49a2-bdee-0601e963f804","resolution":{"observed_at":"2026-08-10T11:03:30.097329Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:29.986190Z","title":"Image-Perfect Imperfections: Safety, Bias, and Authenticity in the Shadow of Text-To-Image Model Evolution","venue":null,"work_id":"54a5a5d2-e0cc-4a70-995f-fdcda228e52e","year":2024},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.239336Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:bd4ea1b95360e380071ea1a6d2092d231a31114cde6418ade085a96b7d241590","observation_id":"cc1edeb2-f072-4d71-9e79-b69876745678","resolution":{"observed_at":"2026-08-10T11:03:30.058739Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:29.924230Z","title":"Fight Fire with Fire: Fine-tuning Hate Detectors using Large Samples of Generated Hate Speech","venue":null,"work_id":"090db00b-0e0d-40c5-a632-dc6f5c1b25b9","year":2021},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.244624Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:85c5cee9fa82a52a001c2e9dcc2275b032deda22aa69c942f30d863f18ca82bd","observation_id":"afb446ee-7248-4722-8ea3-e0267bce9bd5","resolution":{"observed_at":"2026-08-10T11:03:29.928514Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.10305","last_updated":"2025-04-17T08:34:42Z","snapshot_observed_at":"2026-08-14T19:59:05.307983Z","submitted_at":"2023-09-19T04:13:22Z","title":"Baichuan 2: Open Large-scale Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.10305","snapshot_observed_at":"2026-08-10T11:03:29.254537Z","title":"Baichuan 2: Open Large-scale Language Models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.254537Z"},"links":{"cited_paper":"/paper/2309.10305","citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:6683008b919b438b420558f9d8a37e9c244d6c3a46db03e60fa6cc6419e0f04f","observation_id":"7d101714-7083-4b0b-9768-8b2465c0bf0d","resolution":{"observed_at":"2026-08-10T11:03:29.254537Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:29.912670Z","title":"GPT-4chan: This is the worst AI ever.https: //tinyurl.com/2s4jh5p4, 2022","venue":null,"work_id":"4ebac5a0-cbd1-49a6-b930-ed87fc1d0bf3","year":2022},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.263941Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:ffcfff517b8069ae0b19a153734ac9d876e6277c2d178508e51bd065cab76857","observation_id":"5e69bce8-4c74-4ab1-b41a-2832fbc0011e","resolution":{"observed_at":"2026-08-10T11:03:29.916299Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:29.891665Z","title":"OpenAttack: An Open-source Textual Adversarial At- tack Toolkit","venue":null,"work_id":"5315065b-54ad-4bdc-ba9b-9265edc30fb1","year":2021},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.276493Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:71890a57c6a625fc69ff9973b147a9fcbe920568872663d07f86ff00da756e06","observation_id":"854195b2-ce6d-410d-84fb-d210bef4997a","resolution":{"observed_at":"2026-08-10T11:03:29.900311Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:29.877453Z","title":"SecurityNet: Assess- ing Machine Learning Vulnerabilities on Public Models","venue":null,"work_id":"8568413c-0caf-4e63-9d2b-9b38ec5dbf46","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":90,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.293814Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:626ad58948776e5279d0e4ef88d78c069865ddc1ecef2220878a8dd0e5076cf9","observation_id":"19b25f2e-afc1-4a23-8ff5-27052e604aa4","resolution":{"observed_at":"2026-08-10T11:03:29.883125Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.01068","last_updated":"2022-06-21T17:04:40Z","snapshot_observed_at":"2026-08-06T03:13:37.403059Z","submitted_at":"2022-05-02T17:49:50Z","title":"OPT: Open Pre-trained Transformer Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.01068","snapshot_observed_at":"2026-08-10T11:03:29.306584Z","title":"Diab, Xian Li, Xi Victoria Lin, Todor Mihaylov, Myle Ott, Sam Shleifer, Kurt Shuster, Daniel Simig, Punit Singh Koura, Anjali Sridhar, Tianlu Wang, and Luke Zettlemoyer","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.306584Z"},"links":{"cited_paper":"/paper/2205.01068","citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:33c33476384fa1b9a20ea068195aa30527d38382b220671d499463e6e497c964","observation_id":"c9b65d46-ebc3-4c3d-990a-fc43901af425","resolution":{"observed_at":"2026-08-10T11:03:29.306584Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:29.866138Z","title":"Generating Natural Adversarial Examples","venue":null,"work_id":"920b0faa-7af8-46e2-bcb2-51408a46622a","year":2018},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":92,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.319248Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:8fb5afc67d04e433dc0ce00633ae9b814dcfbcd7e1b63376c777d42c62cbe22f","observation_id":"cb337bfa-4dfa-43ab-a3e2-8bd60141f990","resolution":{"observed_at":"2026-08-10T11:03:29.870122Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-18T10:45:11.561476Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-10T11:03:29.331421Z","title":"PromptBench: Towards Evaluating the Robustness of Large Language Models on Ad- versarial Prompts","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":93,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.331421Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:96652294136410e9ce47d8a57e6693df4695f4887a7c01f331f5e1566e8c7946","observation_id":"2d3426a3-2e04-4116-adc9-8c80586222b1","resolution":{"observed_at":"2026-08-10T11:03:29.331421Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:30.381591Z","title":"2, 3, 5, 17, 18","venue":null,"work_id":"050d3bd6-9324-46e0-b719-b0e6a50dd7e2","year":2022},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":94,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:28.910119Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:777c5f0fb37fa3a2ae661ea38b1871608f0f0c3826e17faf51959aa7281c91d8","observation_id":"04791f2d-d5d0-434a-8917-1add3bcf9edc","resolution":{"observed_at":"2026-08-10T11:03:30.390138Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2005.12423","last_updated":"2021-11-10T19:15:03Z","snapshot_observed_at":"2026-07-06T09:23:15.448119Z","submitted_at":"2020-05-25T21:58:09Z","title":"Racism is a Virus: Anti-Asian Hate and Counterspeech in Social Media during the COVID-19 Crisis","version":2},"cited_work":{"arxiv_id":"2005.12423","doi":null,"metadata_source":"pith","pith_arxiv_id":"2005.12423","snapshot_observed_at":"2026-08-10T11:03:29.409572Z","title":"Racism is a Virus: Anti-Asian Hate and Counterspeech in Social Media during the COVID-19 Crisis","venue":"cs.SI","work_id":"745c0298-bfcf-4e06-8dcf-8110e225c35c","year":2020},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":95,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.345163Z"},"links":{"cited_paper":"/paper/2005.12423","citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:0303da4dbd1c5694c139ff6ed8b8a8f3e0515fd93bf86f3117bdf1058b646f93","observation_id":"9cc12a26-8c36-4a77-8b09-b58c130cb60a","resolution":{"observed_at":"2026-08-10T11:03:29.417012Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:29.842044Z","title":null,"venue":null,"work_id":"46b8e8a8-6346-44bf-992d-63d5c244ba12","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":97,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.359653Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:702a3d60635160a9189bd9261e12fe9eb4b43977468f727dd573a532bb576fe1","observation_id":"f44ee80e-a821-4bb2-b532-13a3b18b682a","resolution":{"observed_at":"2026-08-10T11:03:29.845892Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:29.816477Z","title":null,"venue":null,"work_id":"3104f8c9-3d38-44a6-b507-9689549e66d7","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":98,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.363030Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:7410ca6b82de6f3476db779d308fec75a572a59baf0c0b6fcadf2938ad2c4aff","observation_id":"e8b5c314-e0f9-4f8d-847e-cd06e98960e6","resolution":{"observed_at":"2026-08-10T11:03:29.832288Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:29.763164Z","title":null,"venue":null,"work_id":"dc642ac3-dcda-4177-acea-50d4b4a61066","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":99,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.367165Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:74abd2c778c84144bb71ba0ff3034509388eea5c993c3725d745fc10d7a9e641","observation_id":"988831f8-7a0a-4902-a5af-3abb1bb49485","resolution":{"observed_at":"2026-08-10T11:03:29.788575Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:29.711861Z","title":"identity attack","venue":null,"work_id":"d96f8be3-d1bc-4e79-81e0-deb61ed86eba","year":2022},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":100,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.370944Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:5f700aad9db6f6e5c9d6a20821dd624d84f5c06d8def85c5cb2a61b8cdc3ae9a","observation_id":"81c29158-d397-4a44-9abc-f140e901b128","resolution":{"observed_at":"2026-08-10T11:03:29.738939Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T11:03:29.854298Z","title":"toxicity,","venue":null,"work_id":"c7c9ad64-d5aa-486e-88a7-5251c49c0515","year":null},"citing_paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns","version":1},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-10T11:03:29.353557Z"},"links":{"citing_paper":"/paper/2501.16750"},"observation_digest":"sha256:97a5605238786101f7e2d03dbdc5a27f3fca6269270cca7e9cfaf868a9f3853a","observation_id":"c3846f5f-3496-42a5-8729-02c84077f5c3","resolution":{"observed_at":"2026-08-10T11:03:29.858217Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2501.16750","last_updated":"2025-01-28T07:00:45Z","latest_version":1,"primary_category":"cs.CR","snapshot_observed_at":"2026-08-10T20:00:38.555772Z","submitted_at":"2025-01-28T07:00:45Z","title":"HateBench: Benchmarking Hate Speech Detectors on LLM-Generated Content and Hate Campaigns"},"reference_resolution":{"displayed":100,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":43,"verified_exact":3,"verified_fuzzy":53},"total_outbound_references":100},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"thesis":"As of 19 August 2026, this Paper Citation Record lists 100 of 100 outbound references and 1 inbound Pith citation observation for arXiv:2501.16750."}