{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2014:JU7AMMBM3AMJWYI2YWRJW72ZNW","short_pith_number":"pith:JU7AMMBM","schema_version":"1.0","canonical_sha256":"4d3e06302cd8189b611ac5a29b7f596dbdd5735f2d66a3ebf3573bd34e66c51c","source":{"kind":"arxiv","id":"1405.0580","version":1},"attestation_state":"computed","paper":{"title":"Web Content Classification: A Survey","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":[],"primary_cat":"cs.IR","authors_text":"Prabhjot Kaur","submitted_at":"2014-05-03T12:57:50Z","abstract_excerpt":"As the information contained within the web is increasing day by day, organizing this information could be a necessary requirement.The data mining process is to extract information from a data set and transform it into an understandable structure for further use. Classification of web page content is essential to many tasks in web information retrieval such as maintaining web directories and focused crawling.The uncontrolled type of nature of web content presents additional challenges to web page classification as compared to the traditional text classification, but the interconnected nature o"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"1405.0580","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.IR","submitted_at":"2014-05-03T12:57:50Z","cross_cats_sorted":[],"title_canon_sha256":"1f178a9b42cc95d528ccac255588e89070a6ccf1943c5ea319a59393619938fe","abstract_canon_sha256":"b095e589112dbaeb74068515444b3e10b0199c66db8a75248f28ab355ff5cd4d"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-05-18T02:51:24.833025Z","signature_b64":"MWN3M862RUe2uH4bPkoiVijyVXcIWn4GQq04dXj3+teD+QxR1pajL2fMo15SKDKnndXd/fDE6Kw5eyvjwfZVAg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"4d3e06302cd8189b611ac5a29b7f596dbdd5735f2d66a3ebf3573bd34e66c51c","last_reissued_at":"2026-05-18T02:51:24.832490Z","signature_status":"signed_v1","first_computed_at":"2026-05-18T02:51:24.832490Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Web Content Classification: A Survey","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":[],"primary_cat":"cs.IR","authors_text":"Prabhjot Kaur","submitted_at":"2014-05-03T12:57:50Z","abstract_excerpt":"As the information contained within the web is increasing day by day, organizing this information could be a necessary requirement.The data mining process is to extract information from a data set and transform it into an understandable structure for further use. Classification of web page content is essential to many tasks in web information retrieval such as maintaining web directories and focused crawling.The uncontrolled type of nature of web content presents additional challenges to web page classification as compared to the traditional text classification, but the interconnected nature o"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1405.0580","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"1405.0580","created_at":"2026-05-18T02:51:24.832561+00:00"},{"alias_kind":"arxiv_version","alias_value":"1405.0580v1","created_at":"2026-05-18T02:51:24.832561+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1405.0580","created_at":"2026-05-18T02:51:24.832561+00:00"},{"alias_kind":"pith_short_12","alias_value":"JU7AMMBM3AMJ","created_at":"2026-05-18T12:28:35.611951+00:00"},{"alias_kind":"pith_short_16","alias_value":"JU7AMMBM3AMJWYI2","created_at":"2026-05-18T12:28:35.611951+00:00"},{"alias_kind":"pith_short_8","alias_value":"JU7AMMBM","created_at":"2026-05-18T12:28:35.611951+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":0,"internal_anchor_count":0,"sample":[]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/JU7AMMBM3AMJWYI2YWRJW72ZNW","json":"https://pith.science/pith/JU7AMMBM3AMJWYI2YWRJW72ZNW.json","graph_json":"https://pith.science/api/pith-number/JU7AMMBM3AMJWYI2YWRJW72ZNW/graph.json","events_json":"https://pith.science/api/pith-number/JU7AMMBM3AMJWYI2YWRJW72ZNW/events.json","paper":"https://pith.science/paper/JU7AMMBM"},"agent_actions":{"view_html":"https://pith.science/pith/JU7AMMBM3AMJWYI2YWRJW72ZNW","download_json":"https://pith.science/pith/JU7AMMBM3AMJWYI2YWRJW72ZNW.json","view_paper":"https://pith.science/paper/JU7AMMBM","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=1405.0580&json=true","fetch_graph":"https://pith.science/api/pith-number/JU7AMMBM3AMJWYI2YWRJW72ZNW/graph.json","fetch_events":"https://pith.science/api/pith-number/JU7AMMBM3AMJWYI2YWRJW72ZNW/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/JU7AMMBM3AMJWYI2YWRJW72ZNW/action/timestamp_anchor","attest_storage":"https://pith.science/pith/JU7AMMBM3AMJWYI2YWRJW72ZNW/action/storage_attestation","attest_author":"https://pith.science/pith/JU7AMMBM3AMJWYI2YWRJW72ZNW/action/author_attestation","sign_citation":"https://pith.science/pith/JU7AMMBM3AMJWYI2YWRJW72ZNW/action/citation_signature","submit_replication":"https://pith.science/pith/JU7AMMBM3AMJWYI2YWRJW72ZNW/action/replication_record"}},"created_at":"2026-05-18T02:51:24.832561+00:00","updated_at":"2026-05-18T02:51:24.832561+00:00"}