Structured entry data for human tools, AI systems and machine clients.
{
"@context": [
"https://schema.org",
{
"lxw": "https://lxkeys.world/schema/"
}
],
"@type": "Thing",
"identifier": "ENT-00002962",
"name": "Civil Comments",
"alternateName": [],
"additionalType": {
"category": "Data and Evaluation",
"type": "Dataset",
"subtype": "",
"lxkeysEntity": false
},
"description": "Civil Comments is a dataset associated with Jigsaw / Conversation AI, classified in LXKeys.world as Language / Training Dataset.",
"creator": "Jigsaw / Conversation AI",
"url": "https://lxkeys.world/entry.php?id=ENT-00002962&lang=en",
"sameAs": "https://www.kaggle.com/c/jigsaw-unintended-bias-in-toxicity-classification",
"image": "",
"lxkeysWorld": {
"worldId": "ENT-00002962",
"kind": "Data and Evaluation",
"type": "Dataset",
"subtype": "",
"classification": "Language / Training Dataset",
"organization": "Jigsaw / Conversation AI",
"originContext": "Language model evaluation and NLP data",
"firstPublicAppearance": "2019",
"currentStatus": "Active",
"documentationStatus": "Documented",
"documentationIndex": {
"total": 82,
"documentation": 25,
"evidence": 13,
"structure": 25,
"relationships": 19,
"level": "Level V Persistent"
},
"lxCalendarium": {
"start_date_utc": "2023-04-01",
"created_utc": "2026-06-17T01:48:33+00:00",
"created_dypclt": "D-0 Y-2 P-3 C-3 L-22 T-4",
"updated_utc": "2026-06-17T01:48:33+00:00",
"updated_dypclt": "D-0 Y-2 P-3 C-3 L-22 T-4"
},
"facts": [],
"capabilities": [
"Dataset reference",
"Benchmark or training-data context",
"Task-level documentation",
"Model evaluation support",
"Source-based traceability"
],
"limitations": [
"Dataset coverage, licensing, annotation quality and benchmark relevance depend on the source version and use context.",
"Performance claims should be assessed through models evaluated on the dataset rather than inferred from the dataset alone."
],
"tags": [
"Dataset",
"NLP",
"Benchmark",
"Language Understanding"
],
"timeline": [
{
"date": "2019",
"title": "Initial dataset release",
"description": "Civil Comments entered the documented public record in 2019. This event is retained at the precision supported by the Entry’s reviewed source history.",
"source_url": "https://www.kaggle.com/c/jigsaw-unintended-bias-in-toxicity-classification",
"verification_status": "source_backed_official"
}
],
"relationships": [
{
"target": "Jigsaw / Conversation AI",
"type": "Associated organization",
"description": "Jigsaw / Conversation AI is the organization, project community or institutional context associated with Civil Comments.",
"evidence_level": "documentary"
},
{
"target": "Language model evaluation and NLP data",
"type": "Domain context",
"description": "Civil Comments belongs to the Language model evaluation and NLP data layer of the intelligent-entity registry.",
"evidence_level": "documentary"
}
],
"sources": [
{
"label": "Official or reference source",
"url": "https://www.kaggle.com/c/jigsaw-unintended-bias-in-toxicity-classification",
"source_type": "Primary / Reference Source",
"verification_status": "verified",
"authority": "OFFICIAL",
"role": "primary",
"update_enabled": true,
"authority_basis": "curated-corpus-refresh-2026-09-08"
}
],
"canonical": [],
"imageMeta": []
}
}