{
  "schema_version": "0.1.0",
  "record_type": "research-topic",
  "topic_id": "benchmark-datasets",
  "label": "Benchmark Datasets",
  "title": "Benchmark Datasets Research",
  "description": "Benchmark Datasets research papers in the Naser Ezzati-Jivan publication catalog.",
  "introduction": "This topic page groups Naser Ezzati-Jivan research papers related to benchmark datasets. Each linked record provides the paper's problem, method, findings, limitations, keywords, and authoritative source links.",
  "aliases": [
    "benchmark datasets"
  ],
  "search_terms": [
    "Benchmark Datasets"
  ],
  "related_topics": [],
  "canonical_url": "https://threadslab.org/research-publications/topics/benchmark-datasets.html",
  "paper_count": 4,
  "papers": [
    {
      "paper_id": "ai-video-retrieval-semantic-search-timestamp-alignment",
      "title": "AI Video Retrieval: A Semantic Search & Timestamp Alignment System",
      "year": 2025,
      "authors": [
        "Hridoy Rahman",
        "Naser Ezzati-Jivan",
        "Blessing Ogbuokiri"
      ],
      "page_url": "https://threadslab.org/research-publications/papers/ai-video-retrieval-semantic-search-timestamp-alignment/",
      "canonical_source_url": "https://doi.org/10.1109/ACDSA65407.2025.11166430",
      "core_contribution": "The paper implements a timestamp-aware multimodal video-retrieval pipeline that joins speech transcription, sampled-frame captioning, text embeddings, and approximate-nearest-neighbor search.",
      "tags": [
        "multimodal-ai",
        "machine-learning",
        "benchmark-datasets"
      ],
      "keywords": [
        "video retrieval",
        "semantic search",
        "timestamp alignment",
        "AI video search",
        "ACDSA 2025"
      ]
    },
    {
      "paper_id": "decoding-log-parsing-challenges-taxonomy",
      "title": "Decoding Log Parsing Challenges: A Comprehensive Taxonomy for Actionable Solutions",
      "year": 2024,
      "authors": [
        "Issam Sedki",
        "Abdelwahab Hamou-Lhadj",
        "Otmane Ait-Mohamed",
        "Naser Ezzati-Jivan",
        "Mohammed A. Shehab"
      ],
      "page_url": "https://threadslab.org/research-publications/papers/decoding-log-parsing-challenges-taxonomy/",
      "canonical_source_url": "https://doi.org/10.1145/3639478.3643523",
      "core_contribution": "The paper derives a 30-item taxonomy of log event characteristics that induce parsing errors and quantifies the characteristics with the largest impact across eight parsers.",
      "tags": [
        "observability",
        "machine-learning",
        "trace-analysis",
        "benchmark-datasets"
      ],
      "keywords": [
        "log parsing",
        "log event characteristics",
        "LEC taxonomy",
        "LogHub",
        "open coding",
        "Drain",
        "IPLoM",
        "AEL",
        "Spell",
        "LenMa",
        "LogMine",
        "SHISO",
        "ULP",
        "log templates",
        "parsing errors",
        "ICSE 2024"
      ]
    },
    {
      "paper_id": "picturing-ambiguity-winograd-schema",
      "title": "Picturing Ambiguity: A Visual Twist on the Winograd Schema Challenge",
      "year": 2024,
      "authors": [
        "Brendan Park",
        "Madeline Janecek",
        "Naser Ezzati-Jivan",
        "Yifeng Li",
        "Ali Emami"
      ],
      "page_url": "https://threadslab.org/research-publications/papers/picturing-ambiguity-winograd-schema/",
      "canonical_source_url": "https://aclanthology.org/2024.acl-long.22/",
      "core_contribution": "The paper introduces WinoVis, a multimodal benchmark and analysis framework for testing pronoun disambiguation in text-to-image models.",
      "tags": [
        "multimodal-ai",
        "benchmark-datasets",
        "common-sense-reasoning",
        "machine-learning"
      ],
      "keywords": [
        "Winograd Schema Challenge",
        "WinoVis",
        "text-to-image models",
        "pronoun disambiguation",
        "DAAM",
        "Stable Diffusion"
      ]
    },
    {
      "paper_id": "automated-categorizing-similarities-persian-news",
      "title": "New Approach for Automated Categorizing and Finding Similarities in Online Persian News",
      "year": 2010,
      "authors": [
        "Naser Ezzati Jivan",
        "Mahlagha Fazeli",
        "Khadije Sadat Yousefi"
      ],
      "page_url": "https://threadslab.org/research-publications/papers/automated-categorizing-similarities-persian-news/",
      "canonical_source_url": "https://doi.org/10.1007/978-3-642-16032-5_11",
      "core_contribution": "The paper combines automated Persian-news categorization with a web system for retrieving similar news items.",
      "tags": [
        "machine-learning",
        "benchmark-datasets"
      ],
      "keywords": [
        "Persian news",
        "text categorization",
        "document similarity",
        "tf-idf",
        "SVM",
        "Reuters",
        "web crawler",
        "PHP",
        "keyword extraction",
        "semantic similarity"
      ]
    }
  ]
}
