{
 "operator": "Crawl4AI project",
 "slug": "crawl4ai",
 "docs": "https://github.com/unclecode/crawl4ai",
 "crawler_count": 1,
 "crawlers": [
  {
   "slug": "crawl4ai",
   "name": "Crawl4AI",
   "operator": "Crawl4AI project",
   "operator_slug": "crawl4ai",
   "category": "tool",
   "category_label": "Tools and frameworks",
   "robots_token": "Crawl4AI",
   "user_agent_substring": "Crawl4AI",
   "user_agent_example": "Crawl4AI",
   "respects_robots_txt": "undocumented",
   "respects_robots_txt_label": "operator publishes no robots.txt statement",
   "verification_method": "none",
   "verification_label": "no published verification method",
   "published_ip_ranges_url": null,
   "ip_ranges_endpoint": null,
   "ipv4_prefix_count": 0,
   "ipv6_prefix_count": 0,
   "what_it_is": "An open-source LLM-oriented crawler and scraper library, run by whoever installs it. Like Scrapy, the default user-agent identifies the software and says nothing about who is behind the request.",
   "cost_of_blocking": "You block a library, not an operator: the rule catches a researcher and a bulk scraper equally, and anyone who edits one config line is not caught at all.",
   "operator_docs": "https://github.com/unclecode/crawl4ai",
   "html_url": "https://www.pathwren.workers.dev/crawler/crawl4ai.html",
   "json_url": "https://www.pathwren.workers.dev/crawler/crawl4ai.json",
   "last_reviewed": "2026-09-01"
  }
 ],
 "ip_range_endpoints": []
}