{"data":{"slug":"scrapegraphai-scrapegraph-ai","name":"Scrapegraph-ai","tagline":"Python scraper based on AI","github_url":"https://github.com/ScrapeGraphAI/Scrapegraph-ai","owner":"ScrapeGraphAI","repo":"Scrapegraph-ai","owner_avatar_url":"https://avatars.githubusercontent.com/u/171017415?v=4","primary_language":"Python","stars":29618,"forks":2925,"topics":["ai-crawler","ai-scraping","ai-search","crawler","data-extraction","firecrawl-alternative","large-language-model","llm","markdown","rag","scraping","scraping-python","web-crawler","web-crawlers","web-data","web-data-extraction","web-scraper","web-scraping","web-search","webscraping"],"archived":false,"github_pushed_at":"2026-07-20T14:22:20+00:00","maintenance_label":"Active","stars_delta_30d":1203,"url":"https://www.graphcanon.com/tools/scrapegraphai-scrapegraph-ai","markdown_url":"https://www.graphcanon.com/tools/scrapegraphai-scrapegraph-ai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/scrapegraphai-scrapegraph-ai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=scrapegraphai-scrapegraph-ai","description":"Python scraper based on AI","homepage_url":"https://scrapegraphai.com","license":"MIT","open_issues":12,"watchers":171,"ai_summary":"A Python-based AI-powered web scraping tool designed for data extraction and search tasks.","readme_excerpt":"## 🚀 Quick install\n\nThe reference page for Scrapegraph-ai is available on the official page of PyPI: [pypi](https://pypi.org/project/scrapegraphai/).\n\n```bash\npip install scrapegraphai\n\n---\n\n## 📜 License\n\nScrapeGraphAI is licensed under the MIT License. See the [LICENSE](https://github.com/ScrapeGraphAI/Scrapegraph-ai/blob/main/LICENSE) file for more information.","github_created_at":"2024-01-27T16:54:38+00:00","created_at":"2026-07-07T17:32:21.332086+00:00","updated_at":"2026-08-16T18:01:41.666314+00:00","categories":[{"slug":"data-retrieval","name":"Data & Retrieval","url":"https://www.graphcanon.com/categories/data-retrieval","markdown_url":"https://www.graphcanon.com/categories/data-retrieval.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/data-retrieval"},{"slug":"llm-frameworks","name":"LLM Frameworks","url":"https://www.graphcanon.com/categories/llm-frameworks","markdown_url":"https://www.graphcanon.com/categories/llm-frameworks.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/llm-frameworks"}],"tags":[{"slug":"ai-crawler","name":"ai-crawler"},{"slug":"crawler","name":"crawler"},{"slug":"data-extraction","name":"data-extraction"},{"slug":"large-language-model","name":"large-language-model"},{"slug":"llm","name":"llm"},{"slug":"web-crawler","name":"web-crawler"},{"slug":"webscraping","name":"webscraping"}],"trust":{"provenance":{"is_fork":false,"github_id":749126547,"owner_type":"Organization","methodology":"github_public_v1","parent_repo":null,"near_duplicate_slugs":[]},"computed_at":"2026-08-16T18:01:40.801Z","maintenance":{"label":"Active","score":82,"methodology":"github_public_v1","releases_90d":9,"days_since_push":27,"last_release_at":"2026-07-20T14:22:25Z","stars_delta_30d":1203,"open_issues_delta_30d":5},"security_summary":{"status":"no_lockfile","scanner":null,"low_count":0,"high_count":0,"last_scan_at":"2026-07-11T10:58:44.616Z","medium_count":0,"scan_profile":"none","critical_count":0}},"capability_facts":{"scan":{"source":"repo_scan","observed_at":"2026-08-16T18:01:41.298Z"},"deploy":{"source":"dockerfile:Dockerfile","self_host":true,"observed_at":"2026-08-16T18:01:41.298Z","managed_saas":false},"languages":{"value":["python"],"source":"github.language+pyproject.toml","observed_at":"2026-08-16T18:01:41.298Z"},"has_docker":{"value":true,"source":"dockerfile:Dockerfile","observed_at":"2026-08-16T18:01:41.298Z"},"license_spdx":{"value":"MIT","source":"github.license","observed_at":"2026-08-16T18:01:41.298Z"}},"decision_facts":{"hosting":null,"pricing":{"model":"freemium","summary":"Free to use with potential charges for advanced features or premium services, as specified by its license."},"requirements":{"notes":["Developed in Python and can leverage existing libraries and frameworks related to AI and web scraping."],"min_ram_gb":null,"requires_docker":false},"constraints":{"min_ram_gb":null,"pricing_model":"freemium","requires_docker":false},"when_to_use":["When you need advanced AI capabilities to parse and understand the context of scraped web content.","If your project requires integrating with large language models (LLMs) for complex information retrieval from the web.","For tasks that require sophisticated handling of data extraction, where traditional scraping methods are not sufficient."],"when_not_to_use":["If you have simple and straightforward data extraction needs that do not require AI-powered intelligence.","When your project aims to scrape static or relatively unchanging datasets from websites with well-defined structures, as Scrapegraph-ai is more geared towards complex and dynamic scraping tasks where抠"],"source":"enrich:decision_facts","observed_at":"2026-07-11T13:14:43.601Z"},"constraint_facets":{"min_ram_gb":null,"pricing_model":"freemium","requires_docker":false},"decision_summary":[{"label":"Pricing","value":"freemium - Free to use with potential charges for advanced features or premium services, as specified by its license."},{"label":"Requirements","value":"Developed in Python and can leverage existing libraries and frameworks related to AI and web scraping."},{"label":"Adopt for","value":"Scrapegraph-ai is a Python-based scraping tool leveraging AI for smarter data extraction and search tasks."},{"label":"License detail","value":"Scrapegraph-ai operates under the MIT License, offering users flexible rights for modification and distribution of the code."}]}}