{"data":{"node":{"slug":"scrapegraphai-scrapegraph-ai","name":"Scrapegraph-ai","tagline":"Python scraper based on AI","github_url":"https://github.com/ScrapeGraphAI/Scrapegraph-ai","owner":"ScrapeGraphAI","repo":"Scrapegraph-ai","owner_avatar_url":"https://avatars.githubusercontent.com/u/171017415?v=4","primary_language":"Python","stars":29618,"forks":2925,"topics":["ai-crawler","ai-scraping","ai-search","crawler","data-extraction","firecrawl-alternative","large-language-model","llm","markdown","rag","scraping","scraping-python","web-crawler","web-crawlers","web-data","web-data-extraction","web-scraper","web-scraping","web-search","webscraping"],"archived":false,"github_pushed_at":"2026-07-20T14:22:20+00:00","maintenance_label":"Active","stars_delta_30d":1203,"url":"https://www.graphcanon.com/tools/scrapegraphai-scrapegraph-ai","markdown_url":"https://www.graphcanon.com/tools/scrapegraphai-scrapegraph-ai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/scrapegraphai-scrapegraph-ai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=scrapegraphai-scrapegraph-ai"},"categories":[{"slug":"data-retrieval","name":"Data & Retrieval","url":"https://www.graphcanon.com/categories/data-retrieval","markdown_url":"https://www.graphcanon.com/categories/data-retrieval.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/data-retrieval"},{"slug":"llm-frameworks","name":"LLM Frameworks","url":"https://www.graphcanon.com/categories/llm-frameworks","markdown_url":"https://www.graphcanon.com/categories/llm-frameworks.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/llm-frameworks"}],"tags":[{"slug":"ai-crawler","name":"ai-crawler"},{"slug":"crawler","name":"crawler"},{"slug":"data-extraction","name":"data-extraction"},{"slug":"large-language-model","name":"large-language-model"},{"slug":"llm","name":"llm"},{"slug":"web-crawler","name":"web-crawler"},{"slug":"webscraping","name":"webscraping"}],"edges":[{"type":"alternative","direction":"out","explanation":"Scrapegraph-ai and firecrawl both offer tools for extracting data from websites, with Scrapegraph-ai focusing on leveraging large language models to automate scraping processes from web sources and local documents, while firecrawl provides an API-centric approach for searching, scraping, and converting extracted web content into clean formats like Markdown. This alternative relationship stems from","successor_context":null,"tool":{"slug":"firecrawl-firecrawl","name":"firecrawl","tagline":"The API to search, scrape, and interact with the web at scale. 🔥","github_url":"https://github.com/firecrawl/firecrawl","owner":"firecrawl","repo":"firecrawl","owner_avatar_url":"https://avatars.githubusercontent.com/u/135057108?v=4","primary_language":"TypeScript","stars":167794,"forks":9398,"topics":["ai","ai-agents","ai-crawler","ai-scraping","ai-search","crawler","data-extraction","html-to-markdown","llm","markdown","scraper","scraping","web-crawler","web-data","web-data-extraction","web-scraper","web-scraping","web-search","webscraping"],"archived":false,"github_pushed_at":"2026-08-15T19:52:46+00:00","maintenance_label":"Very active","stars_delta_30d":15884,"url":"https://www.graphcanon.com/tools/firecrawl-firecrawl","markdown_url":"https://www.graphcanon.com/tools/firecrawl-firecrawl.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/firecrawl-firecrawl","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=firecrawl-firecrawl"}},{"type":"alternative","direction":"out","explanation":"Both ScrapeGraphAI and trafilatura are tools designed for extracting text and metadata from the web, indicating an alternative approach to similar problems.","successor_context":null,"tool":{"slug":"adbar-trafilatura","name":"trafilatura","tagline":"Python & Command-line tool for web crawling, scraping and text extraction","github_url":"https://github.com/adbar/trafilatura","owner":"adbar","repo":"trafilatura","owner_avatar_url":"https://avatars.githubusercontent.com/u/2125866?v=4","primary_language":"Python","stars":6657,"forks":415,"topics":["article-extractor","corpus-builder","corpus-tools","crawler","html-to-markdown","html2text","llm","news-aggregator","news-crawler","nlp","rag","readability","rss-feed","scraping","tei","text-cleaning","text-extraction","text-mining","text-preprocessing","web-scraping"],"archived":false,"github_pushed_at":"2026-08-15T16:13:21+00:00","maintenance_label":"Very active","stars_delta_30d":343,"url":"https://www.graphcanon.com/tools/adbar-trafilatura","markdown_url":"https://www.graphcanon.com/tools/adbar-trafilatura.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/adbar-trafilatura","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=adbar-trafilatura"}},{"type":"related","direction":"out","explanation":null,"successor_context":null,"tool":{"slug":"graphify-labs-graphify","name":"graphify","tagline":"Turn any code or documentation into a queryable knowledge graph","github_url":"https://github.com/Graphify-Labs/graphify","owner":"Graphify-Labs","repo":"graphify","owner_avatar_url":"https://avatars.githubusercontent.com/u/297659074?v=4","primary_language":"Python","stars":107507,"forks":10441,"topics":["ai-agents","antigravity","ast","claude-code","code-analysis","code-search","codex","cursor","developer-tools","gemini","graphrag","knowledge-graph","leiden","llm","mcp","openclaw","rag","skills","tree-sitter"],"archived":false,"github_pushed_at":"2026-08-17T18:42:58+00:00","maintenance_label":"Very active","stars_delta_30d":16918,"url":"https://www.graphcanon.com/tools/graphify-labs-graphify","markdown_url":"https://www.graphcanon.com/tools/graphify-labs-graphify.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/graphify-labs-graphify","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=graphify-labs-graphify"}},{"type":"integrates_with","direction":"out","explanation":"Due to the nature of ScrapeGraphAI's integration capabilities and its focus on LLMs, it could work well alongside txtai which also handles semantics and language models.","successor_context":null,"tool":{"slug":"neuml-txtai","name":"txtai","tagline":"All-in-one AI framework for semantic search, LLM orchestration and language model workflows","github_url":"https://github.com/neuml/txtai","owner":"neuml","repo":"txtai","owner_avatar_url":"https://avatars.githubusercontent.com/u/59890304?v=4","primary_language":"Python","stars":12890,"forks":873,"topics":["agents","ai","ai-agents","embeddings","information-retrieval","language-model","large-language-models","llm","nlp","python","rag","retrieval-augmented-generation","search","search-engine","semantic-search","sentence-embeddings","transformers","txtai","vector-database","vector-search"],"archived":false,"github_pushed_at":"2026-08-12T13:42:39+00:00","maintenance_label":"Very active","stars_delta_30d":162,"url":"https://www.graphcanon.com/tools/neuml-txtai","markdown_url":"https://www.graphcanon.com/tools/neuml-txtai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/neuml-txtai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=neuml-txtai"}},{"type":"related","direction":"out","explanation":"Both ScrapeGraphAI and browser-use involve automating tasks online, though ScrapeGraphAI focuses on scraping web data using AI.","successor_context":null,"tool":{"slug":"browser-use-browser-use","name":"browser-use","tagline":"Make websites accessible for AI agents. Automate tasks online with ease.","github_url":"https://github.com/browser-use/browser-use","owner":"browser-use","repo":"browser-use","owner_avatar_url":"https://avatars.githubusercontent.com/u/192012301?v=4","primary_language":"Python","stars":109348,"forks":12022,"topics":["ai-agents","ai-tools","browser-automation","browser-use","llm","playwright","python"],"archived":false,"github_pushed_at":"2026-08-15T17:07:06+00:00","maintenance_label":"Very active","stars_delta_30d":4268,"url":"https://www.graphcanon.com/tools/browser-use-browser-use","markdown_url":"https://www.graphcanon.com/tools/browser-use-browser-use.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/browser-use-browser-use","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=browser-use-browser-use"}},{"type":"integrates_with","direction":"out","explanation":"Quivr is an Opiniated RAG for integrating GenAI in applications, which can include the use of web scraping libraries like ScrapeGraphAI to enrich its RAG system with the scraped data.","successor_context":null,"tool":{"slug":"quivrhq-quivr","name":"quivr","tagline":"Opiniated RAG for integrating GenAI in your apps 🧠","github_url":"https://github.com/QuivrHQ/quivr","owner":"QuivrHQ","repo":"quivr","owner_avatar_url":"https://avatars.githubusercontent.com/u/159330290?v=4","primary_language":"Python","stars":39401,"forks":3727,"topics":["ai","api","chatbot","chatgpt","database","docker","framework","frontend","groq","html","javascript","llm","openai","postgresql","privacy","rag","react","security","typescript","vector"],"archived":false,"github_pushed_at":"2025-07-09T12:55:23+00:00","maintenance_label":"Dormant","stars_delta_30d":188,"url":"https://www.graphcanon.com/tools/quivrhq-quivr","markdown_url":"https://www.graphcanon.com/tools/quivrhq-quivr.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/quivrhq-quivr","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=quivrhq-quivr"}},{"type":"alternative","direction":"in","explanation":"Both libraries handle data extraction, with Xberg focusing more on document intelligence while Scrapegraph-ai involves web scraping and integration with LLM.","successor_context":null,"tool":{"slug":"xberg-io-xberg","name":"xberg","tagline":"A polyglot document intelligence framework with a Rust core","github_url":"https://github.com/xberg-io/xberg","owner":"xberg-io","repo":"xberg","owner_avatar_url":"https://avatars.githubusercontent.com/u/241328462?v=4","primary_language":"Rust","stars":9128,"forks":562,"topics":["bun","csharp","document-intelligence","elixir","ffi","golang","java","metadata-extraction","node","pdf-extraction","pdfium","php","python","rag","ruby","rust","table-extraction","tesseract","text-extraction","wasm"],"archived":false,"github_pushed_at":"2026-08-16T16:00:11+00:00","maintenance_label":"Very active","stars_delta_30d":460,"url":"https://www.graphcanon.com/tools/xberg-io-xberg","markdown_url":"https://www.graphcanon.com/tools/xberg-io-xberg.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/xberg-io-xberg","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=xberg-io-xberg"}},{"type":"related","direction":"in","explanation":"Both projects leverage AI for processing structured data, Scrapegraph-AI focuses on web data extraction while LLM4Decompile does binary code decompilation using large language models.","successor_context":null,"tool":{"slug":"albertan017-llm4decompile","name":"LLM4Decompile","tagline":"Decompiling Binary Code with Large Language Models","github_url":"https://github.com/albertan017/LLM4Decompile","owner":"albertan017","repo":"LLM4Decompile","owner_avatar_url":"https://avatars.githubusercontent.com/u/142430876?v=4","primary_language":"Python","stars":6965,"forks":546,"topics":["binary","decompile","large-language-models","reverse-engineering"],"archived":false,"github_pushed_at":"2026-02-12T03:02:03+00:00","maintenance_label":"Slowing","stars_delta_30d":205,"url":"https://www.graphcanon.com/tools/albertan017-llm4decompile","markdown_url":"https://www.graphcanon.com/tools/albertan017-llm4decompile.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/albertan017-llm4decompile","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=albertan017-llm4decompile"}},{"type":"integrates_with","direction":"in","explanation":"ScrapeGraph-AI is focused on integrating AI for web scraping tasks, which could efficiently work with Data-Juicer's data processing capabilities to transform and analyze the scraped data.","successor_context":null,"tool":{"slug":"datajuicer-data-juicer","name":"data-juicer","tagline":"Data processing for and with foundation models","github_url":"https://github.com/datajuicer/data-juicer","owner":"datajuicer","repo":"data-juicer","owner_avatar_url":"https://avatars.githubusercontent.com/u/223222708?v=4","primary_language":"Python","stars":6897,"forks":404,"topics":["data","data-analysis","data-pipeline","data-processing","data-science","data-visualization","foundation-models","instruction-tuning","large-language-models","llm","llms","multi-modal","pre-training","synthetic-data"],"archived":false,"github_pushed_at":"2026-08-13T09:19:31+00:00","maintenance_label":"Very active","stars_delta_30d":166,"url":"https://www.graphcanon.com/tools/datajuicer-data-juicer","markdown_url":"https://www.graphcanon.com/tools/datajuicer-data-juicer.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/datajuicer-data-juicer","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=datajuicer-data-juicer"}},{"type":"alternative","direction":"in","explanation":"Both tools are designed for web data extraction, but Trafilatura focuses on text and metadata extraction using a more traditional approach, while Scrapegraph-AI leverages AI techniques.","successor_context":null,"tool":{"slug":"adbar-trafilatura","name":"trafilatura","tagline":"Python & Command-line tool for web crawling, scraping and text extraction","github_url":"https://github.com/adbar/trafilatura","owner":"adbar","repo":"trafilatura","owner_avatar_url":"https://avatars.githubusercontent.com/u/2125866?v=4","primary_language":"Python","stars":6657,"forks":415,"topics":["article-extractor","corpus-builder","corpus-tools","crawler","html-to-markdown","html2text","llm","news-aggregator","news-crawler","nlp","rag","readability","rss-feed","scraping","tei","text-cleaning","text-extraction","text-mining","text-preprocessing","web-scraping"],"archived":false,"github_pushed_at":"2026-08-15T16:13:21+00:00","maintenance_label":"Very active","stars_delta_30d":343,"url":"https://www.graphcanon.com/tools/adbar-trafilatura","markdown_url":"https://www.graphcanon.com/tools/adbar-trafilatura.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/adbar-trafilatura","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=adbar-trafilatura"}},{"type":"related","direction":"in","explanation":"While both tools involve scraping or extracting information, LangExtract focuses on processing unstructured text into structured data, whereas Scrapegraph-ai extracts web data with integrated LLM capabilities.","successor_context":null,"tool":{"slug":"google-langextract","name":"langextract","tagline":"A Python library for extracting structured information from unstructured text using LLMs.","github_url":"https://github.com/google/langextract","owner":"google","repo":"langextract","owner_avatar_url":"https://avatars.githubusercontent.com/u/1342004?v=4","primary_language":"Python","stars":38400,"forks":2693,"topics":["gemini","gemini-ai","gemini-api","gemini-flash","gemini-pro","information-extration","large-language-models","llm","nlp","python","structured-data"],"archived":false,"github_pushed_at":"2026-08-11T15:31:39+00:00","maintenance_label":"Very active","stars_delta_30d":1241,"url":"https://www.graphcanon.com/tools/google-langextract","markdown_url":"https://www.graphcanon.com/tools/google-langextract.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/google-langextract","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=google-langextract"}},{"type":"integrates_with","direction":"in","explanation":"GPT Researcher could integrate with Scrapegraph AI for enhanced web scraping and data extraction capabilities, improving the depth of research through richer web data acquisition.","successor_context":null,"tool":{"slug":"assafelovic-gpt-researcher","name":"gpt-researcher","tagline":"An autonomous agent that conducts deep research using LLM providers","github_url":"https://github.com/assafelovic/gpt-researcher","owner":"assafelovic","repo":"gpt-researcher","owner_avatar_url":"https://avatars.githubusercontent.com/u/13554167?v=4","primary_language":"Python","stars":28883,"forks":3916,"topics":["agent","ai","automation","deepresearch","llms","mcp","mcp-server","python","research","search","webscraping"],"archived":false,"github_pushed_at":"2026-07-18T07:56:55+00:00","maintenance_label":"Active","stars_delta_30d":737,"url":"https://www.graphcanon.com/tools/assafelovic-gpt-researcher","markdown_url":"https://www.graphcanon.com/tools/assafelovic-gpt-researcher.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/assafelovic-gpt-researcher","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=assafelovic-gpt-researcher"}},{"type":"alternative","direction":"in","explanation":"Both ChatWeb and Scrapegraph-AI are focused on web data extraction for LLM integration, although their methodologies might differ slightly.","successor_context":null,"tool":{"slug":"skywalkerdarren-chatweb","name":"chatWeb","tagline":"ChatWeb can crawl web pages and various document types for content extraction and summarization.","github_url":"https://github.com/SkywalkerDarren/chatWeb","owner":"SkywalkerDarren","repo":"chatWeb","owner_avatar_url":"https://avatars.githubusercontent.com/u/20706299?v=4","primary_language":"Python","stars":916,"forks":137,"topics":["ai","chatgpt","crawler","docx","embedding","faiss","gpt","gpt-35-turbo","news-extractor","newspaper","openai","pdf","pgvector","postgresql","vector-database"],"archived":false,"github_pushed_at":"2026-05-25T16:56:25+00:00","maintenance_label":"Steady","stars_delta_30d":0,"url":"https://www.graphcanon.com/tools/skywalkerdarren-chatweb","markdown_url":"https://www.graphcanon.com/tools/skywalkerdarren-chatweb.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/skywalkerdarren-chatweb","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=skywalkerdarren-chatweb"}},{"type":"alternative","direction":"in","explanation":"'Scrapegraph-ai' and 'Scrapling' both offer AI-powered web scraping; 'Scrapling' emphasizes an adaptive framework for crawling, setting it apart from the Python-based scraper of 'Scrapegraph-ai'.","successor_context":null,"tool":{"slug":"d4vinci-scrapling","name":"Scrapling","tagline":"An adaptive Web Scraping framework","github_url":"https://github.com/D4Vinci/Scrapling","owner":"D4Vinci","repo":"Scrapling","owner_avatar_url":"https://avatars.githubusercontent.com/u/20604835?v=4","primary_language":"Python","stars":71247,"forks":7067,"topics":["ai","ai-scraping","automation","crawler","crawling","crawling-python","data","data-extraction","mcp","mcp-server","playwright","python","scraping","selectors","stealth","web-scraper","web-scraping","web-scraping-python","webscraping","xpath"],"archived":false,"github_pushed_at":"2026-07-25T16:07:12+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/d4vinci-scrapling","markdown_url":"https://www.graphcanon.com/tools/d4vinci-scrapling.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/d4vinci-scrapling","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=d4vinci-scrapling"}}],"neighbours":[{"slug":"graphify-labs-graphify","name":"graphify","tagline":"Turn any code or documentation into a queryable knowledge graph","github_url":"https://github.com/Graphify-Labs/graphify","owner":"Graphify-Labs","repo":"graphify","owner_avatar_url":"https://avatars.githubusercontent.com/u/297659074?v=4","primary_language":"Python","stars":107507,"forks":10441,"topics":["ai-agents","antigravity","ast","claude-code","code-analysis","code-search","codex","cursor","developer-tools","gemini","graphrag","knowledge-graph","leiden","llm","mcp","openclaw","rag","skills","tree-sitter"],"archived":false,"github_pushed_at":"2026-08-17T18:42:58+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/graphify-labs-graphify","markdown_url":"https://www.graphcanon.com/tools/graphify-labs-graphify.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/graphify-labs-graphify","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=graphify-labs-graphify","shared_categories":["data-retrieval"]},{"slug":"d4vinci-scrapling","name":"Scrapling","tagline":"An adaptive Web Scraping framework","github_url":"https://github.com/D4Vinci/Scrapling","owner":"D4Vinci","repo":"Scrapling","owner_avatar_url":"https://avatars.githubusercontent.com/u/20604835?v=4","primary_language":"Python","stars":71247,"forks":7067,"topics":["ai","ai-scraping","automation","crawler","crawling","crawling-python","data","data-extraction","mcp","mcp-server","playwright","python","scraping","selectors","stealth","web-scraper","web-scraping","web-scraping-python","webscraping","xpath"],"archived":false,"github_pushed_at":"2026-07-25T16:07:12+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/d4vinci-scrapling","markdown_url":"https://www.graphcanon.com/tools/d4vinci-scrapling.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/d4vinci-scrapling","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=d4vinci-scrapling","shared_categories":["data-retrieval"]},{"slug":"pathwaycom-llm-app","name":"llm-app","tagline":"Ready-to-run cloud templates for RAG, AI pipelines, and enterprise search with live data.","github_url":"https://github.com/pathwaycom/llm-app","owner":"pathwaycom","repo":"llm-app","owner_avatar_url":"https://avatars.githubusercontent.com/u/25750857?v=4","primary_language":"Jupyter Notebook","stars":59037,"forks":1466,"topics":["chatbot","hugging-face","llm","llm-local","llm-prompting","llm-security","llmops","machine-learning","open-ai","pathway","rag","real-time","retrieval-augmented-generation","vector-database","vector-index"],"archived":false,"github_pushed_at":"2026-07-05T17:59:07+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/pathwaycom-llm-app","markdown_url":"https://www.graphcanon.com/tools/pathwaycom-llm-app.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/pathwaycom-llm-app","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=pathwaycom-llm-app","shared_categories":["data-retrieval","llm-frameworks"]},{"slug":"opendataloader-project-opendataloader-pdf","name":"opendataloader-pdf","tagline":"PDF Parser for AI-ready data","github_url":"https://github.com/opendataloader-project/opendataloader-pdf","owner":"opendataloader-project","repo":"opendataloader-pdf","owner_avatar_url":"https://avatars.githubusercontent.com/u/211280852?v=4","primary_language":"Java","stars":28528,"forks":2724,"topics":["a11y","accessibility","ai","bounding-box","document-parsing","eaa","html","json","markdown","ocr","ocr-recognition","pdf","pdf-accessibility","pdf-converter","pdf-extraction","pdf-parser","pdf-ua","rag","tables","tagged-pdf"],"archived":false,"github_pushed_at":"2026-08-18T04:01:45+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/opendataloader-project-opendataloader-pdf","markdown_url":"https://www.graphcanon.com/tools/opendataloader-project-opendataloader-pdf.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/opendataloader-project-opendataloader-pdf","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=opendataloader-project-opendataloader-pdf","shared_categories":["data-retrieval"]},{"slug":"tencent-weknora","name":"WeKnora","tagline":"Open-source LLM knowledge platform for creating a queryable RAG, autonomous reasoning agent, and self-maintaining Wiki.","github_url":"https://github.com/Tencent/WeKnora","owner":"Tencent","repo":"WeKnora","owner_avatar_url":"https://avatars.githubusercontent.com/u/18461506?v=4","primary_language":"Go","stars":19992,"forks":2877,"topics":["agent","agentic","ai","chatbot","embeddings","evaluation","generative-ai","golang","knowledge-base","llm","multi-tenant","multimodel","ollama","openai","question-answering","rag","reranking","semantic-search","vector-search","wiki"],"archived":false,"github_pushed_at":"2026-08-16T05:31:52+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/tencent-weknora","markdown_url":"https://www.graphcanon.com/tools/tencent-weknora.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/tencent-weknora","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=tencent-weknora","shared_categories":["llm-frameworks"]},{"slug":"yusufkaraaslan-skill-seekers","name":"Skill_Seekers","tagline":"Automation tool for converting documentation and code into Claude AI skills","github_url":"https://github.com/yusufkaraaslan/Skill_Seekers","owner":"yusufkaraaslan","repo":"Skill_Seekers","owner_avatar_url":"https://avatars.githubusercontent.com/u/11597362?v=4","primary_language":"Python","stars":14565,"forks":1477,"topics":["ai-tools","ast-parser","automation","claude-ai","claude-skills","code-analysis","conflict-detection","documentation","documentation-generator","github","github-scraper","mcp","mcp-server","multi-source","ocr","pdf","python","web-scraping"],"archived":false,"github_pushed_at":"2026-07-20T13:09:44+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/yusufkaraaslan-skill-seekers","markdown_url":"https://www.graphcanon.com/tools/yusufkaraaslan-skill-seekers.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/yusufkaraaslan-skill-seekers","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=yusufkaraaslan-skill-seekers","shared_categories":[]},{"slug":"ntegrals-openbrowser","name":"openbrowser","tagline":"AI-powered autonomous web browsing framework","github_url":"https://github.com/ntegrals/openbrowser","owner":"ntegrals","repo":"openbrowser","owner_avatar_url":"https://avatars.githubusercontent.com/u/26648900?v=4","primary_language":"TypeScript","stars":9510,"forks":864,"topics":["ai-agents","automation","claude","playwright","puppeteer","sandbox"],"archived":false,"github_pushed_at":"2026-04-02T12:55:42+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/ntegrals-openbrowser","markdown_url":"https://www.graphcanon.com/tools/ntegrals-openbrowser.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ntegrals-openbrowser","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ntegrals-openbrowser","shared_categories":[]},{"slug":"alirezamika-autoscraper","name":"autoscraper","tagline":"A Smart Automatic Fast and Lightweight Web Scraper for Python","github_url":"https://github.com/alirezamika/autoscraper","owner":"alirezamika","repo":"autoscraper","owner_avatar_url":"https://avatars.githubusercontent.com/u/17881612?v=4","primary_language":"Python","stars":7839,"forks":808,"topics":["ai","artificial-intelligence","automation","crawler","machine-learning","python","scrape","scraper","scraping","web-scraping","webautomation","webscraping"],"archived":false,"github_pushed_at":"2026-07-29T17:17:55+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/alirezamika-autoscraper","markdown_url":"https://www.graphcanon.com/tools/alirezamika-autoscraper.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/alirezamika-autoscraper","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=alirezamika-autoscraper","shared_categories":[]},{"slug":"firecrawl-firecrawl-mcp-server","name":"firecrawl-mcp-server","tagline":"Adding web scraping and search capabilities to LLM clients like Cursor and Claude","github_url":"https://github.com/firecrawl/firecrawl-mcp-server","owner":"firecrawl","repo":"firecrawl-mcp-server","owner_avatar_url":"https://avatars.githubusercontent.com/u/135057108?v=4","primary_language":"JavaScript","stars":7047,"forks":821,"topics":["batch-processing","claude","content-extraction","data-collection","firecrawl","firecrawl-ai","javascript-rendering","llm-tools","mcp","mcp-server","model-context-protocol","search-api","web-crawler","web-scraping"],"archived":false,"github_pushed_at":"2026-07-26T06:47:42+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/firecrawl-firecrawl-mcp-server","markdown_url":"https://www.graphcanon.com/tools/firecrawl-firecrawl-mcp-server.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/firecrawl-firecrawl-mcp-server","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=firecrawl-firecrawl-mcp-server","shared_categories":["data-retrieval","llm-frameworks"]},{"slug":"zipstack-unstract","name":"unstract","tagline":"LLM-Driven Extraction of Unstructured Data for API Deployments and ETL Pipeline Workflows","github_url":"https://github.com/Zipstack/unstract","owner":"Zipstack","repo":"unstract","owner_avatar_url":"https://avatars.githubusercontent.com/u/89070934?v=4","primary_language":"Python","stars":6932,"forks":663,"topics":["ai-agents","data-engineering","document-ai","generative-ai","idp","json-extraction","llm","mcp-server","ocr","pdf-extraction","prompt-engineering","structured-output"],"archived":false,"github_pushed_at":"2026-07-27T22:23:41+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/zipstack-unstract","markdown_url":"https://www.graphcanon.com/tools/zipstack-unstract.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/zipstack-unstract","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=zipstack-unstract","shared_categories":["data-retrieval","llm-frameworks"]},{"slug":"exa-labs-exa-mcp-server","name":"exa-mcp-server","tagline":"Platform for web search and crawling using Model Context Protocol.","github_url":"https://github.com/exa-labs/exa-mcp-server","owner":"exa-labs","repo":"exa-mcp-server","owner_avatar_url":"https://avatars.githubusercontent.com/u/77906174?v=4","primary_language":"TypeScript","stars":4777,"forks":362,"topics":["code-search","codesearch","crawling","mcp","mcp-server","model-context-protocol","web-search","websearch"],"archived":false,"github_pushed_at":"2026-07-24T17:07:01+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/exa-labs-exa-mcp-server","markdown_url":"https://www.graphcanon.com/tools/exa-labs-exa-mcp-server.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/exa-labs-exa-mcp-server","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=exa-labs-exa-mcp-server","shared_categories":["data-retrieval"]},{"slug":"agentset-ai-agentset","name":"agentset","tagline":"The open-source RAG platform with built-in citations and support for deep research","github_url":"https://github.com/agentset-ai/agentset","owner":"agentset-ai","repo":"agentset","owner_avatar_url":"https://avatars.githubusercontent.com/u/200139246?v=4","primary_language":"TypeScript","stars":2066,"forks":185,"topics":["agentic-rag","ai","ai-agents","ai-sdk","chatbots","embeddings","genai","llms","memory","memory-management","rag","vercel-ai-sdk"],"archived":false,"github_pushed_at":"2026-07-16T13:11:34+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/agentset-ai-agentset","markdown_url":"https://www.graphcanon.com/tools/agentset-ai-agentset.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/agentset-ai-agentset","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=agentset-ai-agentset","shared_categories":["data-retrieval"]}]}}