Reddit Subreddit Scraper V2
Pricing
from $1.10 / 1,000 results
Reddit Subreddit Scraper V2
Reddit Subreddit Scraper V2 — cached data source (results may lag up to ~1.5h). Structured JSON, no login, no API key. For real-time data use the V1 actor.
Reddit Subreddit Scraper V2
Pricing
from $1.10 / 1,000 results
Reddit Subreddit Scraper V2 — cached data source (results may lag up to ~1.5h). Structured JSON, no login, no API key. For real-time data use the V1 actor.
You can access the Reddit Subreddit Scraper V2 programmatically from your own applications by using the Apify API. You can also choose the language preference from below. To use the Apify API, you’ll need an Apify account and your API token, found in Integrations settings in Apify Console.
{ "openapi": "3.0.1", "info": { "version": "0.1", "x-build-id": "tEYCK3wO5SqoBSdmC" }, "servers": [ { "url": "https://api.apify.com/v2" } ], "paths": { "/acts/myagizm~reddit-subreddit-scraper-v2/run-sync-get-dataset-items": { "post": { "operationId": "run-sync-get-dataset-items-myagizm-reddit-subreddit-scraper-v2", "x-openai-isConsequential": false, "summary": "Executes an Actor, waits for its completion, and returns Actor's dataset items in response.", "tags": [ "Run Actor" ], "requestBody": { "required": true, "content": { "application/json": { "schema": { "$ref": "#/components/schemas/inputSchema" } } } }, "parameters": [ { "name": "token", "in": "query", "required": true, "schema": { "type": "string" }, "description": "Enter your Apify token here" } ], "responses": { "200": { "description": "OK" } } } }, "/acts/myagizm~reddit-subreddit-scraper-v2/runs": { "post": { "operationId": "runs-sync-myagizm-reddit-subreddit-scraper-v2", "x-openai-isConsequential": false, "summary": "Executes an Actor and returns information about the initiated run in response.", "tags": [ "Run Actor" ], "requestBody": { "required": true, "content": { "application/json": { "schema": { "$ref": "#/components/schemas/inputSchema" } } } }, "parameters": [ { "name": "token", "in": "query", "required": true, "schema": { "type": "string" }, "description": "Enter your Apify token here" } ], "responses": { "200": { "description": "OK", "content": { "application/json": { "schema": { "$ref": "#/components/schemas/runsResponseSchema" } } } } } } }, "/acts/myagizm~reddit-subreddit-scraper-v2/run-sync": { "post": { "operationId": "run-sync-myagizm-reddit-subreddit-scraper-v2", "x-openai-isConsequential": false, "summary": "Executes an Actor, waits for completion, and returns the OUTPUT from Key-value store in response.", "tags": [ "Run Actor" ], "requestBody": { "required": true, "content": { "application/json": { "schema": { "$ref": "#/components/schemas/inputSchema" } } } }, "parameters": [ { "name": "token", "in": "query", "required": true, "schema": { "type": "string" }, "description": "Enter your Apify token here" } ], "responses": { "200": { "description": "OK" } } } } }, "components": { "schemas": { "inputSchema": { "type": "object", "properties": { "subreddits": { "title": "🎯 Subreddits", "type": "array", "description": "Subreddit names to scrape posts from, e.g. \"AskReddit\", \"news\". The r/ prefix is optional (both \"forhire\" and \"r/forhire\" work).", "items": { "type": "string" } }, "sort": { "title": "↕️ Sort", "enum": [ "new", "hot", "top", "rising", "controversial" ], "type": "string", "description": "Listing order for the subreddit feed. For example, choose \"top\" to get the highest-scoring posts, or \"new\" for the most recent.", "default": "hot" }, "time": { "title": "📅 Time range", "enum": [ "hour", "day", "week", "month", "year", "all" ], "type": "string", "description": "Time window applied to Top / Controversial sorts. For example, pick \"week\" to get the top posts of the past 7 days.", "default": "all" }, "includeSelftext": { "title": "💬 Include post text", "type": "boolean", "description": "Fetch each post's self/body text. For example, enable this to capture the full \"[FOR HIRE]\" description on r/forhire posts for lead generation.", "default": true }, "includeMedia": { "title": "🖼️ $ Include media URLs", "type": "boolean", "description": "Extract image/video URLs (i.redd.it, v.redd.it, imgur, galleries). For example, enable this when scraping r/pics to collect every image link alongside the post.", "default": false }, "includeNSFW": { "title": "🔞 Include NSFW", "type": "boolean", "description": "Include posts flagged over-18 / NSFW. For example, leave this off to keep the dataset safe-for-work, or turn it on for complete community coverage.", "default": false }, "postDateLimit": { "title": "📆 Only posts on/after (YYYY-MM-DD)", "type": "string", "description": "Keep only posts created on or after this date. Leave empty for no limit. For example, \"2026-01-01\" keeps only posts from 2026 onward." }, "searchTerms": { "title": "🔎 Keyword filter", "type": "array", "description": "Keep only posts whose title or text contains at least one of these words. For example, [\"hiring\", \"developer\"] keeps only posts mentioning either word.", "items": { "type": "string" } }, "maxItems": { "title": "💯 Max items", "minimum": 1, "maximum": 100, "type": "integer", "description": "Hard cap on the total posts saved across the whole run. For example, set 50 for a quick monitoring pull. A run lasts at most ~10 minutes.", "default": 50 }, "maxPostsPerSource": { "title": "💯 Max posts per subreddit", "minimum": 1, "maximum": 100, "type": "integer", "description": "Cap on posts kept from each subreddit, so one busy community can't dominate. For example, set 25 when scraping 4 subreddits to get a balanced 100-post dataset.", "default": 25 }, "maxPages": { "title": "📄 Max listing pages", "minimum": 1, "maximum": 20, "type": "integer", "description": "Pagination depth per subreddit (25 posts/page). For example, set 5 to reach ~125 posts back into a subreddit's feed.", "default": 1 }, "maxRetries": { "title": "🔁 Max request retries", "minimum": 0, "maximum": 6, "type": "integer", "description": "Retry attempts per failed request. For example, raise to 5 for flaky network conditions; lower to 0 for the fastest possible run.", "default": 3 }, "requestDelayMs": { "title": "⏱️ Delay between requests (ms)", "minimum": 0, "maximum": 10000, "type": "integer", "description": "Politeness delay between requests, in milliseconds. For example, 300 is a good default; raise to 1000 to be gentler on large jobs.", "default": 300 }, "debug": { "title": "🐞 Debug logging", "type": "boolean", "description": "Emit verbose logs. For example, enable this to diagnose why a specific subreddit returns no posts.", "default": false } } }, "runsResponseSchema": { "type": "object", "properties": { "data": { "type": "object", "properties": { "id": { "type": "string" }, "actId": { "type": "string" }, "userId": { "type": "string" }, "startedAt": { "type": "string", "format": "date-time", "example": "2025-01-08T00:00:00.000Z" }, "finishedAt": { "type": "string", "format": "date-time", "example": "2025-01-08T00:00:00.000Z" }, "status": { "type": "string", "example": "READY" }, "meta": { "type": "object", "properties": { "origin": { "type": "string", "example": "API" }, "userAgent": { "type": "string" } } }, "stats": { "type": "object", "properties": { "inputBodyLen": { "type": "integer", "example": 2000 }, "rebootCount": { "type": "integer", "example": 0 }, "restartCount": { "type": "integer", "example": 0 }, "resurrectCount": { "type": "integer", "example": 0 }, "computeUnits": { "type": "integer", "example": 0 } } }, "options": { "type": "object", "properties": { "build": { "type": "string", "example": "latest" }, "timeoutSecs": { "type": "integer", "example": 300 }, "memoryMbytes": { "type": "integer", "example": 1024 }, "diskMbytes": { "type": "integer", "example": 2048 } } }, "buildId": { "type": "string" }, "defaultKeyValueStoreId": { "type": "string" }, "defaultDatasetId": { "type": "string" }, "defaultRequestQueueId": { "type": "string" }, "buildNumber": { "type": "string", "example": "1.0.0" }, "containerUrl": { "type": "string" }, "usage": { "type": "object", "properties": { "ACTOR_COMPUTE_UNITS": { "type": "integer", "example": 0 }, "DATASET_READS": { "type": "integer", "example": 0 }, "DATASET_WRITES": { "type": "integer", "example": 0 }, "KEY_VALUE_STORE_READS": { "type": "integer", "example": 0 }, "KEY_VALUE_STORE_WRITES": { "type": "integer", "example": 1 }, "KEY_VALUE_STORE_LISTS": { "type": "integer", "example": 0 }, "REQUEST_QUEUE_READS": { "type": "integer", "example": 0 }, "REQUEST_QUEUE_WRITES": { "type": "integer", "example": 0 }, "DATA_TRANSFER_INTERNAL_GBYTES": { "type": "integer", "example": 0 }, "DATA_TRANSFER_EXTERNAL_GBYTES": { "type": "integer", "example": 0 }, "PROXY_RESIDENTIAL_TRANSFER_GBYTES": { "type": "integer", "example": 0 }, "PROXY_SERPS": { "type": "integer", "example": 0 } } }, "usageTotalUsd": { "type": "number", "example": 0.00005 }, "usageUsd": { "type": "object", "properties": { "ACTOR_COMPUTE_UNITS": { "type": "integer", "example": 0 }, "DATASET_READS": { "type": "integer", "example": 0 }, "DATASET_WRITES": { "type": "integer", "example": 0 }, "KEY_VALUE_STORE_READS": { "type": "integer", "example": 0 }, "KEY_VALUE_STORE_WRITES": { "type": "number", "example": 0.00005 }, "KEY_VALUE_STORE_LISTS": { "type": "integer", "example": 0 }, "REQUEST_QUEUE_READS": { "type": "integer", "example": 0 }, "REQUEST_QUEUE_WRITES": { "type": "integer", "example": 0 }, "DATA_TRANSFER_INTERNAL_GBYTES": { "type": "integer", "example": 0 }, "DATA_TRANSFER_EXTERNAL_GBYTES": { "type": "integer", "example": 0 }, "PROXY_RESIDENTIAL_TRANSFER_GBYTES": { "type": "integer", "example": 0 }, "PROXY_SERPS": { "type": "integer", "example": 0 } } } } } } } } }}OpenAPI is a standard for designing and describing RESTful APIs, allowing developers to define API structure, endpoints, and data formats in a machine-readable way. It simplifies API development, integration, and documentation.
OpenAPI is effective when used with AI agents and GPTs by standardizing how these systems interact with various APIs, for reliable integrations and efficient communication.
By defining machine-readable API specifications, OpenAPI allows AI models like GPTs to understand and use varied data sources, improving accuracy. This accelerates development, reduces errors, and provides context-aware responses, making OpenAPI a core component for AI applications.
You can download the OpenAPI definitions for Reddit Subreddit Scraper V2 from the options below:
If you’d like to learn more about how OpenAPI powers GPTs, read our blog post.
You can also check out our other API clients: