{"id":5067,"external_id":"plugins_6ab424ad5b448191b8fede2ebf569333","name":"production-data-engineering-copilot","display_name":"Data Engineering Copilot","developer":"Krishna Sathvik","category":"Data & Analytics","listing_language":"en","listing_language_details":{"method":"cld-0.13.0/listing-v1","reliable":true,"detected_at":"2026-10-01T13:22:06Z","input_sha256":"040ba7164a186333d99b9f1b1fb6e11ebf2246478bdf8af495fb2bbd3bd8f4ad","source_fields":["release.description","release.interface.short_description","release.interface.long_description"]},"version":"0.1.0","skill_count":1,"first_seen_at":"2026-09-30T22:02:35.000Z","last_seen_at":"2026-10-02T06:00:01.935Z","last_changed_at":"2026-09-30T22:02:35.000Z","install_count":null,"research_summary":null,"research_reviewed_at":null,"metadata":{"id":"plugins_6ab424ad5b448191b8fede2ebf569333","name":"production-data-engineering-copilot","scope":"GLOBAL","status":"ENABLED","release":{"id":"pluginrel_d255d856c2248191bf1362c3046b509c","skills":[{"name":"data-engineering","interface":{"brand_color":null,"iconography":"chart","display_name":"Data Engineering Copilot","default_prompt":"Debug why this production pipeline is failing or producing wrong data.","icon_large_url":null,"icon_small_url":"https://files.openai.com/content?id=file_0000000000c081f598f2dfb37f7328b4","short_description":"Debug production data systems"},"description":"Production-focused workflow for designing, debugging, recovering, tuning, and operating data pipelines across Spark, Databricks, Delta Lake, SQL, CDC, Structured Streaming, Airflow, backfills, quality, observability, and lakehouse systems.","plugin_release_skill_id":"pluginrsk_6ab424afaf608191b4c67ed67116dcfc"}],"app_ids":[],"version":"0.1.0","keywords":["data-engineering","spark","databricks","delta-lake","cdc","streaming","airflow","data-quality"],"interface":{"category":"Data & Analytics","logo_url":"https://files.openai.com/content?id=file_00000000524481f890727a5fa01bd762","brand_color":null,"website_url":null,"capabilities":["Design reliable batch, CDC, streaming, and lakehouse architectures","Debug failed, slow, duplicated, stale, or backlogged production pipelines","Plan safe replay, backfill, recovery, and checkpoint strategies","Review Delta MERGE logic for keys, ordering, duplicates, updates, and deletes","Diagnose Spark and Databricks performance using plans, shuffle, spill, skew, and file evidence","Design Structured Streaming state, watermark, checkpoint, and sink behavior","Create production-oriented SQL, PySpark, Databricks, and Airflow patterns","Build data-quality gates, reconciliation, freshness checks, and operational runbooks","Evaluate retry safety, idempotency, schema evolution, and downstream contracts","Use current official docs for Spark, Databricks, Delta Lake, Airflow, and runtime-specific behavior"],"logo_url_dark":null,"default_prompt":"Debug why this production pipeline is failing or producing wrong data.","developer_name":"Krishna Sathvik","default_prompts":["Debug why this production pipeline is failing or producing wrong data.","Design a reliable CDC or streaming pipeline for this use case.","Plan a safe backfill or replay without duplicating or losing data."],"screenshot_urls":[],"long_description":"Data Engineering Copilot helps you design, debug, recover, and improve production data systems across Spark, Databricks, Delta Lake, SQL, CDC, Structured Streaming, Airflow, backfills, data quality, observability, and lakehouse architectures. It can diagnose failed or slow pipelines, reason about duplicate and late data, plan safe replay and backfill strategies, review checkpoint and streaming-state risks, tune Spark from execution evidence, and create production-oriented SQL, PySpark, and orchestration patterns. It prioritizes correctness, recoverability, SLA, security, and cost before scaling or redesigning.","composer_icon_url":"https://files.openai.com/content?id=file_0000000010dc81fb9a972b5977355f9a","short_description":"Debug production data systems","plugin_category_id":"data & analytics","privacy_policy_url":null,"terms_of_service_url":null,"composer_icon_dark_url":null},"description":"Production data engineering architecture, debugging, recovery, and reliability workflows.","app_manifest":null,"display_name":"Data Engineering Copilot","app_templates":[],"onboarding_skill_name":null,"requires_local_executor":false},"created_at":"2026-09-23T19:17:18.763916Z","is_template":false,"connector_id":null,"discoverability":"UNLISTED","canonical_app_id":null},"research":null,"package_metadata":{"name":"production-data-engineering-copilot","author":{"name":"Krishna Sathvik"},"sources":[{"path":"plugin.json","sha256":"07f640776e78dcbbc323c88a7fdd8b52613bd2cbaf0398cc4cab784b39a7f600"},{"path":".codex-plugin/plugin.json","sha256":"2a008086fd4646e97fe61742f05b39e721679ddc6d230399917b9d100fe589b7"}],"version":"0.1.0","keywords":["data-engineering","spark","databricks","delta-lake","cdc","streaming","airflow","data-quality"],"artifact_id":11915,"observed_at":"2026-10-02T00:37:04Z","capabilities":["Design reliable batch, CDC, streaming, and lakehouse architectures","Debug failed, slow, duplicated, stale, or backlogged production pipelines","Plan safe replay, backfill, recovery, and checkpoint strategies","Review Delta MERGE logic for keys, ordering, duplicates, updates, and deletes","Diagnose Spark and Databricks performance using plans, shuffle, spill, skew, and file evidence","Design Structured Streaming state, watermark, checkpoint, and sink behavior","Create production-oriented SQL, PySpark, Databricks, and Airflow patterns","Build data-quality gates, reconciliation, freshness checks, and operational runbooks","Evaluate retry safety, idempotency, schema evolution, and downstream contracts","Use current official docs for Spark, Databricks, Delta Lake, Airflow, and runtime-specific behavior"],"field_sources":{"name":0,"author":0,"version":0,"keywords":0,"capabilities":0},"extraction_version":1,"conflicts_or_errors":[]}}