{"users":[{"id":-1,"username":"system","name":"system","avatar_template":"/user_avatar/community.sparkflows.ai/system/{size}/6_2.png","admin":true,"moderator":true,"trust_level":4},{"id":23,"username":"Shad","name":null,"avatar_template":"https://avatars.discourse-cdn.com/v4/letter/s/ecccb3/{size}.png","trust_level":1},{"id":16,"username":"neeraj","name":null,"avatar_template":"https://avatars.discourse-cdn.com/v4/letter/n/e47c2d/{size}.png","moderator":true,"trust_level":1},{"id":13,"username":"Daniel","name":"Daniel","avatar_template":"https://avatars.discourse-cdn.com/v4/letter/d/3ec8ea/{size}.png","trust_level":1},{"id":1,"username":"admin","name":"Admin","avatar_template":"https://avatars.discourse-cdn.com/v4/letter/a/3ec8ea/{size}.png","admin":true,"moderator":true,"trust_level":4}],"primary_groups":[],"flair_groups":[],"topic_list":{"can_create_topic":false,"filter":"latest","more_topics_url":"/c/general/4?page=1","per_page":30,"top_tags":[{"id":3,"name":"faq","slug":"faq"},{"id":9,"name":"troubleshooting","slug":"troubleshooting"}],"topics":[{"fancy_title":"Welcome to Sparkflows! :wave:","id":5,"title":"Welcome to Sparkflows! :wave:","slug":"welcome-to-sparkflows","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2025-12-06T21:06:33.301Z","last_posted_at":"2025-12-06T21:06:33.350Z","bumped":true,"bumped_at":"2025-12-06T21:06:33.350Z","archetype":"regular","unseen":false,"pinned":true,"unpinned":null,"excerpt":"Welcome! \nCreate an account to participate in discussions, personalize your profile, and connect with the community. \nOnce you’ve joined, you can: \n\nIntroduce yourself.\nFollow topics that interest you.\nStart discussions &hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"unicode_title":"Welcome to Sparkflows! 👋","tags":[],"tags_descriptions":{},"views":33,"like_count":0,"has_summary":false,"last_poster_username":"system","category_id":4,"op_like_count":0,"pinned_globally":true,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":-1,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"About the General category","id":3,"title":"About the General category","slug":"about-the-general-category","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2025-12-06T21:06:25.845Z","last_posted_at":null,"bumped":true,"bumped_at":"2025-12-06T21:06:25.845Z","archetype":"regular","unseen":false,"pinned":true,"unpinned":null,"excerpt":null,"visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":8,"like_count":0,"has_summary":false,"last_poster_username":"system","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":-1,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"Sparkflows Single Sign-On (SSO) login issue following the SSL certificate update","id":357,"title":"Sparkflows Single Sign-On (SSO) login issue following the SSL certificate update","slug":"sparkflows-single-sign-on-sso-login-issue-following-the-ssl-certificate-update","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-09-17T07:16:05.625Z","last_posted_at":"2026-09-17T07:16:05.662Z","bumped":true,"bumped_at":"2026-09-17T07:16:05.662Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"When you update the SSL certificate of your Keycloak/SSO Provider with a certificate whose CA isn’t present in the Java keystore, then the authentication silently fails without providing any message. \nIn order to solve t&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":4,"like_count":0,"has_summary":false,"last_poster_username":"Shad","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":23,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"Handling multiple Livy URLs for HA","id":355,"title":"Handling multiple Livy URLs for HA","slug":"handling-multiple-livy-urls-for-ha","posts_count":2,"reply_count":0,"highest_post_number":3,"image_url":null,"created_at":"2026-09-09T05:56:47.066Z","last_posted_at":"2026-09-09T06:10:29.041Z","bumped":true,"bumped_at":"2026-09-09T06:10:29.041Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"We have deployed Livy services in a High Availability (HA) setup. For example, in the Prod environment, Livy is available on two nodes, with one node active and the other in standby mode. \nCurrently, Sparkflows allows on&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":6,"like_count":0,"has_summary":false,"last_poster_username":"Shad","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":23,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"SparkML workflow with xgboost are failing","id":354,"title":"SparkML workflow with xgboost are failing","slug":"sparkml-workflow-with-xgboost-are-failing","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-08-07T11:41:13.429Z","last_posted_at":"2026-08-07T11:41:13.501Z","bumped":true,"bumped_at":"2026-08-07T11:41:13.501Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Problem\nSparkML workflows using XGBoost may fail with the following error: \njava.lang.Exception: Error: Failed to execute the workflow:\n/tmp/libxgboost4j14342423119254343105.so: libgomp.so.1:\ncannot open shared object fi&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":10,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"H2O Model Save to S3 Fails Due to SSL Certificate Hostname Mismatch","id":353,"title":"H2O Model Save to S3 Fails Due to SSL Certificate Hostname Mismatch","slug":"h2o-model-save-to-s3-fails-due-to-ssl-certificate-hostname-mismatch","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-08-06T12:43:30.119Z","last_posted_at":"2026-08-06T12:43:30.182Z","bumped":true,"bumped_at":"2026-08-06T12:43:30.182Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Issue\nWhile saving an H2O model (MOJO) to an S3-compatible object store, the operation fails with the following exception: \norg.apache.hadoop.fs.s3a.AWSClientIOException:\ngetFileStatus on s3a://...\ncom.amazonaws.SdkClien&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":7,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"Job failed, when saving the model in S3","id":352,"title":"Job failed, when saving the model in S3","slug":"job-failed-when-saving-the-model-in-s3","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-08-06T12:39:23.597Z","last_posted_at":"2026-08-06T12:39:23.660Z","bumped":true,"bumped_at":"2026-08-06T12:39:23.660Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Following is the exception I received. \norg.apache.hadoop.fs.s3a.AWSClientIOException: getFileStatus on s3a://mlflow-artifacts/prod/All_Models_Test/h2o_S3_test/6449cf60-a992-423a-bec1-aaba2643891a: com.amazonaws.SdkClien&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":5,"like_count":0,"has_summary":false,"last_poster_username":"Shad","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":23,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"Sparkflows Installer [macOS] Fix: DMG Installer fails with &ldquo;UnsupportedClassVersionError&rdquo; (Homebrew Java 17)","id":350,"title":"Sparkflows Installer [macOS] Fix: DMG Installer fails with \"UnsupportedClassVersionError\" (Homebrew Java 17)","slug":"sparkflows-installer-macos-fix-dmg-installer-fails-with-unsupportedclassversionerror-homebrew-java-17","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":"https://canada1.discourse-cdn.com/flex007/uploads/sparkflows/optimized/1X/60d793c8f197acd53bf4bf8d380706e24678c2ff_2_1024x697.jpeg","created_at":"2026-08-03T08:45:40.432Z","last_posted_at":"2026-08-03T08:45:40.510Z","bumped":true,"bumped_at":"2026-08-03T08:45:40.510Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"The Problem: While trying to install the application via the macOS DMG, the installer failed during the database creation step. A popup error appeared with the following message: \nException in thread “main” java.lang.Uns&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":16,"like_count":0,"has_summary":false,"last_poster_username":"Daniel","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":13,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"H2O Sparkling Water IllegalAccessError on Livy","id":349,"title":"H2O Sparkling Water IllegalAccessError on Livy","slug":"h2o-sparkling-water-illegalaccesserror-on-livy","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-07-30T12:01:49.012Z","last_posted_at":"2026-07-30T12:01:49.135Z","bumped":true,"bumped_at":"2026-07-30T12:01:49.135Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Issue\nDuring workflow execution on Livy, the workflow may fail with the following exception: \nReceived an Exception:\n\njava.lang.IllegalAccessError: class ai.h2o.sparkling.backend.utils.RestCommunication\n(in unnamed modul&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":12,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"Troubleshooting: Expired Client Certificate in k3s.yaml","id":348,"title":"Troubleshooting: Expired Client Certificate in k3s.yaml","slug":"troubleshooting-expired-client-certificate-in-k3s-yaml","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-07-29T10:56:46.605Z","last_posted_at":"2026-07-29T10:56:46.692Z","bumped":true,"bumped_at":"2026-07-29T10:56:46.692Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Checking Kubernetes (k3s.yaml) Client Certificate Expiry\nThe k3s.yaml file contains the client certificate used by kubectl to authenticate with the Kubernetes API server. You can check its validity period using the comma&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":11,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"PySpark Workflow Submission Fails with SparkException: Exception thrown in awaitResult","id":347,"title":"PySpark Workflow Submission Fails with SparkException: Exception thrown in awaitResult","slug":"pyspark-workflow-submission-fails-with-sparkexception-exception-thrown-in-awaitresult","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-07-20T11:53:49.211Z","last_posted_at":"2026-07-20T11:53:49.283Z","bumped":true,"bumped_at":"2026-07-20T11:53:49.283Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Issue\nPySpark Workflow Submission Fails with SparkException: Exception thrown in awaitResult \nError \norg.apache.spark.SparkException: Exception thrown in awaitResult:\n    at org.apache.spark.util.SparkThreadUtils$.awaitR&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[{"id":9,"name":"troubleshooting","slug":"troubleshooting"}],"tags_descriptions":{},"views":12,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"Spark Cluster Workflow Execution: No Browser Results","id":346,"title":"Spark Cluster Workflow Execution: No Browser Results","slug":"spark-cluster-workflow-execution-no-browser-results","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-07-17T12:27:45.282Z","last_posted_at":"2026-07-17T12:27:45.352Z","bumped":true,"bumped_at":"2026-07-17T12:27:45.352Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Problem\nWhen executing workflows on a Spark cluster, the workflow completes or starts running, but no results are displayed in the Browser. \nCause\nSparkflows submits workflow jobs using spark-submit. The Spark driver is &hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[{"id":9,"name":"troubleshooting","slug":"troubleshooting"}],"tags_descriptions":{},"views":8,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"Pipeline Failure: Max retry reached Due to Dynamic Macro Usage","id":345,"title":"Pipeline Failure: Max retry reached Due to Dynamic Macro Usage","slug":"pipeline-failure-max-retry-reached-due-to-dynamic-macro-usage","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-07-16T11:20:49.104Z","last_posted_at":"2026-07-16T11:20:49.167Z","bumped":true,"bumped_at":"2026-07-16T11:20:49.167Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Issue\nPipeline execution fails with a Max retry reached error. \nCause\nWhen Sparkflows macros are used directly inside pipeline nodes, and the macro values change on every execution while the pipeline is scheduled to run &hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[{"id":9,"name":"troubleshooting","slug":"troubleshooting"}],"tags_descriptions":{},"views":8,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"ImportError: Unable to Import AzureOpenAI from OpenAI","id":344,"title":"ImportError: Unable to Import AzureOpenAI from OpenAI","slug":"importerror-unable-to-import-azureopenai-from-openai","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-07-13T10:53:56.980Z","last_posted_at":"2026-07-13T10:53:57.041Z","bumped":true,"bumped_at":"2026-07-13T10:53:57.041Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Error\nImportError: cannot import name &#39;AzureOpenAI&#39; from &#39;openai&#39;\n(/home/incorta/lib/python3.8/site-packages/openai/__init__.py)\n\nReason\nThe application is trying to import the AzureOpenAI client: \nfrom openai import Azu&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[{"id":9,"name":"troubleshooting","slug":"troubleshooting"}],"tags_descriptions":{},"views":11,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"EMR Cluster Job Failure: Unable to Locate ‘delta’ Data Source","id":343,"title":"EMR Cluster Job Failure: Unable to Locate ‘delta’ Data Source","slug":"emr-cluster-job-failure-unable-to-locate-delta-data-source","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-07-10T08:14:49.991Z","last_posted_at":"2026-07-10T08:14:50.059Z","bumped":true,"bumped_at":"2026-07-10T08:14:50.059Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Problem\nWhile running a job on the EMR cluster, the following error is encountered: \nFailed to find the delta data source \nThis indicates that the required Delta Lake JAR is either: \n\nMissing from the EMR cluster environ&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":8,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"Inconsistent Execution of Airflow Pipeline","id":342,"title":"Inconsistent Execution of Airflow Pipeline","slug":"inconsistent-execution-of-airflow-pipeline","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-07-06T10:34:24.597Z","last_posted_at":"2026-07-06T10:34:24.671Z","bumped":true,"bumped_at":"2026-07-06T10:34:24.671Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Problem:\nPipeline is not running consistently in Airflow. \nSolution:\nThere could be different reasons for this issue: \n\n\nThe Background Event Trigger Thread is stalled. \n\n\nThe Websocket doesn’t receive events in timely m&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":6,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"Pipeline Trigger Delay Due to Concurrent Scheduling","id":341,"title":"Pipeline Trigger Delay Due to Concurrent Scheduling","slug":"pipeline-trigger-delay-due-to-concurrent-scheduling","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-07-03T11:26:45.129Z","last_posted_at":"2026-07-03T11:26:45.190Z","bumped":true,"bumped_at":"2026-07-03T11:26:45.190Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Problem:\nScheduler takes a few minutes to trigger all the pipelines if there are too many of them triggered at the same time. \nSolution:\nThere could be different reasons for this issue like: \n\n\nThe RAM allocated to Spark&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":4,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"org.apache.spark.SparkClassNotFoundException- [DATA_SOURCE_NOT_FOUND] Failed to find the data source- delta, When Running Read Incorta Node on Chidori Cluster","id":340,"title":"org.apache.spark.SparkClassNotFoundException- [DATA_SOURCE_NOT_FOUND] Failed to find the data source- delta, When Running Read Incorta Node on Chidori Cluster","slug":"org-apache-spark-sparkclassnotfoundexception-data-source-not-found-failed-to-find-the-data-source-delta-when-running-read-incorta-node-on-chidori-cluster","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-07-03T04:27:36.347Z","last_posted_at":"2026-07-03T04:27:36.428Z","bumped":true,"bumped_at":"2026-07-03T04:27:36.428Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Problem\nWhile submitting a PySpark job on Chidori, the job fails with the following error: \norg.apache.spark.SparkClassNotFoundException: [DATA_SOURCE_NOT_FOUND]\nFailed to find the data source: delta.\n\nCause\nThe Delta La&hellip;","visible":true,"closed":true,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":23,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null},{"extras":null,"description":"Recent Poster","user_id":1,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"Workflow Not Getting Pushed to GitHub","id":339,"title":"Workflow Not Getting Pushed to GitHub","slug":"workflow-not-getting-pushed-to-github","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-07-02T09:11:19.882Z","last_posted_at":"2026-07-02T09:11:19.951Z","bumped":true,"bumped_at":"2026-07-02T09:11:19.951Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Problem\nThe user pushes project changes from the Workflow Editor, but the latest changes are not reflected in the GitHub repository. \nCause\nThe user does not have sufficient write permissions to push changes to the repos&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[{"id":9,"name":"troubleshooting","slug":"troubleshooting"}],"tags_descriptions":{},"views":10,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"ImportError: cannot import name &lsquo;ABCIndex&rsquo; from &lsquo;pandas.core.dtypes.generic&rsquo; While Submitting PySpark Job on Chidori","id":338,"title":"ImportError: cannot import name 'ABCIndex' from 'pandas.core.dtypes.generic' While Submitting PySpark Job on Chidori","slug":"importerror-cannot-import-name-abcindex-from-pandas-core-dtypes-generic-while-submitting-pyspark-job-on-chidori","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-07-02T05:42:59.647Z","last_posted_at":"2026-07-02T05:42:59.717Z","bumped":true,"bumped_at":"2026-07-02T05:42:59.717Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Reason\nCompatible version of Pandas is not available or the existing Pandas installation is corrupted/inconsistent, resulting in the following error: \nImportError: cannot import name &#39;ABCIndex&#39; from &#39;pandas.core.dtypes.g&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[{"id":9,"name":"troubleshooting","slug":"troubleshooting"}],"tags_descriptions":{},"views":16,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"java.lang.IllegalAccessError: ai.h2o.sparkling.backend.utils.RestCommunication When Running H2O Job","id":337,"title":"java.lang.IllegalAccessError: ai.h2o.sparkling.backend.utils.RestCommunication When Running H2O Job","slug":"java-lang-illegalaccesserror-ai-h2o-sparkling-backend-utils-restcommunication-when-running-h2o-job","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-06-19T11:11:30.662Z","last_posted_at":"2026-06-19T11:11:30.733Z","bumped":true,"bumped_at":"2026-06-19T11:11:30.733Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Issue:\nWhen running a H2O job on Livy, the following exception is encountered: \njava.lang.IllegalAccessError: class ai.h2o.sparkling.backend.utils.RestCommunication\n(in unnamed module @0x480846cb) cannot access class\nsun&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[{"id":9,"name":"troubleshooting","slug":"troubleshooting"}],"tags_descriptions":{},"views":19,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"What are the benefits of running Python workloads on EC2 through Sparkflows?","id":336,"title":"What are the benefits of running Python workloads on EC2 through Sparkflows?","slug":"what-are-the-benefits-of-running-python-workloads-on-ec2-through-sparkflows","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-06-08T07:08:27.747Z","last_posted_at":"2026-06-08T07:08:27.814Z","bumped":true,"bumped_at":"2026-06-08T07:08:27.814Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"This capability provides automated lifecycle management of EC2 resources, helping organizations optimize infrastructure usage and costs. By provisioning compute resources only when needed and automatically terminating th&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":11,"like_count":0,"has_summary":false,"last_poster_username":"Daniel","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":13,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"Can Sparkflows execute Python applications on Amazon EC2 instances as part of a pipeline?","id":335,"title":"Can Sparkflows execute Python applications on Amazon EC2 instances as part of a pipeline?","slug":"can-sparkflows-execute-python-applications-on-amazon-ec2-instances-as-part-of-a-pipeline","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-06-08T07:07:37.376Z","last_posted_at":"2026-06-08T07:07:37.435Z","bumped":true,"bumped_at":"2026-06-08T07:07:37.435Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Yes. Sparkflows enables end-to-end execution of Python applications on Amazon EC2 instances within a pipeline. Users can automatically provision and start an EC2 instance, execute Python scripts stored in Amazon S3 using&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":11,"like_count":0,"has_summary":false,"last_poster_username":"Daniel","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":13,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"How can I control the number of concurrent DAG runs in Sparkflows?","id":334,"title":"How can I control the number of concurrent DAG runs in Sparkflows?","slug":"how-can-i-control-the-number-of-concurrent-dag-runs-in-sparkflows","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":"https://canada1.discourse-cdn.com/flex007/uploads/sparkflows/optimized/1X/e5b6b9b60265f74d824d311a48824c3c857c14a1_2_1024x429.png","created_at":"2026-06-08T07:03:06.066Z","last_posted_at":"2026-06-08T07:03:06.133Z","bumped":true,"bumped_at":"2026-06-08T07:03:06.133Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Sparkflows now supports the max_active_runs parameter through the AddDagArguments node, allowing users to define the maximum number of DAG runs that can execute concurrently. The default value is empty, which preserves t&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":6,"like_count":0,"has_summary":false,"last_poster_username":"Daniel","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":13,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"When refreshing schema from S3 seeing PKIX path building failed: sun.security.provider.certpath.SunCertPathBuilderException: unable to find valid certification path to requested target","id":329,"title":"When refreshing schema from S3 seeing PKIX path building failed: sun.security.provider.certpath.SunCertPathBuilderException: unable to find valid certification path to requested target","slug":"when-refreshing-schema-from-s3-seeing-pkix-path-building-failed-sun-security-provider-certpath-suncertpathbuilderexception-unable-to-find-valid-certification-path-to-requested-target","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-06-04T11:19:39.852Z","last_posted_at":"2026-06-04T11:19:39.911Z","bumped":true,"bumped_at":"2026-06-04T11:19:39.911Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"While refreshing the schema from S3, the application was unable to establish a trusted SSL connection because certificate validation failed. The JVM could not find a valid certificate chain for the S3 endpoint, resulting&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[{"id":9,"name":"troubleshooting","slug":"troubleshooting"}],"tags_descriptions":{},"views":13,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"ValidationException: Unsupported Instance Type in EMR RunJobFlow","id":328,"title":"ValidationException: Unsupported Instance Type in EMR RunJobFlow","slug":"validationexception-unsupported-instance-type-in-emr-runjobflow","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-06-02T13:45:17.957Z","last_posted_at":"2026-06-02T13:45:18.030Z","bumped":true,"bumped_at":"2026-06-02T13:45:18.030Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"When submitting a pipeline on Airflow with an incorrect parameter for the Create EMR Cluster step, the logs will capture details as below: \nbotocore.exceptions.ClientError: An error occurred (ValidationException) when ca&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[{"id":9,"name":"troubleshooting","slug":"troubleshooting"}],"tags_descriptions":{},"views":12,"like_count":0,"has_summary":false,"last_poster_username":"neeraj","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":16,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"Troubleshooting Job Failures After Upgrading Python from 3.8 to 3.9","id":327,"title":"Troubleshooting Job Failures After Upgrading Python from 3.8 to 3.9","slug":"troubleshooting-job-failures-after-upgrading-python-from-3-8-to-3-9","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":"https://canada1.discourse-cdn.com/flex007/uploads/sparkflows/original/1X/b4a1daa72267dfaf1c1c16c5a17e7b9019cfb891.png","created_at":"2026-06-01T06:09:52.251Z","last_posted_at":"2026-06-01T06:09:52.323Z","bumped":true,"bumped_at":"2026-06-01T06:09:52.323Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Troubleshooting Job Failures After Upgrading Python from 3.8 to 3.9\nIf you have recently upgraded your environment from Python 3.8 to Python 3.9 and are experiencing job failures, the issue may be related to outdated pac&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":14,"like_count":0,"has_summary":false,"last_poster_username":"Daniel","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":13,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"How to configure S3 in the Sparkflows","id":325,"title":"How to configure S3 in the Sparkflows","slug":"how-to-configure-s3-in-the-sparkflows","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":"https://canada1.discourse-cdn.com/flex007/uploads/sparkflows/original/1X/934e192849adfab86768baba946419473ee03269.png","created_at":"2026-05-27T05:37:22.125Z","last_posted_at":"2026-05-27T05:37:22.187Z","bumped":true,"bumped_at":"2026-05-27T05:37:22.187Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Configuring S3 Storage Access for a Group\nThis configuration can be performed only by a user with administrative privileges. \nRequired Information\nBefore proceeding, ensure the following S3 configuration details are avai&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":15,"like_count":0,"has_summary":false,"last_poster_username":"Daniel","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":13,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"How to Get Pipeline Executions by Pipeline ID","id":320,"title":"How to Get Pipeline Executions by Pipeline ID","slug":"how-to-get-pipeline-executions-by-pipeline-id","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-04-27T10:16:55.499Z","last_posted_at":"2026-04-27T10:16:55.607Z","bumped":true,"bumped_at":"2026-04-27T10:16:55.607Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Get Pipeline Executions by Pipeline ID: A Complete Guide\nModern data platforms rely heavily on pipelines to orchestrate workflows, automate tasks, and ensure seamless data processing. Monitoring these pipelines is just a&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":18,"like_count":0,"has_summary":false,"last_poster_username":"Daniel","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":13,"primary_group_id":null,"flair_group_id":null}]},{"fancy_title":"How to import a workflow via API using cURL command","id":318,"title":"How to import a workflow via API using cURL command","slug":"how-to-import-a-workflow-via-api-using-curl-command","posts_count":1,"reply_count":0,"highest_post_number":1,"image_url":null,"created_at":"2026-04-24T09:08:42.170Z","last_posted_at":"2026-04-24T09:08:42.242Z","bumped":true,"bumped_at":"2026-04-24T09:08:42.242Z","archetype":"regular","unseen":false,"pinned":false,"unpinned":null,"excerpt":"Importing Workflows via API Using cURL\nManaging workflows across projects often requires a quick and reliable way to import multiple workflow definitions. This guide walks you through using a simple API endpoint with cUR&hellip;","visible":true,"closed":false,"archived":false,"bookmarked":null,"liked":null,"tags":[],"tags_descriptions":{},"views":16,"like_count":0,"has_summary":false,"last_poster_username":"Daniel","category_id":4,"op_like_count":0,"pinned_globally":false,"featured_link":null,"is_hot":false,"has_accepted_answer":false,"posters":[{"extras":"latest single","description":"Original Poster, Most Recent Poster","user_id":13,"primary_group_id":null,"flair_group_id":null}]}]}}