a3e51cc990
Creates a top level directory script (as `build/mvn`) to automatically download zinc and the specific version of scala used to easily build spark. This will also download and install maven if the user doesn't already have it and all packages are hosted under the `build/` directory. Tested on both Linux and OSX OS's and both work. All commands pass through to the maven binary so it acts exactly as a traditional maven call would. Author: Brennon York <brennon.york@capitalone.com> Closes #3707 from brennonyork/SPARK-4501 and squashes the following commits: 0e5a0e4 [Brennon York] minor incorrect doc verbage (with -> this) 9b79e38 [Brennon York] fixed merge conflicts with dev/run-tests, properly quoted args in sbt/sbt, fixed bug where relative paths would fail if passed in from build/mvn d2d41b6 [Brennon York] added blurb about leverging zinc with build/mvn b979c58 [Brennon York] updated the merge conflict c5634de [Brennon York] updated documentation to overview build/mvn, updated all points where sbt/sbt was referenced with build/sbt b8437ba [Brennon York] set progress bars for curl and wget when not run on jenkins, no progress bar when run on jenkins, moved sbt script to build/sbt, wrote stub and warning under sbt/sbt which calls build/sbt, modified build/sbt to use the correct directory, fixed bug in build/sbt-launch-lib.bash to correctly pull the sbt version be11317 [Brennon York] added switch to silence download progress only if AMPLAB_JENKINS is set 28d0a99 [Brennon York] updated to remove the python dependency, uses grep instead 7e785a6 [Brennon York] added silent and quiet flags to curl and wget respectively, added single echo output to denote start of a download if download is needed 14a5da0 [Brennon York] removed unnecessary zinc output on startup 1af4a94 [Brennon York] fixed bug with uppercase vs lowercase variable 3e8b9b3 [Brennon York] updated to properly only restart zinc if it was freshly installed a680d12 [Brennon York] Added comments to functions and tested various mvn calls bb8cc9d [Brennon York] removed package files ef017e6 [Brennon York] removed OS complexities, setup generic install_app call, removed extra file complexities, removed help, removed forced install (defaults now), removed double-dash from cli 07bf018 [Brennon York] Updated to specifically handle pulling down the correct scala version f914dea [Brennon York] Beginning final portions of localized scala home 69c4e44 [Brennon York] working linux and osx installers for purely local mvn build 4a1609c [Brennon York] finalizing working linux install for maven to local ./build/apache-maven folder cbfcc68 [Brennon York] Changed the default sbt/sbt to build/sbt and added a build/mvn which will automatically download, install, and execute maven with zinc for easier build capability
221 lines
7.5 KiB
Bash
Executable file
221 lines
7.5 KiB
Bash
Executable file
#!/usr/bin/env bash
|
|
|
|
#
|
|
# Licensed to the Apache Software Foundation (ASF) under one or more
|
|
# contributor license agreements. See the NOTICE file distributed with
|
|
# this work for additional information regarding copyright ownership.
|
|
# The ASF licenses this file to You under the Apache License, Version 2.0
|
|
# (the "License"); you may not use this file except in compliance with
|
|
# the License. You may obtain a copy of the License at
|
|
#
|
|
# http://www.apache.org/licenses/LICENSE-2.0
|
|
#
|
|
# Unless required by applicable law or agreed to in writing, software
|
|
# distributed under the License is distributed on an "AS IS" BASIS,
|
|
# WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
|
|
# See the License for the specific language governing permissions and
|
|
# limitations under the License.
|
|
#
|
|
|
|
# Go to the Spark project root directory
|
|
FWDIR="$(cd "`dirname $0`"/..; pwd)"
|
|
cd "$FWDIR"
|
|
|
|
# Remove work directory
|
|
rm -rf ./work
|
|
|
|
source "$FWDIR/dev/run-tests-codes.sh"
|
|
|
|
CURRENT_BLOCK=$BLOCK_GENERAL
|
|
|
|
function handle_error () {
|
|
echo "[error] Got a return code of $? on line $1 of the run-tests script."
|
|
exit $CURRENT_BLOCK
|
|
}
|
|
|
|
|
|
# Build against the right verison of Hadoop.
|
|
{
|
|
if [ -n "$AMPLAB_JENKINS_BUILD_PROFILE" ]; then
|
|
if [ "$AMPLAB_JENKINS_BUILD_PROFILE" = "hadoop1.0" ]; then
|
|
export SBT_MAVEN_PROFILES_ARGS="-Dhadoop.version=1.0.4"
|
|
elif [ "$AMPLAB_JENKINS_BUILD_PROFILE" = "hadoop2.0" ]; then
|
|
export SBT_MAVEN_PROFILES_ARGS="-Dhadoop.version=2.0.0-mr1-cdh4.1.1"
|
|
elif [ "$AMPLAB_JENKINS_BUILD_PROFILE" = "hadoop2.2" ]; then
|
|
export SBT_MAVEN_PROFILES_ARGS="-Pyarn -Phadoop-2.2 -Dhadoop.version=2.2.0"
|
|
elif [ "$AMPLAB_JENKINS_BUILD_PROFILE" = "hadoop2.3" ]; then
|
|
export SBT_MAVEN_PROFILES_ARGS="-Pyarn -Phadoop-2.3 -Dhadoop.version=2.3.0"
|
|
fi
|
|
fi
|
|
|
|
if [ -z "$SBT_MAVEN_PROFILES_ARGS" ]; then
|
|
export SBT_MAVEN_PROFILES_ARGS="-Pyarn -Phadoop-2.3 -Dhadoop.version=2.3.0"
|
|
fi
|
|
}
|
|
|
|
export SBT_MAVEN_PROFILES_ARGS="$SBT_MAVEN_PROFILES_ARGS -Pkinesis-asl"
|
|
|
|
# Determine Java path and version.
|
|
{
|
|
if test -x "$JAVA_HOME/bin/java"; then
|
|
declare java_cmd="$JAVA_HOME/bin/java"
|
|
else
|
|
declare java_cmd=java
|
|
fi
|
|
|
|
# We can't use sed -r -e due to OS X / BSD compatibility; hence, all the parentheses.
|
|
JAVA_VERSION=$(
|
|
$java_cmd -version 2>&1 \
|
|
| grep -e "^java version" --max-count=1 \
|
|
| sed "s/java version \"\(.*\)\.\(.*\)\.\(.*\)\"/\1\2/"
|
|
)
|
|
|
|
if [ "$JAVA_VERSION" -lt 18 ]; then
|
|
echo "[warn] Java 8 tests will not run because JDK version is < 1.8."
|
|
fi
|
|
}
|
|
|
|
# Only run Hive tests if there are sql changes.
|
|
# Partial solution for SPARK-1455.
|
|
if [ -n "$AMPLAB_JENKINS" ]; then
|
|
git fetch origin master:master
|
|
|
|
sql_diffs=$(
|
|
git diff --name-only master \
|
|
| grep -e "^sql/" -e "^bin/spark-sql" -e "^sbin/start-thriftserver.sh"
|
|
)
|
|
|
|
non_sql_diffs=$(
|
|
git diff --name-only master \
|
|
| grep -v -e "^sql/" -e "^bin/spark-sql" -e "^sbin/start-thriftserver.sh"
|
|
)
|
|
|
|
if [ -n "$sql_diffs" ]; then
|
|
echo "[info] Detected changes in SQL. Will run Hive test suite."
|
|
_RUN_SQL_TESTS=true
|
|
|
|
if [ -z "$non_sql_diffs" ]; then
|
|
echo "[info] Detected no changes except in SQL. Will only run SQL tests."
|
|
_SQL_TESTS_ONLY=true
|
|
fi
|
|
fi
|
|
fi
|
|
|
|
set -o pipefail
|
|
trap 'handle_error $LINENO' ERR
|
|
|
|
echo ""
|
|
echo "========================================================================="
|
|
echo "Running Apache RAT checks"
|
|
echo "========================================================================="
|
|
|
|
CURRENT_BLOCK=$BLOCK_RAT
|
|
|
|
./dev/check-license
|
|
|
|
echo ""
|
|
echo "========================================================================="
|
|
echo "Running Scala style checks"
|
|
echo "========================================================================="
|
|
|
|
CURRENT_BLOCK=$BLOCK_SCALA_STYLE
|
|
|
|
./dev/lint-scala
|
|
|
|
echo ""
|
|
echo "========================================================================="
|
|
echo "Running Python style checks"
|
|
echo "========================================================================="
|
|
|
|
CURRENT_BLOCK=$BLOCK_PYTHON_STYLE
|
|
|
|
./dev/lint-python
|
|
|
|
echo ""
|
|
echo "========================================================================="
|
|
echo "Building Spark"
|
|
echo "========================================================================="
|
|
|
|
CURRENT_BLOCK=$BLOCK_BUILD
|
|
|
|
{
|
|
|
|
# NOTE: echo "q" is needed because sbt on encountering a build file with failure
|
|
# (either resolution or compilation) prompts the user for input either q, r, etc
|
|
# to quit or retry. This echo is there to make it not block.
|
|
# NOTE: Do not quote $BUILD_MVN_PROFILE_ARGS or else it will be interpreted as a
|
|
# single argument!
|
|
# QUESTION: Why doesn't 'yes "q"' work?
|
|
# QUESTION: Why doesn't 'grep -v -e "^\[info\] Resolving"' work?
|
|
# First build with Hive 0.12.0 to ensure patches do not break the Hive 0.12.0 build
|
|
HIVE_12_BUILD_ARGS="$SBT_MAVEN_PROFILES_ARGS -Phive -Phive-thriftserver -Phive-0.12.0"
|
|
echo "[info] Compile with Hive 0.12.0"
|
|
echo -e "q\n" \
|
|
| build/sbt $HIVE_12_BUILD_ARGS clean hive/compile hive-thriftserver/compile \
|
|
| grep -v -e "info.*Resolving" -e "warn.*Merging" -e "info.*Including"
|
|
|
|
# Then build with default Hive version (0.13.1) because tests are based on this version
|
|
echo "[info] Compile with Hive 0.13.1"
|
|
rm -rf lib_managed
|
|
echo "[info] Building Spark with these arguments: $SBT_MAVEN_PROFILES_ARGS"\
|
|
" -Phive -Phive-thriftserver"
|
|
echo -e "q\n" \
|
|
| build/sbt $SBT_MAVEN_PROFILES_ARGS -Phive -Phive-thriftserver package assembly/assembly \
|
|
| grep -v -e "info.*Resolving" -e "warn.*Merging" -e "info.*Including"
|
|
}
|
|
|
|
echo ""
|
|
echo "========================================================================="
|
|
echo "Running Spark unit tests"
|
|
echo "========================================================================="
|
|
|
|
CURRENT_BLOCK=$BLOCK_SPARK_UNIT_TESTS
|
|
|
|
{
|
|
# If the Spark SQL tests are enabled, run the tests with the Hive profiles enabled.
|
|
# This must be a single argument, as it is.
|
|
if [ -n "$_RUN_SQL_TESTS" ]; then
|
|
SBT_MAVEN_PROFILES_ARGS="$SBT_MAVEN_PROFILES_ARGS -Phive -Phive-thriftserver"
|
|
fi
|
|
|
|
if [ -n "$_SQL_TESTS_ONLY" ]; then
|
|
# This must be an array of individual arguments. Otherwise, having one long string
|
|
# will be interpreted as a single test, which doesn't work.
|
|
SBT_MAVEN_TEST_ARGS=("catalyst/test" "sql/test" "hive/test" "mllib/test")
|
|
else
|
|
SBT_MAVEN_TEST_ARGS=("test")
|
|
fi
|
|
|
|
echo "[info] Running Spark tests with these arguments: $SBT_MAVEN_PROFILES_ARGS ${SBT_MAVEN_TEST_ARGS[@]}"
|
|
|
|
# NOTE: echo "q" is needed because sbt on encountering a build file with failure
|
|
# (either resolution or compilation) prompts the user for input either q, r, etc
|
|
# to quit or retry. This echo is there to make it not block.
|
|
# NOTE: Do not quote $SBT_MAVEN_PROFILES_ARGS or else it will be interpreted as a
|
|
# single argument!
|
|
# "${SBT_MAVEN_TEST_ARGS[@]}" is cool because it's an array.
|
|
# QUESTION: Why doesn't 'yes "q"' work?
|
|
# QUESTION: Why doesn't 'grep -v -e "^\[info\] Resolving"' work?
|
|
echo -e "q\n" \
|
|
| build/sbt $SBT_MAVEN_PROFILES_ARGS "${SBT_MAVEN_TEST_ARGS[@]}" \
|
|
| grep -v -e "info.*Resolving" -e "warn.*Merging" -e "info.*Including"
|
|
}
|
|
|
|
echo ""
|
|
echo "========================================================================="
|
|
echo "Running PySpark tests"
|
|
echo "========================================================================="
|
|
|
|
CURRENT_BLOCK=$BLOCK_PYSPARK_UNIT_TESTS
|
|
|
|
./python/run-tests
|
|
|
|
echo ""
|
|
echo "========================================================================="
|
|
echo "Detecting binary incompatibilities with MiMa"
|
|
echo "========================================================================="
|
|
|
|
CURRENT_BLOCK=$BLOCK_MIMA
|
|
|
|
./dev/mima
|