%------------------------------------------------------------------------------
% File : Vampire---5.0.1
% Problem : NUM893_1 : TPTP v9.3.1. Released v5.0.0.
% Transfm : none
% Format : tptp:raw
% Command : run_vampire /export/starexec/sandbox2/benchmark/theBenchmark.p 300 THM
% Computer : n026.cluster.edu
% Model : x86_64 x86_64
% CPU : Intel(R) Xeon(R) CPU E5-2620 v4 2.10GHz
% Memory : 8046.5625MB
% OS : Linux 6.8.0-71-generic
% CPULimit : 300s
% WCLimit : 300s
% DateTime : Tue Sep 29 12:17:23 PM UTC 2026
% Result : Theorem 4.20s 1.83s
% Output : Refutation 4.20s
% Verified :
% SZS Type : Refutation
% Derivation depth : 11
% Number of leaves : 1
% Syntax : Number of formulae : 12 ( 1 unt; 0 typ; 0 def)
% Number of atoms : 41 ( 40 equ)
% Maximal formula atoms : 6 ( 3 avg)
% Number of connectives : 48 ( 19 ~; 13 |; 11 &)
% ( 4 <=>; 0 =>; 0 <=; 1 <~>)
% Maximal formula depth : 8 ( 6 avg)
% Maximal term depth : 3 ( 1 avg)
% Number arithmetic : 105 ( 0 atm; 68 fun; 8 num; 29 var)
% Number of types : 1 ( 0 usr; 1 ari; 0 dat; 0 cdt)
% Number of type conns : 0 ( 0 >; 0 *; 0 +; 0 <<)
% Number of predicates : 2 ( 0 usr; 1 prp; 0-2 aty)
% Number of functors : 6 ( 2 usr; 3 con; 0-2 aty)
% Number of variables : 29 ( 17 !; 12 ?; 29 :)
% Comments :
%------------------------------------------------------------------------------
tff(func_def_5,type,
'$inst0': $int ).
tff(func_def_6,type,
'$inst1': $int ).
tff(f1,conjecture,
? [X2: $int,X1: $int,X0: $int] :
( ( ( $difference(X2,X0) = X1 )
& ( $difference(X2,X1) = X0 ) )
<=> ( $sum(X0,X1) = X2 ) ),
file('/export/starexec/sandbox2/benchmark/theBenchmark.p',sum_same_as_difference) ).
tff(f2,negated_conjecture,
~ ? [X2: $int,X1: $int,X0: $int] :
( ( ( $difference(X2,X0) = X1 )
& ( $difference(X2,X1) = X0 ) )
<=> ( $sum(X0,X1) = X2 ) ),
inference(negated_conjecture,[status(cth)],[f1]) ).
tff(f3,plain,
~ ? [X2: $int,X1: $int,X0: $int] :
( ( ( $sum(X2,$uminus(X0)) = X1 )
& ( $sum(X2,$uminus(X1)) = X0 ) )
<=> ( $sum(X0,X1) = X2 ) ),
inference(theory_normalization,[],[f2]) ).
tff(f16,plain,
~ ? [X1: $int,X2: $int,X0: $int] :
( ( ( $sum(X0,$uminus(X1)) = X2 )
& ( $sum(X0,$uminus(X2)) = X1 ) )
<=> ( $sum(X2,X1) = X0 ) ),
inference(rectify,[],[f3]) ).
tff(f17,plain,
! [X1: $int,X2: $int,X0: $int] :
( ( ( $sum(X0,$uminus(X1)) = X2 )
& ( $sum(X0,$uminus(X2)) = X1 ) )
<~> ( $sum(X2,X1) = X0 ) ),
inference(ennf_transformation,[],[f16]) ).
tff(f18,plain,
! [X1: $int,X2: $int,X0: $int] :
( ( ( $sum(X2,X1) != X0 )
| ( $sum(X0,$uminus(X1)) != X2 )
| ( $sum(X0,$uminus(X2)) != X1 ) )
& ( ( $sum(X2,X1) = X0 )
| ( ( $sum(X0,$uminus(X1)) = X2 )
& ( $sum(X0,$uminus(X2)) = X1 ) ) ) ),
inference(nnf_transformation,[],[f17]) ).
tff(f19,plain,
! [X1: $int,X2: $int,X0: $int] :
( ( ( $sum(X2,X1) != X0 )
| ( $sum(X0,$uminus(X1)) != X2 )
| ( $sum(X0,$uminus(X2)) != X1 ) )
& ( ( $sum(X2,X1) = X0 )
| ( ( $sum(X0,$uminus(X1)) = X2 )
& ( $sum(X0,$uminus(X2)) = X1 ) ) ) ),
inference(flattening,[],[f18]) ).
tff(f20,plain,
! [X0: $int,X1: $int,X2: $int] :
( ( ( $sum(X1,X0) != X2 )
| ( $sum(X2,$uminus(X0)) != X1 )
| ( $sum(X2,$uminus(X1)) != X0 ) )
& ( ( $sum(X1,X0) = X2 )
| ( ( $sum(X2,$uminus(X0)) = X1 )
& ( $sum(X2,$uminus(X1)) = X0 ) ) ) ),
inference(rectify,[],[f19]) ).
tff(f23,plain,
! [X2: $int,X0: $int,X1: $int] :
( ( $sum(X1,X0) != X2 )
| ( $sum(X2,$uminus(X0)) != X1 )
| ( $sum(X2,$uminus(X1)) != X0 ) ),
inference(cnf_transformation,[],[f20]) ).
tff(f24,plain,
! [X0: $int,X1: $int] :
( ( $sum($sum(X1,X0),$uminus(X1)) != X0 )
| ( $sum($sum(X1,X0),$uminus(X0)) != X1 ) ),
inference(equality_resolution,[],[f23]) ).
tff(f152,plain,
( ( 0 != $sum($sum(0,0),$uminus(0)) )
| ( 0 != $sum($sum(0,0),$uminus(0)) ) ),
inference(instantiation,[],[f24]) ).
tff(f153,plain,
$false,
inference(interpreted_simplification,[],[f152]) ).
%------------------------------------------------------------------------------
%----ORIGINAL SYSTEM OUTPUT
% 0.00/0.04 % Problem : NUM893_1 : TPTP v9.3.1. Released v5.0.0.
% 0.00/0.07 % Command : run_vampire /export/starexec/sandbox2/benchmark/theBenchmark.p 300 THM
% 0.18/0.42 % Computer : n026.cluster.edu
% 0.18/0.42 % Model : x86_64 x86_64
% 0.18/0.42 % CPU : Intel(R) Xeon(R) CPU E5-2620 v4 @ 2.10GHz
% 0.18/0.42 % Memory : 8046.5625MB
% 0.18/0.42 % OS : Linux 6.8.0-71-generic
% 0.18/0.42 % CPULimit : 300
% 0.18/0.42 % WCLimit : 300
% 0.18/0.42 % DateTime : Sun Sep 27 21:41:11 UTC 2026
% 0.18/0.43 % CPUTime :
% 0.18/0.43 Running run_vampire /export/starexec/sandbox2/benchmark/theBenchmark.p 300 THM
% 0.22/0.48 Running first-order theorem proving
% 0.22/0.48 Running: /export/starexec/sandbox2/solver/bin/vampire --input_syntax tptp --output_axiom_names on --mode casc -m 16384 --cores 7 -t 300 /export/starexec/sandbox2/benchmark/theBenchmark.p
% 4.20/1.83 % (3228610)Detected arithmetic, will pick strategies from an ALASCA-aware ARI schedule.
% 4.20/1.83 % (3228622)dis+10_3_slsqr=1,4:to=lpo:sil=128000:thi=strong:si=on:uwa=off:s2agt=20:slsqc=1:slsq=on:random_seed=2247109685:i=201:slsql=off:asg=cautious:rtra=on:gtg=all:ss=axioms:sgt=16_2999 on theBenchmark for (2999ds/201Mi)
% 4.20/1.83 % (3228622)First to succeed.
% 4.20/1.83 % (3228622)Solution written to "/export/starexec/sandbox2/tmp/vampire-proof-3228610"
% 4.20/1.83 % (3228626)lrs+10_1_tgt=ground:sas=z3:si=on:random_seed=3251770312:i=33:rtra=on_2999 on theBenchmark for (2999ds/33Mi)
% 4.20/1.83 % (3228621)dis+1002_1_to=kbo:sil=128000:tgt=ground:sas=z3:si=on:spb=units:tha=off:random_seed=1916277748:i=307:kws=precedence:nm=0:rtra=on_2999 on theBenchmark for (2999ds/307Mi)
% 4.20/1.83 % (3228623)lrs+1002_4:1_to=lpo:sil=64000:si=on:br=off:random_seed=2089774171:s2a=on:i=7:rtra=on:inst=on_2999 on theBenchmark for (2999ds/7Mi)
% 4.20/1.83 % (3228624)dis+21_64_to=kbo:sil=128000:si=on:sp=weighted_frequency:uwa=alasca_can_abstract:random_seed=1517235080:i=4:rtra=on_2999 on theBenchmark for (2999ds/4Mi)
% 4.20/1.83 % (3228620)dis+1002_16:1_to=lpo:sil=64000:sas=z3:si=on:norm_ineq=on:gve=force:uwa=one_side_constant:random_seed=2668251540:i=12:doe=on:rtra=on:gtg=exists_top:ss=axioms_2999 on theBenchmark for (2999ds/12Mi)
% 4.20/1.83 % (3228623)Also succeeded, but the first one will report.
% 4.20/1.83 % (3228624)Instruction limit reached!
% 4.20/1.83 % (3228624)------------------------------
% 4.20/1.83 % (3228624)Version: Vampire 5.0.1 (Release build, commit ea8961452 on 2026-07-16 15:14:34 +0200)
% 4.20/1.83 % (3228624)Linked with Z3 4.14.0.0 3c47fd96cf5645d0c42b2c819d9e9a84380aa721 z3-4.8.4-9178-g3c47fd96c
% 4.20/1.83 % (3228624)CaDiCaL version: 2.1.3
% 4.20/1.83 % (3228624)Termination reason: Instruction limit
% 4.20/1.83 % (3228624)Termination phase: Saturation
% 4.20/1.83 % (3228624)Time elapsed: 0.005 s
% 4.20/1.83 % (3228624)Peak memory usage: 88 MB
% 4.20/1.83 % (3228624)Instructions burned: 4 (million)
% 4.20/1.83 % (3228625)lrs+10_1_to=lpo:sas=z3:si=on:tha=off:random_seed=1722270541:i=46:rtra=on_2999 on theBenchmark for (2999ds/46Mi)
% 4.20/1.83 % (3228626)Also succeeded, but the first one will report.
% 4.20/1.83 % (3228620)Refutation not found, incomplete strategy
% 4.20/1.83 % (3228620)------------------------------
% 4.20/1.83 % (3228620)Version: Vampire 5.0.1 (Release build, commit ea8961452 on 2026-07-16 15:14:34 +0200)
% 4.20/1.83 % (3228620)Linked with Z3 4.14.0.0 3c47fd96cf5645d0c42b2c819d9e9a84380aa721 z3-4.8.4-9178-g3c47fd96c
% 4.20/1.83 % (3228620)CaDiCaL version: 2.1.3
% 4.20/1.83 % (3228620)Termination reason: Refutation not found, incomplete strategy
% 4.20/1.83 % (3228620)Time elapsed: 0.039 s
% 4.20/1.83 % (3228620)Peak memory usage: 112 MB
% 4.20/1.83 % (3228620)Instructions burned: 6 (million)
% 4.20/1.83 % (3228625)Instruction limit reached!
% 4.20/1.83 % (3228625)------------------------------
% 4.20/1.83 % (3228625)Version: Vampire 5.0.1 (Release build, commit ea8961452 on 2026-07-16 15:14:34 +0200)
% 4.20/1.83 % (3228625)Linked with Z3 4.14.0.0 3c47fd96cf5645d0c42b2c819d9e9a84380aa721 z3-4.8.4-9178-g3c47fd96c
% 4.20/1.83 % (3228625)CaDiCaL version: 2.1.3
% 4.20/1.83 % (3228625)Termination reason: Instruction limit
% 4.20/1.83 % (3228625)Termination phase: Saturation
% 4.20/1.83 % (3228625)Time elapsed: 0.071 s
% 4.20/1.83 % (3228625)Peak memory usage: 112 MB
% 4.20/1.83 % (3228625)Instructions burned: 46 (million)
% 4.20/1.83 % (3228636)dis+11_3_anc=none:drc=ordering:si=on:urr=ec_only:bce=on:tha=off:sac=on:random_seed=902871292:st=5:i=14:sd=10:rtra=on:ss=axioms:rawr=on_2997 on theBenchmark for (2997ds/14Mi)
% 4.20/1.83 % (3228636)Instruction limit reached!
% 4.20/1.83 % (3228636)------------------------------
% 4.20/1.83 % (3228636)Version: Vampire 5.0.1 (Release build, commit ea8961452 on 2026-07-16 15:14:34 +0200)
% 4.20/1.83 % (3228636)Linked with Z3 4.14.0.0 3c47fd96cf5645d0c42b2c819d9e9a84380aa721 z3-4.8.4-9178-g3c47fd96c
% 4.20/1.83 % (3228636)CaDiCaL version: 2.1.3
% 4.20/1.83 % (3228636)Termination reason: Instruction limit
% 4.20/1.83 % (3228636)Termination phase: Saturation
% 4.20/1.83 % (3228636)Time elapsed: 0.010 s
% 4.20/1.83 % (3228636)Peak memory usage: 88 MB
% 4.20/1.83 % (3228636)Instructions burned: 14 (million)
% 4.20/1.83 % (3228622)Refutation found. Thanks to Tanya!
% 4.20/1.83 % SZS status Theorem for theBenchmark
% 4.20/1.83 % SZS output start Proof for theBenchmark
% See solution above
% 4.20/1.83 % (3228622)------------------------------
% 4.20/1.83 % (3228622)Version: Vampire 5.0.1 (Release build, commit ea8961452 on 2026-07-16 15:14:34 +0200)
% 4.20/1.83 % (3228622)Linked with Z3 4.14.0.0 3c47fd96cf5645d0c42b2c819d9e9a84380aa721 z3-4.8.4-9178-g3c47fd96c
% 4.20/1.83 % (3228622)CaDiCaL version: 2.1.3
% 4.20/1.83 % (3228622)Termination reason: Refutation
% 4.20/1.83 % (3228622)Time elapsed: 0.028 s
% 4.20/1.83 % (3228622)Peak memory usage: 116 MB
% 4.20/1.83 % (3228622)Instructions burned: 17 (million)
% 4.20/1.83 % (3228622)------------------------------
% 4.20/1.83 % (3228622)------------------------------
% 4.20/1.83 % (3228610)Success in time 0.472 s
% 4.20/1.83 % Vampire exiting
%------------------------------------------------------------------------------